tickmarkr 1.67.0 → 1.69.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (46) hide show
  1. package/dist/adapters/kimi.d.ts +25 -1
  2. package/dist/adapters/kimi.js +80 -0
  3. package/dist/adapters/types.d.ts +6 -0
  4. package/dist/cli/commands/eval.d.ts +4 -0
  5. package/dist/cli/commands/eval.js +26 -0
  6. package/dist/cli/commands/status.d.ts +20 -0
  7. package/dist/cli/commands/status.js +9 -9
  8. package/dist/cli/index.d.ts +1 -1
  9. package/dist/cli/index.js +3 -1
  10. package/dist/eval/canary.d.ts +46 -0
  11. package/dist/eval/canary.js +113 -0
  12. package/dist/eval/dispatch.d.ts +31 -0
  13. package/dist/eval/dispatch.js +207 -0
  14. package/dist/eval/fixtures.d.ts +22 -0
  15. package/dist/eval/fixtures.js +85 -0
  16. package/dist/eval/report.d.ts +35 -0
  17. package/dist/eval/report.js +82 -0
  18. package/dist/eval/selfcheck.d.ts +22 -0
  19. package/dist/eval/selfcheck.js +177 -0
  20. package/dist/run/daemon.js +130 -83
  21. package/dist/run/git.d.ts +2 -0
  22. package/dist/run/git.js +9 -0
  23. package/dist/run/interactive-seed.d.ts +15 -0
  24. package/dist/run/interactive-seed.js +25 -0
  25. package/dist/tui/app.d.ts +3 -0
  26. package/dist/tui/app.js +14 -2
  27. package/dist/tui/views/consult-dossier.d.ts +49 -0
  28. package/dist/tui/views/consult-dossier.js +169 -0
  29. package/dist/tui/views/runs-view.d.ts +70 -0
  30. package/dist/tui/views/runs-view.js +387 -0
  31. package/fixtures/eval/canary/solution/a.txt +1 -0
  32. package/fixtures/eval/canary/spec.md +8 -0
  33. package/fixtures/eval/canary/start/a.txt +1 -0
  34. package/fixtures/eval/sample/solution/a.txt +1 -0
  35. package/fixtures/eval/sample/spec.md +8 -0
  36. package/fixtures/eval/sample/start/a.txt +1 -0
  37. package/fixtures/gsd-sample/07-live-check/07-01-PLAN.md +42 -0
  38. package/fixtures/gsd-sample/07-live-check/07-02-PLAN.md +21 -0
  39. package/fixtures/gsd-sample/07-live-check/07-03-PLAN.md +18 -0
  40. package/fixtures/gsd-sample/07-live-check/07-03-SUMMARY.md +1 -0
  41. package/fixtures/missing-mandatory-gate.native.md +10 -0
  42. package/fixtures/sample-pin.prd.md +18 -0
  43. package/fixtures/sample.native.md +35 -0
  44. package/fixtures/sample.prd.md +22 -0
  45. package/fixtures/speckit-sample/tasks.md +20 -0
  46. package/package.json +3 -2
@@ -0,0 +1,70 @@
1
+ import { type RunGraph, type Task, type TaskStatus } from "../../graph/schema.js";
2
+ import { type JournalEvent, type TelemetryRow, type WorkerResultCause } from "../../run/journal.js";
3
+ import { type GateState } from "../../cli/commands/status.js";
4
+ import type { View } from "../app.js";
5
+ export type AttemptRecord = {
6
+ attempt: number;
7
+ channel: string;
8
+ outcome: "clean" | "failed" | "in-flight" | "aborted";
9
+ /** typed worker-result failure reason when the attempt did not finish cleanly */
10
+ cause?: WorkerResultCause;
11
+ /** worker-result summary or consult-verdict reason, when present */
12
+ note?: string;
13
+ };
14
+ export type RunsTask = {
15
+ task: Task;
16
+ status: TaskStatus;
17
+ states: GateState[];
18
+ channel: string;
19
+ ctx?: number;
20
+ activity?: string;
21
+ attempts: AttemptRecord[];
22
+ };
23
+ /** Everything the Runs cockpit renders — loaded by the caller, never by the render path. */
24
+ export type RunsViewData = {
25
+ runId?: string;
26
+ events: JournalEvent[];
27
+ graph: RunGraph;
28
+ /** Whether the journal is comparable with the loaded graph. Defaults to true. */
29
+ comparable?: boolean;
30
+ /** Per-tier per-task pricing from config, rendered ONLY through the shared costSignal formatter. */
31
+ pricing?: Record<string, number>;
32
+ /** This run's observed usage rows (telemetry.jsonl), folded per channel by the cost ticker. */
33
+ telemetry?: TelemetryRow[];
34
+ };
35
+ /** Options for creating a Runs view. Backward-compatible: plain `createRunsView(data)` still works. */
36
+ export type RunsViewOptions = {
37
+ data?: RunsViewData;
38
+ repoRoot?: string;
39
+ /** Notify the caller of a single-line notice change (confirmation, refusal, progress, result). */
40
+ onNotice?: (message: string | null) => void;
41
+ /** Refresh injected data after a mutation. Called after a successful approval. */
42
+ reload?: () => RunsViewData | Promise<RunsViewData>;
43
+ };
44
+ export type RunsView = View & {
45
+ /** Handle a decoded key name ("up" | "down" | "a" | "y"). */
46
+ key(name: string): void;
47
+ /** Index of the cursor in the task list. */
48
+ cursor: number;
49
+ /** The currently selected task, if any. */
50
+ selectedTask(): RunsTask | undefined;
51
+ /** Promise of an in-flight approval, exposed for tests. */
52
+ approval?: Promise<void>;
53
+ };
54
+ /** Build the injected raw data into task rows. Pure function, no filesystem access. */
55
+ export declare function foldRunsTasks(data: RunsViewData): RunsTask[];
56
+ export type CostTickerRow = {
57
+ /** adapter:model */
58
+ key: string;
59
+ /** verbatim output of the shared cost-signal formatter */
60
+ signal: string;
61
+ /** distinct tasks dispatched on this channel */
62
+ tasks: number;
63
+ /** summed observed token usage for this channel; undefined when nothing was metered */
64
+ tokens?: number;
65
+ };
66
+ /** Sum observed token usage per channel for the loaded run. Dispatch events name the channel's
67
+ * full assignment (billing channel + tier); telemetry rows contribute the metered tokens. A
68
+ * telemetry-only channel degrades costSignal to "api metered" — never an invented price. */
69
+ export declare function foldCostTicker(data: RunsViewData): CostTickerRow[];
70
+ export declare function createRunsView(data?: RunsViewData, opts?: Omit<RunsViewOptions, "data">): RunsView;
@@ -0,0 +1,387 @@
1
+ // T1 (v1.68): Runs cockpit — journal timeline, per-task gate ladder, attempt history.
2
+ // Pure render over INJECTED data: this module never touches the filesystem. The caller hands in
3
+ // a journal event list and a compiled graph; the view folds them into task cards and a dispatch-
4
+ // ordered attempt history itself. The task card language and gate ladder reuse the existing status
5
+ // command helpers verbatim rather than reimplementing them.
6
+ import { GLYPHS, dim, fail, legend, ok, statusRow, warn } from "../../brand.js";
7
+ import { channelKey } from "../../adapters/types.js";
8
+ import { GATE_NAMES, } from "../../graph/schema.js";
9
+ import { foldActivity } from "../../run/activity.js";
10
+ import { formatJournalNarration } from "../../run/journal.js";
11
+ import { costSignal } from "../../cli/commands/fleet-picker.js";
12
+ import { approve } from "../../cli/commands/approve.js";
13
+ import { gateChain, gateStates, defaultGateStates, humanGateSuffix, shortGoal, failedGates, } from "../../cli/commands/status.js";
14
+ import { renderDossierPlaceholder } from "./consult-dossier.js";
15
+ const channelOf = (assignment) => {
16
+ const a = assignment;
17
+ return typeof a?.adapter === "string" && typeof a.model === "string" ? `${a.adapter}:${a.model}` : "unknown channel";
18
+ };
19
+ const taskVerdict = (st) => st === "done" ? "pass" : st === "failed" ? "fail" : st === "human" ? "warn" : "neutral";
20
+ /** Local replay of task statuses from events only — mirrors Journal.replayStatuses but keeps the
21
+ * view filesystem-free. */
22
+ function replayStatuses(events) {
23
+ const s = new Map();
24
+ for (const e of events) {
25
+ if (!e.taskId)
26
+ continue;
27
+ if (e.event === "task-dispatch")
28
+ s.set(e.taskId, "running");
29
+ else if (e.event === "task-done")
30
+ s.set(e.taskId, "done");
31
+ else if (e.event === "task-failed")
32
+ s.set(e.taskId, "failed");
33
+ else if (e.event === "task-human")
34
+ s.set(e.taskId, "human");
35
+ else if (e.event === "task-approved")
36
+ s.set(e.taskId, "pending");
37
+ }
38
+ for (const [id, st] of s)
39
+ if (st === "running")
40
+ s.set(id, "pending");
41
+ return s;
42
+ }
43
+ /** Fold one task's dispatch-ordered attempt history. Each task-dispatch opens an attempt; a
44
+ * worker-result sets the attempt's outcome; task-done/task-failed/task-approved finish it. A new
45
+ * dispatch while the previous attempt is still in-flight marks it aborted. */
46
+ function foldAttempts(events, taskId) {
47
+ const attempts = [];
48
+ let open;
49
+ const finish = () => {
50
+ if (open) {
51
+ attempts.push(open);
52
+ open = undefined;
53
+ }
54
+ };
55
+ for (const e of events) {
56
+ if (e.taskId !== taskId)
57
+ continue;
58
+ if (e.event === "task-dispatch" || e.event === "escalation") {
59
+ if (open) {
60
+ if (open.outcome === "in-flight")
61
+ open.outcome = "aborted";
62
+ finish();
63
+ }
64
+ const attempt = (Number.isInteger(e.data.attempt) ? e.data.attempt : 0) + 1;
65
+ open = { attempt, channel: channelOf(e.data.assignment), outcome: "in-flight" };
66
+ }
67
+ else if (e.event === "worker-result" && open) {
68
+ const ok = e.data.ok === true;
69
+ const finished = e.data.finished === true;
70
+ if (ok && finished) {
71
+ open.outcome = "clean";
72
+ }
73
+ else {
74
+ open.outcome = "failed";
75
+ open.cause = e.data.cause ?? undefined;
76
+ open.note = typeof e.data.summary === "string" ? e.data.summary : undefined;
77
+ }
78
+ }
79
+ else if (e.event === "task-done" || e.event === "task-approved") {
80
+ if (open && open.outcome === "in-flight")
81
+ open.outcome = "clean";
82
+ finish();
83
+ }
84
+ else if (e.event === "task-failed") {
85
+ if (open && open.outcome === "in-flight")
86
+ open.outcome = "failed";
87
+ finish();
88
+ }
89
+ }
90
+ finish();
91
+ return attempts;
92
+ }
93
+ /** Build the injected raw data into task rows. Pure function, no filesystem access. */
94
+ export function foldRunsTasks(data) {
95
+ const { events, graph, comparable = true } = data;
96
+ const statuses = replayStatuses(events);
97
+ const assignments = new Map();
98
+ const contexts = new Map();
99
+ for (const e of events) {
100
+ if (e.event === "task-dispatch" && e.taskId) {
101
+ assignments.set(e.taskId, channelOf(e.data.assignment));
102
+ }
103
+ if (e.event === "context-sample" && e.taskId && typeof e.data.tokens === "number" && Number.isFinite(e.data.tokens)) {
104
+ contexts.set(e.taskId, e.data.tokens);
105
+ }
106
+ }
107
+ const activityTasks = graph.tasks.map((t) => ({
108
+ id: t.id,
109
+ gates: t.gates,
110
+ deps: t.deps,
111
+ status: statuses.get(t.id) ?? t.status,
112
+ }));
113
+ const activity = foldActivity(comparable ? events : [], activityTasks);
114
+ return graph.tasks.map((t) => {
115
+ const status = statuses.get(t.id) ?? t.status;
116
+ const states = comparable ? gateStates(t, events) : defaultGateStates(t);
117
+ const channel = assignments.get(t.id) ?? "-";
118
+ return {
119
+ task: t,
120
+ status,
121
+ states,
122
+ channel,
123
+ ctx: contexts.get(t.id),
124
+ activity: activity.cells.get(t.id),
125
+ attempts: foldAttempts(events, t.id),
126
+ };
127
+ });
128
+ }
129
+ const CARD_LINES = 3; // 2 content lines + 1 blank separator
130
+ const usageTotal = (u) => u.input + u.output + (u.cacheRead ?? 0) + (u.cacheWrite ?? 0) + (u.reasoning ?? 0);
131
+ /** Sum observed token usage per channel for the loaded run. Dispatch events name the channel's
132
+ * full assignment (billing channel + tier); telemetry rows contribute the metered tokens. A
133
+ * telemetry-only channel degrades costSignal to "api metered" — never an invented price. */
134
+ export function foldCostTicker(data) {
135
+ const pricing = data.pricing ?? {};
136
+ const rows = new Map();
137
+ const rowFor = (key) => {
138
+ let r = rows.get(key);
139
+ if (!r) {
140
+ r = { tasks: new Set(), tokens: 0, metered: false };
141
+ rows.set(key, r);
142
+ }
143
+ return r;
144
+ };
145
+ for (const e of data.events) {
146
+ if (e.event !== "task-dispatch" || !e.taskId)
147
+ continue;
148
+ const a = e.data.assignment;
149
+ if (!a || typeof a.adapter !== "string" || typeof a.model !== "string")
150
+ continue;
151
+ const r = rowFor(channelKey(a));
152
+ r.tasks.add(e.taskId);
153
+ if ((a.channel === "sub" || a.channel === "api") && typeof a.tier === "string")
154
+ r.assignment = a;
155
+ }
156
+ for (const t of data.telemetry ?? []) {
157
+ const r = rowFor(channelKey(t));
158
+ if (t.tokens) {
159
+ r.tokens += usageTotal(t.tokens);
160
+ r.metered = true;
161
+ }
162
+ if (!r.assignment && (t.channel === "sub" || t.channel === "api")) {
163
+ r.assignment = { adapter: t.adapter, model: t.model, channel: t.channel, tier: "" };
164
+ }
165
+ }
166
+ return [...rows.entries()].map(([key, r]) => ({
167
+ key,
168
+ signal: r.assignment ? costSignal(r.assignment, pricing) : "channel unknown",
169
+ tasks: r.tasks.size,
170
+ tokens: r.metered ? r.tokens : undefined,
171
+ }));
172
+ }
173
+ /** Last tip-verify event wins; none recorded yet ⇒ pending — never a false pass or fail. */
174
+ function tipVerifyState(events) {
175
+ for (let i = events.length - 1; i >= 0; i--) {
176
+ const e = events[i];
177
+ if (e.event === "tip-verify" || e.event === "tip-verify-failed") {
178
+ return {
179
+ state: e.event === "tip-verify" ? "passed" : "failed",
180
+ gate: typeof e.data.gate === "string" ? e.data.gate : undefined,
181
+ };
182
+ }
183
+ }
184
+ return { state: "pending" };
185
+ }
186
+ function tipVerifyLine(events) {
187
+ const tip = tipVerifyState(events);
188
+ const word = tip.state === "passed" ? ok("passed") : tip.state === "failed" ? fail("failed") : dim("pending");
189
+ return legend(` tip-verify: ${tip.gate ? `${tip.gate} ` : ""}`) + word;
190
+ }
191
+ /** Cost values stay plain (CLI-DESIGN "data plain"): costSignal's output verbatim, never colorized. */
192
+ function costTickerPanel(data) {
193
+ const rows = foldCostTicker(data);
194
+ const lines = ["", dim("── cost ticker ──")];
195
+ if (rows.length === 0) {
196
+ lines.push(" no channel usage observed yet");
197
+ return lines;
198
+ }
199
+ const keyW = Math.max(...rows.map((r) => r.key.length));
200
+ for (const r of rows) {
201
+ const parts = [r.signal, `${r.tasks} task${r.tasks === 1 ? "" : "s"}`];
202
+ if (r.tokens !== undefined)
203
+ parts.push(`${r.tokens} tokens`);
204
+ lines.push(` ${r.key.padEnd(keyW)} ${parts.join(" · ")}`);
205
+ }
206
+ return lines;
207
+ }
208
+ export function createRunsView(data, opts) {
209
+ let currentData = data;
210
+ let tasks = currentData ? foldRunsTasks(currentData) : [];
211
+ let cursor = 0;
212
+ let confirming = null;
213
+ let notice = null;
214
+ let busy = false;
215
+ let approvalPromise;
216
+ const repoRoot = opts?.repoRoot;
217
+ const onNotice = opts?.onNotice ?? (() => { });
218
+ const reload = opts?.reload;
219
+ const setNotice = (message) => {
220
+ notice = message;
221
+ onNotice(message);
222
+ };
223
+ const refresh = async () => {
224
+ if (!reload)
225
+ return;
226
+ currentData = await reload();
227
+ tasks = currentData ? foldRunsTasks(currentData) : [];
228
+ };
229
+ const runApprove = async (taskId) => {
230
+ if (!repoRoot || !currentData?.runId) {
231
+ setNotice("approval not available — no run loaded");
232
+ return;
233
+ }
234
+ busy = true;
235
+ setNotice("approving…");
236
+ try {
237
+ const message = await approve([currentData.runId, taskId], repoRoot);
238
+ await refresh();
239
+ setNotice(message);
240
+ }
241
+ catch (e) {
242
+ setNotice(e instanceof Error ? e.message : String(e));
243
+ }
244
+ finally {
245
+ busy = false;
246
+ approvalPromise = undefined;
247
+ }
248
+ };
249
+ const view = {
250
+ id: "runs",
251
+ label: "Runs",
252
+ cursor: 0,
253
+ key(name) {
254
+ if (busy)
255
+ return;
256
+ if (confirming) {
257
+ const taskId = confirming.taskId;
258
+ confirming = null;
259
+ if (name === "y") {
260
+ approvalPromise = runApprove(taskId);
261
+ return;
262
+ }
263
+ setNotice("approval cancelled");
264
+ return;
265
+ }
266
+ if (!tasks.length)
267
+ return;
268
+ if (name === "up")
269
+ cursor = Math.max(cursor - 1, 0);
270
+ else if (name === "down")
271
+ cursor = Math.min(cursor + 1, tasks.length - 1);
272
+ else if (name === "a") {
273
+ if (!repoRoot || !currentData?.runId) {
274
+ setNotice("approval not available — no run loaded");
275
+ return;
276
+ }
277
+ const task = tasks[cursor];
278
+ if (!task)
279
+ return;
280
+ if (task.status === "human") {
281
+ confirming = { taskId: task.task.id };
282
+ setNotice(`approve ${task.task.id}? releases the parked human gate and resumes the run. [y] confirm [any key] cancel`);
283
+ return;
284
+ }
285
+ approvalPromise = runApprove(task.task.id);
286
+ }
287
+ this.cursor = cursor;
288
+ },
289
+ selectedTask() {
290
+ return tasks[cursor];
291
+ },
292
+ render(props) {
293
+ if (!currentData || !currentData.runId) {
294
+ return emptyState();
295
+ }
296
+ const last = currentData.events.at(-1);
297
+ const now = last ? formatJournalNarration(last) : undefined;
298
+ const lines = [];
299
+ if (now)
300
+ lines.push(legend(` now: ${now}`));
301
+ else
302
+ lines.push(legend(" now: —"));
303
+ lines.push(tipVerifyLine(currentData.events));
304
+ lines.push(legend(` gates: ${GATE_NAMES.join(" · ")}`));
305
+ lines.push("");
306
+ const { rows } = props;
307
+ const ticker = costTickerPanel(currentData);
308
+ // Reserve space for: now, tip-verify, gates legend, blank, detail panel (blank + divider +
309
+ // content + attempts), cost ticker panel, optional notice line, trailing blank
310
+ const detailReserve = 6 + ticker.length + (notice ? 1 : 0);
311
+ const cardBudget = Math.max(0, rows - lines.length - detailReserve);
312
+ const visibleCount = Math.max(1, Math.floor(cardBudget / CARD_LINES));
313
+ const maxStart = Math.max(0, tasks.length - visibleCount);
314
+ let start = cursor - Math.floor(visibleCount / 2);
315
+ if (start < 0)
316
+ start = 0;
317
+ if (start > maxStart)
318
+ start = maxStart;
319
+ const end = Math.min(tasks.length, start + visibleCount);
320
+ const width = Math.min(props.cols, 100);
321
+ const idW = Math.max(2, ...tasks.map((c) => c.task.id.length));
322
+ const goalAvail = Math.max(8, width - (idW + 12));
323
+ const goals = tasks.map((c) => shortGoal(c.task.goal, goalAvail));
324
+ const goalW = Math.max(8, ...goals.map((s) => s.length));
325
+ const indent = " ".repeat(idW + 5);
326
+ for (let i = start; i < end; i++) {
327
+ const t = tasks[i];
328
+ const goal = goals[i];
329
+ const pointer = i === cursor ? `${GLYPHS.pointer} ` : " ";
330
+ const f = failedGates(t.states);
331
+ const human = humanGateSuffix(t.task, t.status, t.states);
332
+ const statusCell = renderStatusCell(t.status, f, human);
333
+ const line1 = ` ${pointer}${statusRow(taskVerdict(t.status), `${t.task.id.padEnd(idW)} ${goal.padEnd(goalW)} ${statusCell}`)}`;
334
+ const detailParts = [t.activity ?? t.channel];
335
+ if (t.ctx !== undefined)
336
+ detailParts.push(`ctx ${t.ctx}`);
337
+ const line2 = `${indent}${gateChain(t.states, true)} ${dim(detailParts.join(" · "))}`;
338
+ lines.push(line1, line2, "");
339
+ }
340
+ lines.push(...detailPanel(tasks[cursor]));
341
+ lines.push(...ticker);
342
+ if (notice)
343
+ lines.push(legend(` ${notice}`));
344
+ return lines;
345
+ },
346
+ };
347
+ Object.defineProperty(view, "approval", {
348
+ get: () => approvalPromise,
349
+ configurable: true,
350
+ });
351
+ return view;
352
+ }
353
+ function renderStatusCell(status, failed, human) {
354
+ const stWord = status === "done" ? ok(String(status))
355
+ : status === "failed" ? fail(String(status))
356
+ : status === "human" ? warn(String(status))
357
+ : String(status);
358
+ const dot = dim(" · ");
359
+ return stWord +
360
+ (failed.length ? dot + fail(failed.join(", ")) : "") +
361
+ (human ? dot + warn("awaiting approval") : "");
362
+ }
363
+ function detailPanel(task) {
364
+ if (!task)
365
+ return [];
366
+ const lines = ["", dim(`── ${task.task.id} — attempts & consult dossier ──`)];
367
+ if (task.attempts.length === 0) {
368
+ lines.push(` no attempts recorded for ${task.task.id}`);
369
+ return lines;
370
+ }
371
+ for (const a of task.attempts) {
372
+ const reason = a.cause ? ` — ${a.cause}` : a.note ? ` — ${a.note}` : "";
373
+ const statusWord = a.outcome === "clean" ? "done" : a.outcome === "failed" ? "failed" : "in flight";
374
+ lines.push(` attempt ${a.attempt} ${a.channel} ${statusWord}${reason}`);
375
+ }
376
+ lines.push("");
377
+ lines.push(dim(renderDossierPlaceholder()));
378
+ return lines;
379
+ }
380
+ function emptyState() {
381
+ return [
382
+ "",
383
+ " no run loaded — run `tickmarkr run` or `tickmarkr resume <runId>` to start one.",
384
+ " the Runs cockpit reads the active or most recent run's journal; nothing to show yet.",
385
+ "",
386
+ ];
387
+ }
@@ -0,0 +1 @@
1
+ canary-verified
@@ -0,0 +1,8 @@
1
+ <!-- tickmarkr:spec -->
2
+ # Judge canary
3
+
4
+ ## T1: Canary verification marker
5
+ - goal: The repository contains a verified canary marker
6
+ - shape: implement
7
+ - acceptance:
8
+ - judge: The file a.txt contains the exact text "canary-verified".
@@ -0,0 +1 @@
1
+ canary-start
@@ -0,0 +1 @@
1
+ hello from solution
@@ -0,0 +1,8 @@
1
+ <!-- tickmarkr:spec -->
2
+ # Eval fixture sample
3
+
4
+ ## T1: Solution text is present
5
+ - goal: The fixture reference content is produced
6
+ - shape: implement
7
+ - acceptance:
8
+ - command: [ "$(cat a.txt)" = "hello from solution" ]
@@ -0,0 +1 @@
1
+ hello from start
@@ -0,0 +1,42 @@
1
+ ---
2
+ phase: 07-live-check
3
+ plan: "01"
4
+ type: execute
5
+ wave: 1
6
+ depends_on: []
7
+ files_modified:
8
+ - src/**
9
+ autonomous: true
10
+ must_haves:
11
+ truths:
12
+ - "truth A holds in the shipped artifact"
13
+ artifacts:
14
+ - path: "src/out.js"
15
+ provides: "the thing"
16
+ ---
17
+
18
+ <objective>
19
+ Implement the first objective sentence. Additional prose that is not the title.
20
+ </objective>
21
+
22
+ <context>
23
+ @.planning/PROJECT.md
24
+ @$HOME/.claude/get-shit-done/workflows/execute-plan.md
25
+ @~/somewhere/outside.md
26
+ </context>
27
+
28
+ <tasks>
29
+
30
+ <task type="auto">
31
+ <name>Task 1: Build the widget</name>
32
+ <action>do the thing</action>
33
+ <verify>look at it</verify>
34
+ <done>widget builds green</done>
35
+ </task>
36
+
37
+ <task type="auto">
38
+ <name>Task 2: Wire the widget</name>
39
+ <done>widget is wired and demoed</done>
40
+ </task>
41
+
42
+ </tasks>
@@ -0,0 +1,21 @@
1
+ ---
2
+ phase: 07-live-check
3
+ plan: "02"
4
+ type: execute
5
+ wave: 2
6
+ depends_on: ["07-01"]
7
+ files_modified:
8
+ - docs/**
9
+ autonomous: true
10
+ ---
11
+
12
+ <objective>
13
+ Document the widget. Trailing sentence.
14
+ </objective>
15
+
16
+ <tasks>
17
+ <task type="checkpoint:human-action" gate="blocking">
18
+ <name>Task 1: Operator reviews the docs</name>
19
+ <done>operator approved the docs page</done>
20
+ </task>
21
+ </tasks>
@@ -0,0 +1,18 @@
1
+ ---
2
+ phase: 07-live-check
3
+ plan: "03"
4
+ type: execute
5
+ depends_on: ["07-02"]
6
+ autonomous: true
7
+ ---
8
+
9
+ <objective>
10
+ Already-finished work. Done earlier.
11
+ </objective>
12
+
13
+ <tasks>
14
+ <task type="auto">
15
+ <name>Task 1: old work</name>
16
+ <done>previously completed</done>
17
+ </task>
18
+ </tasks>
@@ -0,0 +1 @@
1
+ # done earlier
@@ -0,0 +1,10 @@
1
+ <!-- tickmarkr:spec -->
2
+
3
+ ## T1: Missing mandatory build gate
4
+ - gates:
5
+ - test
6
+ - lint
7
+ - evidence
8
+ - scope
9
+ - acceptance:
10
+ - compilation fails loudly naming build
@@ -0,0 +1,18 @@
1
+ # Sample PRD: pinned delivery
2
+
3
+ ## T1: Implement pinned compiler work
4
+ - shape: implement
5
+ - complexity: 7
6
+ - files: src/compiler.ts, src/types.ts
7
+ - pin: claude-code sonnet
8
+ - acceptance:
9
+ - compiler returns a graph
10
+
11
+ ## T2: Verify pinned compiler work
12
+ - shape: tests
13
+ - deps: T1
14
+ - files: tests/compiler.test.ts
15
+ - humanGate: true
16
+ - acceptance:
17
+ - tests reject malformed input
18
+ - reviewer approves the output
@@ -0,0 +1,35 @@
1
+ <!-- drovr:spec -->
2
+ # Native spec: compiler delivery
3
+
4
+ ## T1: Build the native compiler
5
+ - goal: Compile the complete native task surface
6
+ - shape: implement
7
+ - deps: none
8
+ - files: src/compile/native.ts, src/compile/index.ts
9
+ - context: docs/native.md, src/graph/schema.ts
10
+ - complexity: 8
11
+ - humanGate: true
12
+ - pin: claude-code opus
13
+ - floor: frontier
14
+ - gates:
15
+ - build
16
+ - test
17
+ - lint
18
+ - evidence
19
+ - scope
20
+ - acceptance
21
+ - acceptance:
22
+ - every native field reaches the graph
23
+ - malformed fields fail loudly
24
+
25
+ ## T2: Test native detection
26
+ - goal: Keep native and generic markdown routing distinct
27
+ - shape: tests
28
+ - deps: T1
29
+ - files: tests/compile/native.test.ts
30
+ - context: fixtures/sample.native.md
31
+ - complexity: 3
32
+ - humanGate: false
33
+ - acceptance:
34
+ - marked markdown selects native
35
+ - marker-less markdown stays PRD
@@ -0,0 +1,22 @@
1
+ # Sample PRD: tiny greeter
2
+
3
+ ## T1: Implement greet(name) in src/greet.js
4
+ - shape: implement
5
+ - complexity: 3
6
+ - files: src/**
7
+ - acceptance:
8
+ - greet("drover") returns a string containing "drover"
9
+ - module exports a single function
10
+
11
+ ## T2: Cover greet with tests
12
+ - deps: T1
13
+ - files: test/**
14
+ - acceptance:
15
+ - tests fail when greet drops the name
16
+ - npm test exits 0
17
+
18
+ ## T3: Document usage in README
19
+ - deps: T1
20
+ - humanGate: true
21
+ - acceptance:
22
+ - README shows an executable example