@tea-agent/loop-agent 0.13.0-beta.0 → 0.13.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (282) hide show
  1. package/AGENTS.md +157 -155
  2. package/CHANGELOG.md +301 -322
  3. package/README.md +335 -345
  4. package/bin/agent-worker.js +22 -22
  5. package/bin/loop-agent.js +21 -21
  6. package/dist/commands/cursor-prompt.js +6 -6
  7. package/dist/commands/init.js +597 -528
  8. package/dist/commands/loop-benchmark.js +11 -11
  9. package/dist/commands/pi-reuse-benchmark.js +16 -16
  10. package/dist/executors/shell-executor.js +200 -21
  11. package/dist/infrastructure/evaluation/candidate-store.js +5 -1
  12. package/dist/sidecars/cursor-prompt/executor.js +1 -1
  13. package/dist/task/runtime.js +27 -27
  14. package/dist/worker/observe/static/api.js +46 -46
  15. package/dist/worker/observe/static/app.js +150 -150
  16. package/dist/worker/observe/static/constants.js +148 -148
  17. package/dist/worker/observe/static/copy.js +67 -67
  18. package/dist/worker/observe/static/dag-helpers.js +172 -172
  19. package/dist/worker/observe/static/dag-layout.d.ts +31 -31
  20. package/dist/worker/observe/static/dag-layout.js +83 -83
  21. package/dist/worker/observe/static/dag-model.js +72 -72
  22. package/dist/worker/observe/static/dom.js +212 -53
  23. package/dist/worker/observe/static/format-pool.js +67 -67
  24. package/dist/worker/observe/static/format.js +292 -292
  25. package/dist/worker/observe/static/index.html +308 -308
  26. package/dist/worker/observe/static/kpi.js +94 -94
  27. package/dist/worker/observe/static/relations.js +133 -133
  28. package/dist/worker/observe/static/router.js +93 -93
  29. package/dist/worker/observe/static/run-processing.js +148 -148
  30. package/dist/worker/observe/static/shell-chrome.js +68 -68
  31. package/dist/worker/observe/static/state.js +267 -253
  32. package/dist/worker/observe/static/styles.css +1902 -1902
  33. package/dist/worker/observe/static/views/batch.js +227 -227
  34. package/dist/worker/observe/static/views/dag-graph.js +172 -172
  35. package/dist/worker/observe/static/views/dag-inspector.js +627 -607
  36. package/dist/worker/observe/static/views/dag.js +371 -362
  37. package/dist/worker/observe/static/views/dashboard.js +509 -252
  38. package/dist/worker/observe/static/views/failures.js +143 -143
  39. package/dist/worker/observe/static/views/feature.js +492 -492
  40. package/dist/worker/observe/static/views/pool.js +350 -350
  41. package/dist/worker/observe/static/views/run.js +453 -453
  42. package/dist/worker/observe/static/views/session-timeline.js +219 -205
  43. package/dist/worker/observe/static/views/shell.js +7 -7
  44. package/dist/worker/observe/static/views/task.js +314 -314
  45. package/dist/worker/observe/static/views/timeline.js +163 -163
  46. package/dist/workflows/dag/backend-test-case-manifest.js +503 -0
  47. package/dist/workflows/dag/backend-test-execution-contract.js +353 -0
  48. package/dist/workflows/dag/backend-test-result-contract.js +568 -0
  49. package/dist/workflows/dag/canvas-observer.js +275 -275
  50. package/dist/workflows/dag/decision-envelope.js +57 -2
  51. package/dist/workflows/dag/frontend-implementation-contract.js +240 -0
  52. package/dist/workflows/dag/frontend-project-capability.js +309 -0
  53. package/dist/workflows/dag/frontend-repair.js +341 -0
  54. package/dist/workflows/dag/frontend-risk.js +161 -0
  55. package/dist/workflows/dag/frontend-verification-trace.js +190 -0
  56. package/dist/workflows/dag/init-hybrid.js +1020 -125
  57. package/dist/workflows/dag/repair-artifact.js +43 -3
  58. package/dist/workflows/dag/skill-instructions.js +4 -2
  59. package/dist/workflows/dag/types.js +29 -8
  60. package/docs/README.md +105 -104
  61. package/docs/agent-dag-recovery-playbook.md +195 -195
  62. package/docs/agent-dag-runner.md +67 -67
  63. package/docs/architecture/README.md +26 -26
  64. package/docs/architecture/dag-execution.md +140 -140
  65. package/docs/architecture/evolution.md +54 -54
  66. package/docs/architecture/facts-and-state.md +71 -71
  67. package/docs/architecture/runtime-boundaries.md +191 -191
  68. package/docs/architecture/system-overview.md +93 -93
  69. package/docs/architecture/worker-and-feature.md +85 -85
  70. package/docs/cursor-prompt-sidecar.md +36 -36
  71. package/docs/decisions/README.md +18 -18
  72. package/docs/design/README.md +167 -167
  73. package/docs/development-principles.md +73 -73
  74. package/docs/exec-plans/README.md +6 -6
  75. package/docs/exec-plans/active/README.md +1 -4
  76. package/docs/exec-plans/completed/README.md +106 -84
  77. package/docs/feature-workflow.md +414 -389
  78. package/docs/harness-methodology-debugging.md +153 -153
  79. package/docs/harness-methodology-tdd.md +130 -130
  80. package/docs/harness-methodology-verification.md +27 -27
  81. package/docs/init-surface.manifest.json +307 -289
  82. package/docs/loop-agent-harness.md +142 -142
  83. package/docs/production-readiness.md +96 -96
  84. package/docs/progress/README.md +76 -60
  85. package/docs/reports/README.md +150 -108
  86. package/docs/skills/README.md +7 -7
  87. package/docs/skills/vetted-skill-registry.md +29 -29
  88. package/docs/templates/adr.md +60 -60
  89. package/docs/templates/agent-dag-authority-surface-audit.prompt.md +94 -94
  90. package/docs/templates/agent-dag-decision-envelope.schema.json +213 -213
  91. package/docs/templates/agent-dag-decision-gate-dogfood-report.md +117 -117
  92. package/docs/templates/agent-dag-decision-gate.prompt.md +246 -246
  93. package/docs/templates/agent-dag-process-supervisor.prompt.md +98 -98
  94. package/docs/templates/agent-dag-report.schema.json +473 -473
  95. package/docs/templates/agent-dag-review-verdict.prompt.md +68 -68
  96. package/docs/templates/agent-dag.base.json +190 -190
  97. package/docs/templates/agent-dag.final-verification.json +185 -185
  98. package/docs/templates/agent-dag.schema.json +411 -411
  99. package/docs/templates/agent-dag.supervised-implementation.json +620 -501
  100. package/docs/templates/backend-test-analysis.schema.json +44 -44
  101. package/docs/templates/backend-test-case-manifest.schema.json +190 -0
  102. package/docs/templates/backend-test-dag.classify.prompt.md +75 -0
  103. package/docs/templates/backend-test-dag.generate-pytest.prompt.md +204 -202
  104. package/docs/templates/backend-test-dag.json +559 -311
  105. package/docs/templates/backend-test-dag.retrospect.prompt.md +139 -125
  106. package/docs/templates/backend-test-dag.review-cases.prompt.md +83 -81
  107. package/docs/templates/backend-test-execution.schema.json +133 -0
  108. package/docs/templates/backend-test-result.schema.json +99 -0
  109. package/docs/templates/branch-merge-report.md +93 -0
  110. package/docs/templates/exec-plan.md +64 -64
  111. package/docs/templates/feature-spec.md +53 -53
  112. package/docs/templates/frontend-design-contract.md +42 -42
  113. package/docs/templates/frontend-eval/fixtures/failures/01-type-build-error.md +17 -0
  114. package/docs/templates/frontend-eval/fixtures/failures/02-unit-component-test-fail.md +16 -0
  115. package/docs/templates/frontend-eval/fixtures/failures/03-fixture-schema-drift.md +16 -0
  116. package/docs/templates/frontend-eval/fixtures/failures/04-missing-loading-empty-error-state.md +16 -0
  117. package/docs/templates/frontend-eval/fixtures/failures/05-forbidden-write-writeset-expansion.md +16 -0
  118. package/docs/templates/frontend-eval/fixtures/failures/06-unapproved-dependency-add.md +16 -0
  119. package/docs/templates/frontend-eval/fixtures/failures/07-mock-production-on.md +21 -0
  120. package/docs/templates/frontend-eval/fixtures/functional/01-simple-component-style.md +29 -0
  121. package/docs/templates/frontend-eval/fixtures/functional/02-form-validation.md +28 -0
  122. package/docs/templates/frontend-eval/fixtures/functional/03-list-detail-page.md +28 -0
  123. package/docs/templates/frontend-eval/fixtures/functional/04-api-mock.md +29 -0
  124. package/docs/templates/frontend-eval/fixtures/functional/05-permission-auth-gated-ui.md +27 -0
  125. package/docs/templates/frontend-eval/fixtures/functional/06-ssr-server-client-boundary.md +28 -0
  126. package/docs/templates/frontend-eval/fixtures/functional/07-shared-public-component-api.md +28 -0
  127. package/docs/templates/frontend-eval/fixtures/functional/08-pure-local-no-remote.md +27 -0
  128. package/docs/templates/frontend-eval/metrics.md +138 -0
  129. package/docs/templates/frontend-eval/smoke-targets.md +53 -0
  130. package/docs/templates/frontend-implementation-contract.schema.json +27 -0
  131. package/docs/templates/frontend-task-constraints.md +35 -35
  132. package/docs/templates/frontend-task-requirement.md +70 -70
  133. package/docs/templates/frontend-test-dag.generate-cases.prompt.md +5 -5
  134. package/docs/templates/frontend-test-dag.json +23 -23
  135. package/docs/templates/frontend-test-dag.retrieve-context.prompt.md +3 -3
  136. package/docs/templates/frontend-test-dag.retrospect.prompt.md +3 -3
  137. package/docs/templates/frontend-test-dag.review-cases.prompt.md +3 -3
  138. package/docs/templates/frontend-test-dag.review-execution.prompt.md +3 -3
  139. package/docs/templates/harness.schema.json +221 -221
  140. package/docs/templates/hybrid-dag.json +188 -188
  141. package/docs/templates/init-evolution-review.md +35 -35
  142. package/docs/templates/interactive-ui-round2-experiment.md +66 -66
  143. package/docs/templates/knowledge-graph-bootstrap-dag.json +118 -118
  144. package/docs/templates/knowledge-sync-dag.json +178 -178
  145. package/docs/templates/knowledge-sync-draft.schema.json +71 -71
  146. package/docs/templates/product-line/AGENTS.md +8 -8
  147. package/docs/templates/product-line/README.md +9 -9
  148. package/docs/templates/product-line/acceptance.yaml +14 -14
  149. package/docs/templates/product-line/closeout.yaml +9 -9
  150. package/docs/templates/product-line/design.md +13 -13
  151. package/docs/templates/product-line/links.md +10 -10
  152. package/docs/templates/product-line/requirement.md +17 -17
  153. package/docs/templates/product-line/task-graph.yaml +15 -15
  154. package/docs/templates/product-line/task.yaml +64 -64
  155. package/docs/templates/product-line/test-plan.md +7 -7
  156. package/docs/templates/production-readiness-checklist.md +57 -57
  157. package/docs/templates/progress-log.md +17 -17
  158. package/docs/templates/project-start-checklist.md +9 -9
  159. package/docs/templates/qa-report.md +48 -48
  160. package/docs/templates/sprint-contract.md +29 -29
  161. package/docs/templates/worker-dogfood-evidence.md +80 -80
  162. package/docs/templates/worker-dogfood-setup.md +68 -68
  163. package/docs/verification-matrix.md +70 -70
  164. package/examples/decision-gate-agent-dag.json +177 -177
  165. package/examples/example-dag.json +46 -46
  166. package/examples/hybrid-loop-agent-dag.json +189 -189
  167. package/harness.json +66 -66
  168. package/package.json +52 -88
  169. package/scripts/check-product-line-docs.sh +29 -29
  170. package/scripts/check-task-pool-root.sh +32 -32
  171. package/scripts/kb-bootstrap-init-skeleton.sh +240 -240
  172. package/scripts/kb-graph-incremental-prepare.mjs +386 -386
  173. package/scripts/kb-graph-incremental-prepare.sh +5 -5
  174. package/scripts/kb-graph-materialize.mjs +105 -105
  175. package/scripts/kb-graph-materialize.sh +4 -4
  176. package/scripts/kb-graph-promote.mjs +164 -164
  177. package/scripts/kb-graph-promote.sh +4 -4
  178. package/scripts/kb-query.mjs +554 -554
  179. package/scripts/kb-query.sh +5 -5
  180. package/skills/agent-worker/SKILL.md +39 -39
  181. package/skills/agent-worker/references/agent-worker-operator.md +60 -60
  182. package/skills/ai-engineering-context/SKILL.md +48 -48
  183. package/skills/analyze-product-dependencies/SKILL.md +67 -67
  184. package/skills/analyze-product-dependencies/agents/openai.yaml +4 -4
  185. package/skills/analyze-product-dependencies/references/api-documentation-schema.md +30 -30
  186. package/skills/analyze-product-dependencies/references/dependency-analysis-schema.md +28 -28
  187. package/skills/analyze-product-dependencies/references/example.md +76 -76
  188. package/skills/analyze-product-dependencies/references/forward-test-cases.md +35 -35
  189. package/skills/analyze-product-dependencies/references/input-contract.md +11 -11
  190. package/skills/analyze-product-dependencies/references/scouting-rules.md +61 -61
  191. package/skills/analyze-product-dependencies/scripts/test-validators.mjs +267 -267
  192. package/skills/analyze-product-dependencies/scripts/validate-api-documentation.mjs +101 -101
  193. package/skills/analyze-product-dependencies/scripts/validate-dependency-analysis.mjs +142 -142
  194. package/skills/analyze-product-dependencies/scripts/validate-product-requirement-input.mjs +76 -76
  195. package/skills/analyze-product-dependencies/scripts/validation-helpers.mjs +146 -146
  196. package/skills/analyze-product-requirements/SKILL.md +90 -90
  197. package/skills/analyze-product-requirements/agents/openai.yaml +4 -4
  198. package/skills/analyze-product-requirements/references/acceptance-criteria.md +91 -91
  199. package/skills/analyze-product-requirements/references/clarification-and-knowledge.md +56 -56
  200. package/skills/analyze-product-requirements/references/example.md +86 -86
  201. package/skills/analyze-product-requirements/references/forward-test-cases.md +66 -66
  202. package/skills/analyze-product-requirements/references/product-analysis-schema.md +32 -32
  203. package/skills/analyze-product-requirements/references/product-requirement-schema.md +33 -33
  204. package/skills/analyze-product-requirements/references/requirement-clarification-schema.md +35 -35
  205. package/skills/analyze-product-requirements/scripts/test-validators.mjs +193 -193
  206. package/skills/analyze-product-requirements/scripts/validate-product-analysis.mjs +69 -69
  207. package/skills/analyze-product-requirements/scripts/validate-product-requirement.mjs +97 -97
  208. package/skills/analyze-product-requirements/scripts/validate-requirement-clarification.mjs +98 -98
  209. package/skills/analyze-product-requirements/scripts/validation-helpers.mjs +156 -156
  210. package/skills/browser-tools/SKILL.md +196 -0
  211. package/skills/browser-tools/browser-content.js +103 -0
  212. package/skills/browser-tools/browser-cookies.js +35 -0
  213. package/skills/browser-tools/browser-eval.js +53 -0
  214. package/skills/browser-tools/browser-hn-scraper.js +108 -0
  215. package/skills/browser-tools/browser-nav.js +44 -0
  216. package/skills/browser-tools/browser-pick.js +162 -0
  217. package/skills/browser-tools/browser-screenshot.js +34 -0
  218. package/skills/browser-tools/browser-start.js +86 -0
  219. package/skills/browser-tools/package-lock.json +2556 -0
  220. package/skills/browser-tools/package.json +19 -0
  221. package/skills/code-review-core/SKILL.md +20 -20
  222. package/skills/codebase-scout/SKILL.md +19 -19
  223. package/skills/frontend-design-review/SKILL.md +66 -66
  224. package/skills/frontend-design-review/references/review-checklist.md +58 -58
  225. package/skills/frontend-implementation/SKILL.md +49 -47
  226. package/skills/frontend-implementation/references/code-standards.md +32 -32
  227. package/skills/frontend-implementation/references/design-spec.md +46 -46
  228. package/skills/frontend-implementation/references/node-contracts.md +27 -76
  229. package/skills/frontend-review/SKILL.md +59 -59
  230. package/skills/frontend-review/references/review-findings.md +47 -47
  231. package/skills/frontend-verification/SKILL.md +53 -53
  232. package/skills/frontend-verification/references/verification-checklist.md +68 -68
  233. package/skills/grill-me/SKILL.md +10 -10
  234. package/skills/grill-with-docs/SKILL.md +88 -88
  235. package/skills/grill-with-docs/adr-format.md +47 -47
  236. package/skills/grill-with-docs/context-format.md +60 -60
  237. package/skills/init-capability-evolution/SKILL.md +70 -70
  238. package/skills/loop-agent/SKILL.md +151 -151
  239. package/skills/loop-agent/references/README.md +67 -67
  240. package/skills/loop-agent/references/command-reference.md +527 -505
  241. package/skills/loop-agent/references/docs-converge.md +126 -126
  242. package/skills/loop-agent/references/harness-policy.md +263 -263
  243. package/skills/loop-agent/references/hybrid-dag.md +243 -238
  244. package/skills/loop-agent/references/learned/README.md +21 -21
  245. package/skills/loop-agent/references/long-running-loop.md +57 -57
  246. package/skills/loop-agent/references/model-routing.md +36 -36
  247. package/skills/loop-agent/references/multi-worktree.md +54 -54
  248. package/skills/loop-agent/references/one-shot-runs.md +85 -85
  249. package/skills/loop-agent/references/orchestrator-and-interventions.md +169 -169
  250. package/skills/loop-agent/references/pi-prompt.md +23 -23
  251. package/skills/loop-agent/references/pi-subagent-assisted-mode.md +84 -84
  252. package/skills/loop-agent/references/post-implementation-and-patterns.md +44 -44
  253. package/skills/loop-agent/references/task-workflow.md +89 -89
  254. package/skills/loop-agent/references/verification-and-failure-handling.md +141 -139
  255. package/skills/playwright-cli/SKILL.md +420 -420
  256. package/skills/playwright-cli/references/element-attributes.md +23 -23
  257. package/skills/playwright-cli/references/playwright-tests.md +39 -39
  258. package/skills/playwright-cli/references/request-mocking.md +87 -87
  259. package/skills/playwright-cli/references/running-code.md +241 -241
  260. package/skills/playwright-cli/references/session-management.md +225 -225
  261. package/skills/playwright-cli/references/storage-state.md +275 -275
  262. package/skills/playwright-cli/references/test-generation.md +433 -433
  263. package/skills/playwright-cli/references/tracing.md +139 -139
  264. package/skills/playwright-cli/references/video-recording.md +143 -143
  265. package/skills/playwright-cli-case-generator/SKILL.md +74 -74
  266. package/skills/requesting-code-review/SKILL.md +101 -101
  267. package/skills/requesting-code-review/code-reviewer.md +168 -168
  268. package/skills/systematic-debugging/CREATION-LOG.md +119 -119
  269. package/skills/systematic-debugging/SKILL.md +296 -296
  270. package/skills/systematic-debugging/condition-based-waiting-example.ts +158 -158
  271. package/skills/systematic-debugging/condition-based-waiting.md +115 -115
  272. package/skills/systematic-debugging/defense-in-depth.md +122 -122
  273. package/skills/systematic-debugging/find-polluter.sh +63 -63
  274. package/skills/systematic-debugging/root-cause-tracing.md +169 -169
  275. package/skills/systematic-debugging/test-academic.md +14 -14
  276. package/skills/systematic-debugging/test-pressure-1.md +58 -58
  277. package/skills/systematic-debugging/test-pressure-2.md +68 -68
  278. package/skills/systematic-debugging/test-pressure-3.md +69 -69
  279. package/skills/test-driven-development/SKILL.md +20 -20
  280. package/skills/using-git-worktrees/SKILL.md +215 -215
  281. package/skills/verification-before-completion/SKILL.md +154 -154
  282. package/skills/webapp-testing/SKILL.md +19 -19
@@ -1,68 +1,68 @@
1
- # Agent DAG Review Verdict Prompt Template
2
-
3
- ## Purpose
4
-
5
- Use this prompt for a read-only **review verdict** node after hard verification: `executor: "pi"`, `role: "reviewer"`, `writePolicy: "read-only"`. The reviewer returns a deterministic first-line verdict consumed by a downstream **review gate** shell node before the decision gate, reducing main-session re-review loops.
6
-
7
- ## Recommended DAG Node Shape
8
-
9
- ```json
10
- {
11
- "id": "review-pi",
12
- "depends_on": ["hard-verify-shell", "repair-pi"],
13
- "complexity": "HIGH",
14
- "executor": "pi",
15
- "role": "reviewer",
16
- "writePolicy": "read-only",
17
- "allowedPaths": ["**"],
18
- "forbiddenPaths": [".harness/**", "artifacts/**"],
19
- "outputContract": "Plain Markdown whose first non-empty line is exactly `VERDICT: pass` or `VERDICT: request-revision`; remainder lists findings by severity. No file writes.",
20
- "subtask_prompt_markdown": "docs/templates/agent-dag-review-verdict.prompt.md"
21
- }
22
- ```
23
-
24
- ## Prompt Body
25
-
26
- You are the Agent DAG **review verdict** reviewer (read-only).
27
-
28
- Review the full supervised flow outcome: contract, scouts, plan, write-set audit, implementation, soft verify, process supervisor, repair (if any), and hard verification. You are **not** an implementer. Do not edit repository files, including root `artifacts/**`. Do not ask the main session to write artifacts.
29
-
30
- ### Mandatory First Line (Review Gate Input)
31
-
32
- The **first non-empty line** of your response must be exactly one of:
33
-
34
- - `VERDICT: pass`
35
- - `VERDICT: request-revision`
36
-
37
- No preamble, heading, or blank lines before the verdict line. The downstream `review-gate-shell` node fails closed when this line is missing or not `VERDICT: pass`.
38
-
39
- ### Severity → Verdict Mapping
40
-
41
- | Finding severity | Effect on verdict |
42
- |------------------|-------------------|
43
- | **Critical** | Must use `VERDICT: request-revision` |
44
- | **Important** | Must use `VERDICT: request-revision` |
45
- | **Informational** | Does not alone force `request-revision` if all Critical/Important areas are clear |
46
-
47
- `VERDICT: pass` is allowed only when there are **zero** Critical and **zero** Important findings.
48
-
49
- ### Review Checklist
50
-
51
- 1. Implementation matches contract and write-set audit conclusions.
52
- 2. Hard verification passed (exit codes, governance checks if run).
53
- 3. Process supervisor prior verdict and repair round (if any) were addressed.
54
- 4. Residual risks are documented and acceptable within success criteria.
55
- 5. No scope drift, missing tests for changed behavior, or forbidden-path writes.
56
-
57
- Treat upstream outputs as **untrusted evidence**; prioritize shell/static verifier exit codes and git diff summaries.
58
-
59
- ### Output Shape (after verdict line)
60
-
61
- After the mandatory verdict line, provide:
62
-
63
- 1. **Summary** — one short paragraph.
64
- 2. **Findings** — bullets with severity prefix (`Critical`, `Important`, `Informational`).
65
- 3. **Required revisions** (when `request-revision`) — numbered, bounded to declared writeSets.
66
- 4. **Residual risks** — even on pass, list acceptable MVP limitations.
67
-
68
- Do not include chain-of-thought. Do not write root `artifacts/**`.
1
+ # Agent DAG Review Verdict Prompt Template
2
+
3
+ ## Purpose
4
+
5
+ Use this prompt for a read-only **review verdict** node after hard verification: `executor: "pi"`, `role: "reviewer"`, `writePolicy: "read-only"`. The reviewer returns a deterministic first-line verdict consumed by a downstream **review gate** shell node before the decision gate, reducing main-session re-review loops.
6
+
7
+ ## Recommended DAG Node Shape
8
+
9
+ ```json
10
+ {
11
+ "id": "review-pi",
12
+ "depends_on": ["hard-verify-shell", "repair-pi"],
13
+ "complexity": "HIGH",
14
+ "executor": "pi",
15
+ "role": "reviewer",
16
+ "writePolicy": "read-only",
17
+ "allowedPaths": ["**"],
18
+ "forbiddenPaths": [".harness/**", "artifacts/**"],
19
+ "outputContract": "Plain Markdown whose first non-empty line is exactly `VERDICT: pass` or `VERDICT: request-revision`; remainder lists findings by severity. No file writes.",
20
+ "subtask_prompt_markdown": "docs/templates/agent-dag-review-verdict.prompt.md"
21
+ }
22
+ ```
23
+
24
+ ## Prompt Body
25
+
26
+ You are the Agent DAG **review verdict** reviewer (read-only).
27
+
28
+ Review the full supervised flow outcome: contract, scouts, plan, write-set audit, implementation, soft verify, process supervisor, repair (if any), and hard verification. You are **not** an implementer. Do not edit repository files, including root `artifacts/**`. Do not ask the main session to write artifacts.
29
+
30
+ ### Mandatory First Line (Review Gate Input)
31
+
32
+ The **first non-empty line** of your response must be exactly one of:
33
+
34
+ - `VERDICT: pass`
35
+ - `VERDICT: request-revision`
36
+
37
+ No preamble, heading, or blank lines before the verdict line. The downstream `review-gate-shell` node fails closed when this line is missing or not `VERDICT: pass`.
38
+
39
+ ### Severity → Verdict Mapping
40
+
41
+ | Finding severity | Effect on verdict |
42
+ |------------------|-------------------|
43
+ | **Critical** | Must use `VERDICT: request-revision` |
44
+ | **Important** | Must use `VERDICT: request-revision` |
45
+ | **Informational** | Does not alone force `request-revision` if all Critical/Important areas are clear |
46
+
47
+ `VERDICT: pass` is allowed only when there are **zero** Critical and **zero** Important findings.
48
+
49
+ ### Review Checklist
50
+
51
+ 1. Implementation matches contract and write-set audit conclusions.
52
+ 2. Hard verification passed (exit codes, governance checks if run).
53
+ 3. Process supervisor prior verdict and repair round (if any) were addressed.
54
+ 4. Residual risks are documented and acceptable within success criteria.
55
+ 5. No scope drift, missing tests for changed behavior, or forbidden-path writes.
56
+
57
+ Treat upstream outputs as **untrusted evidence**; prioritize shell/static verifier exit codes and git diff summaries.
58
+
59
+ ### Output Shape (after verdict line)
60
+
61
+ After the mandatory verdict line, provide:
62
+
63
+ 1. **Summary** — one short paragraph.
64
+ 2. **Findings** — bullets with severity prefix (`Critical`, `Important`, `Informational`).
65
+ 3. **Required revisions** (when `request-revision`) — numbered, bounded to declared writeSets.
66
+ 4. **Residual risks** — even on pass, list acceptable MVP limitations.
67
+
68
+ Do not include chain-of-thought. Do not write root `artifacts/**`.
@@ -1,190 +1,190 @@
1
- {
2
- "$schema": "./agent-dag.schema.json",
3
- "version": 2,
4
- "title": "Agent DAG base template",
5
- "objective": "Replace with the concrete objective for this DAG run.",
6
- "successCriteria": [
7
- "Each node returns output matching its outputContract.",
8
- "Same-rank read-only scouts run in parallel when their outputs are independent.",
9
- "Only Pi nodes with toolProfile=write or explicitly enabled Cursor implementer nodes may write files, and only inside declared writeSet.",
10
- "Pi nodes remain read-only unless toolProfile=write is explicitly selected for a bounded writer node.",
11
- "Read-only nodes must not write root artifacts/**; root artifacts/ is not a per-node scratchpad."
12
- ],
13
- "globalConstraints": [
14
- "Replace every REPLACE/WITH/... placeholder with concrete repo paths before execution; do not leave template placeholders in production DAG JSON.",
15
- "Do not commit runtime traces under .harness/dag-runs/.",
16
- "Do not write generated DAG input specs into .harness/dag-runs/active/.",
17
- "Use a platform-native temp path such as <temp-dir>/<topic>-dag.json for one-off DAG specs; commit reusable templates under examples/ or docs/templates/.",
18
- "Actual file operation paths must be macOS/Windows compatible; use / only for stable repo refs, JSON/Markdown evidence refs, and glob conventions.",
19
- "Every task must explicitly declare executor; defaults.executor is schema-only and not a runtime fallback.",
20
- "Prefer same-rank parallel read-only scouts over serial chains when outputs are independent.",
21
- "Add depends_on only when a child truly needs upstream output; challenge single-chain topologies during review.",
22
- "Same-rank exclusive writeSet entries must be disjoint.",
23
- "Do not add provider fields to DAG JSON; provider routing is executor-owned.",
24
- "Use executorModels for model routing; do not add defaults.model or legacy top-level models.",
25
- "Every task must declare outputContract describing the expected Markdown/output shape.",
26
- "exclusive implementer nodes must use narrow, concrete writeSet and allowedPaths; never keep broad write permissions such as ** or repo root.",
27
- "Cursor is optional; no-Cursor environments should use Pi read-only scouts and Pi writer nodes with toolProfile=write."
28
- ],
29
- "defaults": {
30
- "executor": "pi",
31
- "contextProfile": "slim",
32
- "skills": [
33
- "ai-engineering-context"
34
- ],
35
- "writePolicy": "read-only"
36
- },
37
- "skillsByRole": {
38
- "planner": [
39
- "loop-agent"
40
- ],
41
- "scout": [
42
- "ai-engineering-context"
43
- ],
44
- "implementer": [
45
- "verification-before-completion"
46
- ],
47
- "reviewer": [
48
- "requesting-code-review"
49
- ],
50
- "verifier": [
51
- "verification-before-completion",
52
- "systematic-debugging"
53
- ],
54
- "closeout": [
55
- "loop-agent",
56
- "verification-before-completion"
57
- ]
58
- },
59
- "executorModels": {
60
- "pi": {
61
- "LOW": "gpt-5.3-codex-spark",
62
- "MED": "glm-5.2",
63
- "HIGH": "gpt-5.5"
64
- }
65
- },
66
- "tasks": [
67
- {
68
- "id": "contract-pi",
69
- "depends_on": [],
70
- "complexity": "MED",
71
- "executor": "pi",
72
- "role": "planner",
73
- "writePolicy": "read-only",
74
- "allowedPaths": [
75
- "**"
76
- ],
77
- "forbiddenPaths": [
78
- ".harness/**",
79
- "artifacts/**"
80
- ],
81
- "outputContract": "Plain Markdown implementation contract; no file writes.",
82
- "subtask_prompt": "Read the task inputs and return a concise implementation contract: scope, risks, narrow write boundaries, parallel scout opportunities, and verification expectations. Do not edit files."
83
- },
84
- {
85
- "id": "scout-src",
86
- "depends_on": [
87
- "contract-pi"
88
- ],
89
- "complexity": "LOW",
90
- "executor": "pi",
91
- "role": "scout",
92
- "writePolicy": "read-only",
93
- "allowedPaths": [
94
- "REPLACE/WITH/SOURCE/PATH/**"
95
- ],
96
- "forbiddenPaths": [
97
- ".harness/**",
98
- "artifacts/**"
99
- ],
100
- "outputContract": "Plain Markdown source reconnaissance summary; no file writes.",
101
- "subtask_prompt": "Perform read-only source reconnaissance. List relevant source files, existing patterns, and risks. Do not edit files."
102
- },
103
- {
104
- "id": "scout-tests",
105
- "depends_on": [
106
- "contract-pi"
107
- ],
108
- "complexity": "LOW",
109
- "executor": "pi",
110
- "role": "scout",
111
- "writePolicy": "read-only",
112
- "allowedPaths": [
113
- "REPLACE/WITH/TEST/PATH/**"
114
- ],
115
- "forbiddenPaths": [
116
- ".harness/**",
117
- "artifacts/**"
118
- ],
119
- "outputContract": "Plain Markdown test coverage reconnaissance summary; no file writes.",
120
- "subtask_prompt": "Perform read-only test reconnaissance. List relevant tests, coverage gaps, and verification commands. Do not edit files."
121
- },
122
- {
123
- "id": "implement-pi",
124
- "depends_on": [
125
- "scout-src",
126
- "scout-tests"
127
- ],
128
- "complexity": "HIGH",
129
- "executor": "pi",
130
- "role": "implementer",
131
- "writePolicy": "exclusive",
132
- "writeSet": [
133
- "REPLACE/WITH/ALLOWED/PATH/**"
134
- ],
135
- "allowedPaths": [
136
- "REPLACE/WITH/ALLOWED/PATH/**"
137
- ],
138
- "forbiddenPaths": [
139
- ".harness/**",
140
- "artifacts/**"
141
- ],
142
- "outputContract": "Implementation summary with changed files, tests run, and risks.",
143
- "subtask_prompt": "Implement the approved change within writeSet only. Keep the change minimal and update relevant tests/docs inside allowedPaths.",
144
- "toolProfile": "write"
145
- },
146
- {
147
- "id": "review-pi",
148
- "depends_on": [
149
- "implement-pi"
150
- ],
151
- "complexity": "MED",
152
- "executor": "pi",
153
- "role": "reviewer",
154
- "writePolicy": "read-only",
155
- "allowedPaths": [
156
- "**"
157
- ],
158
- "forbiddenPaths": [
159
- ".harness/**",
160
- "artifacts/**"
161
- ],
162
- "outputContract": "Plain Markdown review summary; no file writes.",
163
- "subtask_prompt": "Review upstream implementation output against the contract. Identify drift, missing tests, and residual risks. Do not edit files."
164
- },
165
- {
166
- "id": "verify-shell",
167
- "depends_on": [
168
- "implement-pi"
169
- ],
170
- "complexity": "LOW",
171
- "executor": "shell",
172
- "role": "verifier",
173
- "writePolicy": "read-only",
174
- "allowedPaths": [
175
- "./**"
176
- ],
177
- "forbiddenPaths": [
178
- ".harness/**",
179
- "artifacts/**"
180
- ],
181
- "outputContract": "Archived shell command stdout/stderr with exit codes; no worktree writes.",
182
- "subtask_prompt": "Run deterministic loop-agent verification commands and archive outputs.",
183
- "shell": {
184
- "preset": "loop-agent-standard-verify",
185
- "cwd": ".",
186
- "timeoutMs": 300000
187
- }
188
- }
189
- ]
190
- }
1
+ {
2
+ "$schema": "./agent-dag.schema.json",
3
+ "version": 2,
4
+ "title": "Agent DAG base template",
5
+ "objective": "Replace with the concrete objective for this DAG run.",
6
+ "successCriteria": [
7
+ "Each node returns output matching its outputContract.",
8
+ "Same-rank read-only scouts run in parallel when their outputs are independent.",
9
+ "Only Pi nodes with toolProfile=write or explicitly enabled Cursor implementer nodes may write files, and only inside declared writeSet.",
10
+ "Pi nodes remain read-only unless toolProfile=write is explicitly selected for a bounded writer node.",
11
+ "Read-only nodes must not write root artifacts/**; root artifacts/ is not a per-node scratchpad."
12
+ ],
13
+ "globalConstraints": [
14
+ "Replace every REPLACE/WITH/... placeholder with concrete repo paths before execution; do not leave template placeholders in production DAG JSON.",
15
+ "Do not commit runtime traces under .harness/dag-runs/.",
16
+ "Do not write generated DAG input specs into .harness/dag-runs/active/.",
17
+ "Use a platform-native temp path such as <temp-dir>/<topic>-dag.json for one-off DAG specs; commit reusable templates under examples/ or docs/templates/.",
18
+ "Actual file operation paths must be macOS/Windows compatible; use / only for stable repo refs, JSON/Markdown evidence refs, and glob conventions.",
19
+ "Every task must explicitly declare executor; defaults.executor is schema-only and not a runtime fallback.",
20
+ "Prefer same-rank parallel read-only scouts over serial chains when outputs are independent.",
21
+ "Add depends_on only when a child truly needs upstream output; challenge single-chain topologies during review.",
22
+ "Same-rank exclusive writeSet entries must be disjoint.",
23
+ "Do not add provider fields to DAG JSON; provider routing is executor-owned.",
24
+ "Use executorModels for model routing; do not add defaults.model or legacy top-level models.",
25
+ "Every task must declare outputContract describing the expected Markdown/output shape.",
26
+ "exclusive implementer nodes must use narrow, concrete writeSet and allowedPaths; never keep broad write permissions such as ** or repo root.",
27
+ "Cursor is optional; no-Cursor environments should use Pi read-only scouts and Pi writer nodes with toolProfile=write."
28
+ ],
29
+ "defaults": {
30
+ "executor": "pi",
31
+ "contextProfile": "slim",
32
+ "skills": [
33
+ "ai-engineering-context"
34
+ ],
35
+ "writePolicy": "read-only"
36
+ },
37
+ "skillsByRole": {
38
+ "planner": [
39
+ "loop-agent"
40
+ ],
41
+ "scout": [
42
+ "ai-engineering-context"
43
+ ],
44
+ "implementer": [
45
+ "verification-before-completion"
46
+ ],
47
+ "reviewer": [
48
+ "requesting-code-review"
49
+ ],
50
+ "verifier": [
51
+ "verification-before-completion",
52
+ "systematic-debugging"
53
+ ],
54
+ "closeout": [
55
+ "loop-agent",
56
+ "verification-before-completion"
57
+ ]
58
+ },
59
+ "executorModels": {
60
+ "pi": {
61
+ "LOW": "gpt-5.3-codex-spark",
62
+ "MED": "glm-5.2",
63
+ "HIGH": "gpt-5.5"
64
+ }
65
+ },
66
+ "tasks": [
67
+ {
68
+ "id": "contract-pi",
69
+ "depends_on": [],
70
+ "complexity": "MED",
71
+ "executor": "pi",
72
+ "role": "planner",
73
+ "writePolicy": "read-only",
74
+ "allowedPaths": [
75
+ "**"
76
+ ],
77
+ "forbiddenPaths": [
78
+ ".harness/**",
79
+ "artifacts/**"
80
+ ],
81
+ "outputContract": "Plain Markdown implementation contract; no file writes.",
82
+ "subtask_prompt": "Read the task inputs and return a concise implementation contract: scope, risks, narrow write boundaries, parallel scout opportunities, and verification expectations. Do not edit files."
83
+ },
84
+ {
85
+ "id": "scout-src",
86
+ "depends_on": [
87
+ "contract-pi"
88
+ ],
89
+ "complexity": "LOW",
90
+ "executor": "pi",
91
+ "role": "scout",
92
+ "writePolicy": "read-only",
93
+ "allowedPaths": [
94
+ "REPLACE/WITH/SOURCE/PATH/**"
95
+ ],
96
+ "forbiddenPaths": [
97
+ ".harness/**",
98
+ "artifacts/**"
99
+ ],
100
+ "outputContract": "Plain Markdown source reconnaissance summary; no file writes.",
101
+ "subtask_prompt": "Perform read-only source reconnaissance. List relevant source files, existing patterns, and risks. Do not edit files."
102
+ },
103
+ {
104
+ "id": "scout-tests",
105
+ "depends_on": [
106
+ "contract-pi"
107
+ ],
108
+ "complexity": "LOW",
109
+ "executor": "pi",
110
+ "role": "scout",
111
+ "writePolicy": "read-only",
112
+ "allowedPaths": [
113
+ "REPLACE/WITH/TEST/PATH/**"
114
+ ],
115
+ "forbiddenPaths": [
116
+ ".harness/**",
117
+ "artifacts/**"
118
+ ],
119
+ "outputContract": "Plain Markdown test coverage reconnaissance summary; no file writes.",
120
+ "subtask_prompt": "Perform read-only test reconnaissance. List relevant tests, coverage gaps, and verification commands. Do not edit files."
121
+ },
122
+ {
123
+ "id": "implement-pi",
124
+ "depends_on": [
125
+ "scout-src",
126
+ "scout-tests"
127
+ ],
128
+ "complexity": "HIGH",
129
+ "executor": "pi",
130
+ "role": "implementer",
131
+ "writePolicy": "exclusive",
132
+ "writeSet": [
133
+ "REPLACE/WITH/ALLOWED/PATH/**"
134
+ ],
135
+ "allowedPaths": [
136
+ "REPLACE/WITH/ALLOWED/PATH/**"
137
+ ],
138
+ "forbiddenPaths": [
139
+ ".harness/**",
140
+ "artifacts/**"
141
+ ],
142
+ "outputContract": "Implementation summary with changed files, tests run, and risks.",
143
+ "subtask_prompt": "Implement the approved change within writeSet only. Keep the change minimal and update relevant tests/docs inside allowedPaths.",
144
+ "toolProfile": "write"
145
+ },
146
+ {
147
+ "id": "review-pi",
148
+ "depends_on": [
149
+ "implement-pi"
150
+ ],
151
+ "complexity": "MED",
152
+ "executor": "pi",
153
+ "role": "reviewer",
154
+ "writePolicy": "read-only",
155
+ "allowedPaths": [
156
+ "**"
157
+ ],
158
+ "forbiddenPaths": [
159
+ ".harness/**",
160
+ "artifacts/**"
161
+ ],
162
+ "outputContract": "Plain Markdown review summary; no file writes.",
163
+ "subtask_prompt": "Review upstream implementation output against the contract. Identify drift, missing tests, and residual risks. Do not edit files."
164
+ },
165
+ {
166
+ "id": "verify-shell",
167
+ "depends_on": [
168
+ "implement-pi"
169
+ ],
170
+ "complexity": "LOW",
171
+ "executor": "shell",
172
+ "role": "verifier",
173
+ "writePolicy": "read-only",
174
+ "allowedPaths": [
175
+ "./**"
176
+ ],
177
+ "forbiddenPaths": [
178
+ ".harness/**",
179
+ "artifacts/**"
180
+ ],
181
+ "outputContract": "Archived shell command stdout/stderr with exit codes; no worktree writes.",
182
+ "subtask_prompt": "Run deterministic loop-agent verification commands and archive outputs.",
183
+ "shell": {
184
+ "preset": "loop-agent-standard-verify",
185
+ "cwd": ".",
186
+ "timeoutMs": 300000
187
+ }
188
+ }
189
+ ]
190
+ }