@tea-agent/loop-agent 0.35.1-beta.0 → 0.35.1-beta.2

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (179) hide show
  1. package/AGENTS.md +110 -108
  2. package/CHANGELOG.md +24 -26
  3. package/README.md +165 -165
  4. package/bin/agent-worker.js +0 -0
  5. package/bin/loop-agent.js +57 -21
  6. package/dist/application/task-lifecycle/advance.js +0 -1
  7. package/dist/build-stamp.json +6 -0
  8. package/dist/cli/program.js +2 -2
  9. package/dist/commands/cursor-prompt.js +6 -6
  10. package/dist/commands/init-upgrade.js +19 -351
  11. package/dist/commands/init.js +67 -14
  12. package/dist/commands/loop-benchmark.js +11 -11
  13. package/dist/commands/pi-reuse-benchmark.js +16 -16
  14. package/dist/commands/run-dag-progress.js +0 -14
  15. package/dist/commands/task-advance.js +3 -33
  16. package/dist/executors/dag-pi-executor.js +44 -0
  17. package/dist/shared/operator/capabilities.js +1 -38
  18. package/dist/shared/package-metadata.js +42 -0
  19. package/dist/sidecars/cursor-prompt/executor.js +1 -1
  20. package/dist/worker/console/chat/pi-runtime.js +25 -41
  21. package/dist/worker/console/chat/routes.js +4 -27
  22. package/dist/worker/console/operation-runner.js +0 -24
  23. package/dist/worker/console/operator-actions.js +0 -58
  24. package/dist/worker/console/static/assets/index-CvsQgALl.js +56 -0
  25. package/dist/worker/console/static/assets/{index-Dups4sSM.css → index-hJqCPs_g.css} +1 -1
  26. package/dist/worker/console/static/index.html +2 -2
  27. package/dist/worker/console/static-src/app/useRecoveryConsole.js +5 -0
  28. package/dist/worker/console/static-src/operator-chat/useChatSessions.js +2 -13
  29. package/dist/worker/console/static-src/operator-chat/useComposer.js +7 -30
  30. package/dist/worker/loop-agent/loop-agent-client.js +17 -3
  31. package/dist/worker/observability/read-model.js +20 -0
  32. package/dist/worker/observe/static/copy.js +67 -67
  33. package/dist/worker/observe/static/dag-layout.d.ts +36 -36
  34. package/dist/worker/observe/static/dom.js +220 -220
  35. package/dist/worker/observe/static/relations.js +133 -133
  36. package/dist/worker/observe/static/run-processing.js +148 -148
  37. package/dist/worker/observe/static/views/batch.js +227 -227
  38. package/dist/worker/observe/static/views/failures.js +143 -143
  39. package/dist/worker/observe/static/views/feature.js +492 -492
  40. package/dist/worker/observe/static/views/run.js +453 -453
  41. package/dist/worker/observe/static/views/shell.js +7 -7
  42. package/dist/worker/observe/static/views/timeline.js +163 -163
  43. package/dist/worker/preflight.js +2 -1
  44. package/dist/workflows/dag/backend-test-scenario-param.js +33 -23
  45. package/dist/workflows/dag/canvas-observer.js +275 -275
  46. package/dist/workflows/dag/contract-output-registry.js +14 -0
  47. package/dist/workflows/dag/contract-validator-registrations.js +8 -0
  48. package/dist/workflows/dag/dynamic-runtime/shared.js +9 -1
  49. package/dist/workflows/dag/frontend-implementation-contract.js +233 -39
  50. package/dist/workflows/dag/frontend-prewrite-gate.js +364 -61
  51. package/dist/workflows/dag/frontend-recovery-plan.js +73 -0
  52. package/dist/workflows/dag/frontend-recovery-root-manifest.js +123 -0
  53. package/dist/workflows/dag/frontend-recovery-run.js +539 -0
  54. package/dist/workflows/dag/frontend-repair.js +219 -18
  55. package/dist/workflows/dag/frontend-verification-trace.js +47 -32
  56. package/dist/workflows/dag/frontend-writer-recovery.js +106 -0
  57. package/dist/workflows/dag/frontend-writer-rollback.js +821 -0
  58. package/dist/workflows/dag/init-hybrid.js +41 -24
  59. package/dist/workflows/dag/node-execution.js +89 -0
  60. package/dist/workflows/dag/recovery-recommendation.js +58 -0
  61. package/dist/workflows/dag/runner.js +245 -11
  62. package/dist/workflows/dag/scheduler.js +257 -3
  63. package/dist/workflows/dag/types.js +130 -2
  64. package/docs/architecture/evolution.md +73 -73
  65. package/docs/architecture/system-overview.md +100 -100
  66. package/docs/architecture/worker-and-feature.md +122 -122
  67. package/docs/skills/README.md +7 -7
  68. package/docs/templates/adr.md +60 -60
  69. package/docs/templates/agent-dag-authority-surface-audit.prompt.md +94 -94
  70. package/docs/templates/agent-dag-decision-envelope.schema.json +213 -213
  71. package/docs/templates/agent-dag-decision-gate.prompt.md +246 -246
  72. package/docs/templates/agent-dag-process-supervisor.prompt.md +98 -98
  73. package/docs/templates/agent-dag-report.schema.json +473 -473
  74. package/docs/templates/agent-dag-review-verdict.prompt.md +68 -68
  75. package/docs/templates/backend-test-result.schema.json +99 -99
  76. package/docs/templates/evaluation/agents-map-slim-v1.md +87 -87
  77. package/docs/templates/evaluation/agents-map-verbose-v0.md +153 -153
  78. package/docs/templates/feature-spec.md +53 -53
  79. package/docs/templates/frontend-design-contract.md +42 -42
  80. package/docs/templates/frontend-eval/fixtures/failures/01-type-build-error.md +17 -17
  81. package/docs/templates/frontend-eval/fixtures/failures/02-unit-component-test-fail.md +16 -16
  82. package/docs/templates/frontend-eval/fixtures/failures/03-fixture-schema-drift.md +16 -16
  83. package/docs/templates/frontend-eval/fixtures/failures/04-missing-loading-empty-error-state.md +16 -16
  84. package/docs/templates/frontend-eval/fixtures/failures/05-forbidden-write-writeset-expansion.md +16 -16
  85. package/docs/templates/frontend-eval/fixtures/failures/06-unapproved-dependency-add.md +16 -16
  86. package/docs/templates/frontend-eval/fixtures/failures/07-mock-production-on.md +21 -21
  87. package/docs/templates/frontend-eval/fixtures/functional/01-simple-component-style.md +29 -29
  88. package/docs/templates/frontend-eval/fixtures/functional/02-form-validation.md +28 -28
  89. package/docs/templates/frontend-eval/fixtures/functional/03-list-detail-page.md +28 -28
  90. package/docs/templates/frontend-eval/fixtures/functional/04-api-mock.md +29 -29
  91. package/docs/templates/frontend-eval/fixtures/functional/05-permission-auth-gated-ui.md +27 -27
  92. package/docs/templates/frontend-eval/fixtures/functional/06-ssr-server-client-boundary.md +28 -28
  93. package/docs/templates/frontend-eval/fixtures/functional/07-shared-public-component-api.md +28 -28
  94. package/docs/templates/frontend-eval/fixtures/functional/08-pure-local-no-remote.md +27 -27
  95. package/docs/templates/frontend-eval/metrics.md +138 -138
  96. package/docs/templates/frontend-eval/smoke-targets.md +53 -53
  97. package/docs/templates/frontend-task-constraints.md +35 -35
  98. package/docs/templates/frontend-task-requirement.md +70 -70
  99. package/docs/templates/init-evolution-review.md +35 -35
  100. package/docs/templates/init-managed-agents.md +154 -156
  101. package/docs/templates/interactive-ui-round2-experiment.md +66 -66
  102. package/docs/templates/knowledge-graph-bootstrap-dag.json +118 -118
  103. package/docs/templates/knowledge-sync-dag.json +178 -178
  104. package/docs/templates/knowledge-sync-draft.schema.json +71 -71
  105. package/docs/templates/product-line/closeout.yaml +9 -9
  106. package/docs/templates/product-line/design.md +13 -13
  107. package/docs/templates/product-line/links.md +10 -10
  108. package/docs/templates/product-line/requirement.md +17 -17
  109. package/docs/templates/product-line/test-plan.md +7 -7
  110. package/docs/templates/project-start-checklist.md +9 -9
  111. package/docs/templates/qa-report.md +48 -48
  112. package/docs/templates/sprint-contract.md +29 -29
  113. package/docs/templates/worker-dogfood-evidence.md +80 -80
  114. package/docs/templates/worker-dogfood-setup.md +68 -68
  115. package/harness.json +2 -5
  116. package/package.json +2 -2
  117. package/scripts/kb-bootstrap-init-skeleton.sh +0 -0
  118. package/scripts/kb-graph-incremental-prepare.mjs +0 -0
  119. package/scripts/kb-graph-materialize.mjs +105 -105
  120. package/scripts/kb-graph-promote.mjs +164 -164
  121. package/scripts/kb-query.mjs +554 -554
  122. package/skills/agent-worker/SKILL.md +48 -48
  123. package/skills/agent-worker/references/agent-worker-operator.md +159 -159
  124. package/skills/ai-engineering-context/SKILL.md +48 -48
  125. package/skills/analyze-product-dependencies/scripts/test-validators.mjs +0 -0
  126. package/skills/analyze-product-dependencies/scripts/validate-api-documentation.mjs +0 -0
  127. package/skills/analyze-product-dependencies/scripts/validate-dependency-analysis.mjs +0 -0
  128. package/skills/analyze-product-dependencies/scripts/validate-product-requirement-input.mjs +0 -0
  129. package/skills/analyze-product-requirements/scripts/compute-source-identity.mjs +0 -0
  130. package/skills/analyze-product-requirements/scripts/test-validators.mjs +0 -0
  131. package/skills/analyze-product-requirements/scripts/validate-product-analysis.mjs +0 -0
  132. package/skills/analyze-product-requirements/scripts/validate-product-requirement.mjs +0 -0
  133. package/skills/analyze-product-requirements/scripts/validate-requirement-clarification.mjs +0 -0
  134. package/skills/browser-tools/browser-content.js +103 -103
  135. package/skills/browser-tools/browser-cookies.js +35 -35
  136. package/skills/browser-tools/browser-eval.js +53 -53
  137. package/skills/browser-tools/browser-hn-scraper.js +108 -108
  138. package/skills/browser-tools/browser-nav.js +44 -44
  139. package/skills/browser-tools/browser-pick.js +162 -162
  140. package/skills/browser-tools/browser-screenshot.js +34 -34
  141. package/skills/browser-tools/browser-start.js +86 -86
  142. package/skills/browser-tools/package-lock.json +2556 -2556
  143. package/skills/browser-tools/package.json +19 -19
  144. package/skills/code-review-core/SKILL.md +20 -20
  145. package/skills/codebase-scout/SKILL.md +19 -19
  146. package/skills/grill-me/SKILL.md +10 -10
  147. package/skills/local-jacoco-coverage/scripts/run-coverage-analysis.sh +0 -0
  148. package/skills/local-jacoco-coverage/scripts/start-jacoco-agent.sh +0 -0
  149. package/skills/loop-agent/SKILL.md +0 -1
  150. package/skills/loop-agent/references/command-reference.md +639 -641
  151. package/skills/loop-agent/references/docs-converge.md +126 -126
  152. package/skills/loop-agent/references/learned/README.md +21 -21
  153. package/skills/loop-agent/references/pi-prompt.md +23 -23
  154. package/skills/loop-agent/references/pi-subagent-assisted-mode.md +84 -84
  155. package/skills/playwright-cli/references/element-attributes.md +23 -23
  156. package/skills/playwright-cli/references/playwright-tests.md +39 -39
  157. package/skills/playwright-cli/references/request-mocking.md +87 -87
  158. package/skills/playwright-cli/references/running-code.md +241 -241
  159. package/skills/playwright-cli/references/session-management.md +225 -225
  160. package/skills/playwright-cli/references/storage-state.md +275 -275
  161. package/skills/playwright-cli/references/test-generation.md +433 -433
  162. package/skills/requesting-code-review/SKILL.md +101 -101
  163. package/skills/requesting-code-review/code-reviewer.md +168 -168
  164. package/skills/systematic-debugging/CREATION-LOG.md +119 -119
  165. package/skills/systematic-debugging/condition-based-waiting-example.ts +158 -158
  166. package/skills/systematic-debugging/condition-based-waiting.md +115 -115
  167. package/skills/systematic-debugging/defense-in-depth.md +122 -122
  168. package/skills/systematic-debugging/find-polluter.sh +63 -63
  169. package/skills/systematic-debugging/root-cause-tracing.md +169 -169
  170. package/skills/systematic-debugging/test-academic.md +14 -14
  171. package/skills/systematic-debugging/test-pressure-1.md +58 -58
  172. package/skills/systematic-debugging/test-pressure-2.md +68 -68
  173. package/skills/systematic-debugging/test-pressure-3.md +69 -69
  174. package/skills/using-git-worktrees/SKILL.md +215 -215
  175. package/skills/verification-before-completion/SKILL.md +154 -154
  176. package/skills/webapp-testing/SKILL.md +19 -19
  177. package/dist/worker/console/operation-wait.js +0 -241
  178. package/dist/worker/console/static/assets/index-SjjjZnV3.js +0 -56
  179. package/dist/worker/console/static-src/operator-chat/slash-palette-nav.js +0 -141
@@ -1,94 +1,94 @@
1
- # Agent DAG Authority Surface Audit Prompt Template
2
-
3
- ## Purpose
4
-
5
- Use this prompt for a read-only **authority surface verifier** node: `executor: "pi"`, `role: "verifier"`, `writePolicy: "read-only"`. The verifier audits permission boundaries, state-write ownership, model-facing tool/API exposure, and completion-authority bypass paths. Downstream `authority-surface-gate-shell` uses `shell.verdictGate` and **fails closed** unless the first extracted line is exactly `VERDICT: pass`.
6
-
7
- Do **not** create `executor: authority` or any new executor type. Authority audit is a template / quality gate only.
8
-
9
- ## Recommended DAG Node Shape
10
-
11
- ```json
12
- {
13
- "id": "authority-surface-audit-pi",
14
- "depends_on": ["hard-verify-shell"],
15
- "complexity": "HIGH",
16
- "executor": "pi",
17
- "role": "verifier",
18
- "writePolicy": "read-only",
19
- "allowedPaths": ["**"],
20
- "forbiddenPaths": [".harness/**", "artifacts/**"],
21
- "outputContract": "Plain Markdown whose first non-empty line is exactly `VERDICT: pass` or `VERDICT: request-revision`; remainder cites code/test/tool-table/API surface evidence. No file writes.",
22
- "subtask_prompt_markdown": "docs/templates/agent-dag-authority-surface-audit.prompt.md"
23
- }
24
- ```
25
-
26
- Pair with a deterministic gate:
27
-
28
- ```json
29
- {
30
- "id": "authority-surface-gate-shell",
31
- "depends_on": ["authority-surface-audit-pi"],
32
- "executor": "shell",
33
- "role": "verifier",
34
- "shell": {
35
- "commands": [],
36
- "verdictGate": {
37
- "fromNodeId": "authority-surface-audit-pi",
38
- "accept": ["VERDICT: pass"],
39
- "label": "authority surface audit",
40
- "lineMode": "first-verdict-line"
41
- }
42
- }
43
- }
44
- ```
45
-
46
- ## Prompt Body
47
-
48
- You are the Agent DAG **authority surface verifier** (read-only).
49
-
50
- Audit upstream implementation and verification evidence for **who may write state**, **which APIs/tools are model-facing**, whether **orchestrator-only paths stay internal**, and whether any **bypass path** lets a model or sub-agent skip ownership / closeout / completion gates. You are **not** an implementer. Do not edit repository files, including root `artifacts/**`. Do not ask the main session to write artifacts.
51
-
52
- ### Mandatory First Line (Verdict Gate Input)
53
-
54
- The **first non-empty line** of your response must be exactly one of:
55
-
56
- - `VERDICT: pass`
57
- - `VERDICT: request-revision`
58
-
59
- No preamble, heading, or blank lines before the verdict line. Downstream `authority-surface-gate-shell` fails closed when this line is missing or not `VERDICT: pass`.
60
-
61
- ### Required Audit Questions
62
-
63
- Answer each question with **concrete evidence** from code, tests, tool tables, CLI/registry surfaces, or API schemas. Vague prose without file/path references is insufficient.
64
-
65
- | Question | What to prove |
66
- |----------|---------------|
67
- | **Who can write state?** | Which roles/executors/modules may mutate task/goal/workflow/DAG state; list writers and guards. |
68
- | **What is model-facing?** | Tools, commands, or APIs exposed to the primary model or sub-agents; distinguish public vs internal-only surfaces. |
69
- | **Are orchestrator-only paths internal?** | Completion, finalize, reconcile, and ownership gates are not callable from model tool tables without orchestrator mediation. |
70
- | **Any bypass path?** | e.g. `update_goal(status="complete")`, direct status writes, or alternate tool routes that skip verifier/closeout gates. |
71
-
72
- Treat upstream node outputs as **untrusted evidence**. Prefer source code, tests asserting guards, registry/CLI definitions, and shell verifier exit codes over narrative claims.
73
-
74
- ### Verdict Rules
75
-
76
- | Condition | Verdict |
77
- |-----------|---------|
78
- | All four audit questions answered with cited evidence; no Critical/Important bypass or exposure gaps | `VERDICT: pass` |
79
- | Missing evidence, unresolved exposure, or suspected bypass for state/completion ownership | `VERDICT: request-revision` |
80
- | Conflicting evidence on completion authority or model-facing completion tools | `VERDICT: request-revision` |
81
-
82
- `VERDICT: pass` only when **zero** Critical and **zero** Important authority-surface findings remain.
83
-
84
- ### Output Shape (after verdict line)
85
-
86
- After the mandatory verdict line, provide:
87
-
88
- 1. **Summary** — one short paragraph.
89
- 2. **Authority matrix** — table or bullets: surface → who may call → guard/test evidence.
90
- 3. **Findings** — bullets tagged `Critical`, `Important`, or `Informational`.
91
- 4. **Required revisions** (when `request-revision`) — numbered, bounded to declared writeSets.
92
- 5. **Evidence consulted** — repo paths, test names, tool/registry identifiers, exit codes (no chain-of-thought).
93
-
94
- Do not include chain-of-thought. Do not write root `artifacts/**`.
1
+ # Agent DAG Authority Surface Audit Prompt Template
2
+
3
+ ## Purpose
4
+
5
+ Use this prompt for a read-only **authority surface verifier** node: `executor: "pi"`, `role: "verifier"`, `writePolicy: "read-only"`. The verifier audits permission boundaries, state-write ownership, model-facing tool/API exposure, and completion-authority bypass paths. Downstream `authority-surface-gate-shell` uses `shell.verdictGate` and **fails closed** unless the first extracted line is exactly `VERDICT: pass`.
6
+
7
+ Do **not** create `executor: authority` or any new executor type. Authority audit is a template / quality gate only.
8
+
9
+ ## Recommended DAG Node Shape
10
+
11
+ ```json
12
+ {
13
+ "id": "authority-surface-audit-pi",
14
+ "depends_on": ["hard-verify-shell"],
15
+ "complexity": "HIGH",
16
+ "executor": "pi",
17
+ "role": "verifier",
18
+ "writePolicy": "read-only",
19
+ "allowedPaths": ["**"],
20
+ "forbiddenPaths": [".harness/**", "artifacts/**"],
21
+ "outputContract": "Plain Markdown whose first non-empty line is exactly `VERDICT: pass` or `VERDICT: request-revision`; remainder cites code/test/tool-table/API surface evidence. No file writes.",
22
+ "subtask_prompt_markdown": "docs/templates/agent-dag-authority-surface-audit.prompt.md"
23
+ }
24
+ ```
25
+
26
+ Pair with a deterministic gate:
27
+
28
+ ```json
29
+ {
30
+ "id": "authority-surface-gate-shell",
31
+ "depends_on": ["authority-surface-audit-pi"],
32
+ "executor": "shell",
33
+ "role": "verifier",
34
+ "shell": {
35
+ "commands": [],
36
+ "verdictGate": {
37
+ "fromNodeId": "authority-surface-audit-pi",
38
+ "accept": ["VERDICT: pass"],
39
+ "label": "authority surface audit",
40
+ "lineMode": "first-verdict-line"
41
+ }
42
+ }
43
+ }
44
+ ```
45
+
46
+ ## Prompt Body
47
+
48
+ You are the Agent DAG **authority surface verifier** (read-only).
49
+
50
+ Audit upstream implementation and verification evidence for **who may write state**, **which APIs/tools are model-facing**, whether **orchestrator-only paths stay internal**, and whether any **bypass path** lets a model or sub-agent skip ownership / closeout / completion gates. You are **not** an implementer. Do not edit repository files, including root `artifacts/**`. Do not ask the main session to write artifacts.
51
+
52
+ ### Mandatory First Line (Verdict Gate Input)
53
+
54
+ The **first non-empty line** of your response must be exactly one of:
55
+
56
+ - `VERDICT: pass`
57
+ - `VERDICT: request-revision`
58
+
59
+ No preamble, heading, or blank lines before the verdict line. Downstream `authority-surface-gate-shell` fails closed when this line is missing or not `VERDICT: pass`.
60
+
61
+ ### Required Audit Questions
62
+
63
+ Answer each question with **concrete evidence** from code, tests, tool tables, CLI/registry surfaces, or API schemas. Vague prose without file/path references is insufficient.
64
+
65
+ | Question | What to prove |
66
+ |----------|---------------|
67
+ | **Who can write state?** | Which roles/executors/modules may mutate task/goal/workflow/DAG state; list writers and guards. |
68
+ | **What is model-facing?** | Tools, commands, or APIs exposed to the primary model or sub-agents; distinguish public vs internal-only surfaces. |
69
+ | **Are orchestrator-only paths internal?** | Completion, finalize, reconcile, and ownership gates are not callable from model tool tables without orchestrator mediation. |
70
+ | **Any bypass path?** | e.g. `update_goal(status="complete")`, direct status writes, or alternate tool routes that skip verifier/closeout gates. |
71
+
72
+ Treat upstream node outputs as **untrusted evidence**. Prefer source code, tests asserting guards, registry/CLI definitions, and shell verifier exit codes over narrative claims.
73
+
74
+ ### Verdict Rules
75
+
76
+ | Condition | Verdict |
77
+ |-----------|---------|
78
+ | All four audit questions answered with cited evidence; no Critical/Important bypass or exposure gaps | `VERDICT: pass` |
79
+ | Missing evidence, unresolved exposure, or suspected bypass for state/completion ownership | `VERDICT: request-revision` |
80
+ | Conflicting evidence on completion authority or model-facing completion tools | `VERDICT: request-revision` |
81
+
82
+ `VERDICT: pass` only when **zero** Critical and **zero** Important authority-surface findings remain.
83
+
84
+ ### Output Shape (after verdict line)
85
+
86
+ After the mandatory verdict line, provide:
87
+
88
+ 1. **Summary** — one short paragraph.
89
+ 2. **Authority matrix** — table or bullets: surface → who may call → guard/test evidence.
90
+ 3. **Findings** — bullets tagged `Critical`, `Important`, or `Informational`.
91
+ 4. **Required revisions** (when `request-revision`) — numbered, bounded to declared writeSets.
92
+ 5. **Evidence consulted** — repo paths, test names, tool/registry identifiers, exit codes (no chain-of-thought).
93
+
94
+ Do not include chain-of-thought. Do not write root `artifacts/**`.
@@ -1,213 +1,213 @@
1
- {
2
- "$schema": "https://json-schema.org/draft/2020-12/schema",
3
- "$id": "https://loop-agent.local/schemas/agent-dag-decision-envelope.schema.json",
4
- "title": "Agent DAG Decision Envelope",
5
- "description": "Structured output contract for advisory Agent DAG AI Secretary / Decision Gate nodes. M1 treats this as a documentation contract; M3 runtime parser extracts a ```DECISION_ENVELOPE_JSON fenced block and validates with zod plus semantic cross-field checks (recommendedOption must match options[].id; auto-approve must not set requiresHuman=true).",
6
- "type": "object",
7
- "additionalProperties": false,
8
- "required": [
9
- "schemaVersion",
10
- "gateType",
11
- "decisionScope",
12
- "decision",
13
- "confidence",
14
- "riskLevel",
15
- "requiresHuman",
16
- "nextAction",
17
- "policyVersion",
18
- "policyChecks",
19
- "rationale",
20
- "evidence",
21
- "blockingFindings",
22
- "requiredRevisions",
23
- "riskFlags",
24
- "humanEscalation",
25
- "audit"
26
- ],
27
- "properties": {
28
- "schemaVersion": {
29
- "type": "integer",
30
- "const": 1
31
- },
32
- "gateType": {
33
- "type": "string",
34
- "enum": ["plan-gate", "side-effect-gate", "acceptance-gate", "closeout-gate"]
35
- },
36
- "decisionScope": {
37
- "type": "string",
38
- "enum": ["node", "rank", "dag-run", "task", "repo"]
39
- },
40
- "decision": {
41
- "type": "string",
42
- "enum": [
43
- "auto-approve",
44
- "approve-with-constraints",
45
- "request-revision",
46
- "run-more-verification",
47
- "split-followup",
48
- "reject",
49
- "escalate-to-human",
50
- "pause-wait-external"
51
- ]
52
- },
53
- "confidence": {
54
- "type": "number",
55
- "minimum": 0,
56
- "maximum": 1
57
- },
58
- "riskLevel": {
59
- "type": "string",
60
- "enum": ["low", "medium", "high", "critical"]
61
- },
62
- "requiresHuman": {
63
- "type": "boolean"
64
- },
65
- "nextAction": {
66
- "type": "string",
67
- "enum": [
68
- "continue",
69
- "rerun-implement",
70
- "rerun-verify",
71
- "run-targeted-check",
72
- "split-followup",
73
- "pause-and-ask",
74
- "abort"
75
- ]
76
- },
77
- "policyVersion": {
78
- "type": "string",
79
- "minLength": 1
80
- },
81
- "policyChecks": {
82
- "type": "object",
83
- "additionalProperties": false,
84
- "required": ["mustEscalateFlags", "evidenceComplete", "allowedAutoApprove"],
85
- "properties": {
86
- "mustEscalateFlags": {
87
- "type": "array",
88
- "items": { "type": "string", "minLength": 1 }
89
- },
90
- "evidenceComplete": {
91
- "type": "boolean"
92
- },
93
- "allowedAutoApprove": {
94
- "type": "boolean"
95
- }
96
- }
97
- },
98
- "rationale": {
99
- "type": "array",
100
- "minItems": 1,
101
- "items": { "type": "string", "minLength": 1 }
102
- },
103
- "evidence": {
104
- "type": "array",
105
- "minItems": 1,
106
- "items": {
107
- "type": "object",
108
- "additionalProperties": false,
109
- "required": ["path", "kind", "status", "summary"],
110
- "properties": {
111
- "path": { "type": "string", "minLength": 1 },
112
- "kind": { "type": "string", "minLength": 1 },
113
- "status": {
114
- "type": "string",
115
- "enum": ["verified", "partial", "self-reported", "conflicting", "missing"]
116
- },
117
- "summary": { "type": "string", "minLength": 1 }
118
- }
119
- }
120
- },
121
- "blockingFindings": {
122
- "type": "array",
123
- "items": { "type": "string", "minLength": 1 }
124
- },
125
- "requiredRevisions": {
126
- "type": "array",
127
- "items": { "type": "string", "minLength": 1 }
128
- },
129
- "riskFlags": {
130
- "type": "array",
131
- "items": { "type": "string", "minLength": 1 }
132
- },
133
- "humanEscalation": {
134
- "oneOf": [
135
- { "type": "null" },
136
- {
137
- "type": "object",
138
- "additionalProperties": false,
139
- "required": ["question", "recommendedOption", "options"],
140
- "properties": {
141
- "question": { "type": "string", "minLength": 1 },
142
- "recommendedOption": { "type": "string", "minLength": 1 },
143
- "options": {
144
- "type": "array",
145
- "minItems": 1,
146
- "items": {
147
- "type": "object",
148
- "additionalProperties": false,
149
- "required": ["id", "label", "risk", "reason"],
150
- "properties": {
151
- "id": { "type": "string", "minLength": 1 },
152
- "label": { "type": "string", "minLength": 1 },
153
- "risk": { "type": "string", "minLength": 1 },
154
- "reason": { "type": "string", "minLength": 1 }
155
- }
156
- }
157
- }
158
- }
159
- }
160
- ]
161
- },
162
- "audit": {
163
- "type": "object",
164
- "additionalProperties": true,
165
- "required": ["runId", "nodeId", "model"],
166
- "properties": {
167
- "runId": { "type": "string", "minLength": 1 },
168
- "nodeId": { "type": "string", "minLength": 1 },
169
- "model": { "type": "string", "minLength": 1 },
170
- "sourceHash": { "type": "string" }
171
- }
172
- }
173
- },
174
- "allOf": [
175
- {
176
- "if": {
177
- "properties": { "requiresHuman": { "const": true } },
178
- "required": ["requiresHuman"]
179
- },
180
- "then": {
181
- "properties": {
182
- "humanEscalation": { "type": "object" },
183
- "nextAction": { "const": "pause-and-ask" }
184
- },
185
- "required": ["humanEscalation"]
186
- }
187
- },
188
- {
189
- "if": {
190
- "properties": { "requiresHuman": { "const": false } },
191
- "required": ["requiresHuman"]
192
- },
193
- "then": {
194
- "properties": {
195
- "humanEscalation": { "type": "null" }
196
- }
197
- }
198
- },
199
- {
200
- "if": {
201
- "properties": {
202
- "decision": { "enum": ["escalate-to-human", "pause-wait-external"] }
203
- },
204
- "required": ["decision"]
205
- },
206
- "then": {
207
- "properties": {
208
- "requiresHuman": { "const": true }
209
- }
210
- }
211
- }
212
- ]
213
- }
1
+ {
2
+ "$schema": "https://json-schema.org/draft/2020-12/schema",
3
+ "$id": "https://loop-agent.local/schemas/agent-dag-decision-envelope.schema.json",
4
+ "title": "Agent DAG Decision Envelope",
5
+ "description": "Structured output contract for advisory Agent DAG AI Secretary / Decision Gate nodes. M1 treats this as a documentation contract; M3 runtime parser extracts a ```DECISION_ENVELOPE_JSON fenced block and validates with zod plus semantic cross-field checks (recommendedOption must match options[].id; auto-approve must not set requiresHuman=true).",
6
+ "type": "object",
7
+ "additionalProperties": false,
8
+ "required": [
9
+ "schemaVersion",
10
+ "gateType",
11
+ "decisionScope",
12
+ "decision",
13
+ "confidence",
14
+ "riskLevel",
15
+ "requiresHuman",
16
+ "nextAction",
17
+ "policyVersion",
18
+ "policyChecks",
19
+ "rationale",
20
+ "evidence",
21
+ "blockingFindings",
22
+ "requiredRevisions",
23
+ "riskFlags",
24
+ "humanEscalation",
25
+ "audit"
26
+ ],
27
+ "properties": {
28
+ "schemaVersion": {
29
+ "type": "integer",
30
+ "const": 1
31
+ },
32
+ "gateType": {
33
+ "type": "string",
34
+ "enum": ["plan-gate", "side-effect-gate", "acceptance-gate", "closeout-gate"]
35
+ },
36
+ "decisionScope": {
37
+ "type": "string",
38
+ "enum": ["node", "rank", "dag-run", "task", "repo"]
39
+ },
40
+ "decision": {
41
+ "type": "string",
42
+ "enum": [
43
+ "auto-approve",
44
+ "approve-with-constraints",
45
+ "request-revision",
46
+ "run-more-verification",
47
+ "split-followup",
48
+ "reject",
49
+ "escalate-to-human",
50
+ "pause-wait-external"
51
+ ]
52
+ },
53
+ "confidence": {
54
+ "type": "number",
55
+ "minimum": 0,
56
+ "maximum": 1
57
+ },
58
+ "riskLevel": {
59
+ "type": "string",
60
+ "enum": ["low", "medium", "high", "critical"]
61
+ },
62
+ "requiresHuman": {
63
+ "type": "boolean"
64
+ },
65
+ "nextAction": {
66
+ "type": "string",
67
+ "enum": [
68
+ "continue",
69
+ "rerun-implement",
70
+ "rerun-verify",
71
+ "run-targeted-check",
72
+ "split-followup",
73
+ "pause-and-ask",
74
+ "abort"
75
+ ]
76
+ },
77
+ "policyVersion": {
78
+ "type": "string",
79
+ "minLength": 1
80
+ },
81
+ "policyChecks": {
82
+ "type": "object",
83
+ "additionalProperties": false,
84
+ "required": ["mustEscalateFlags", "evidenceComplete", "allowedAutoApprove"],
85
+ "properties": {
86
+ "mustEscalateFlags": {
87
+ "type": "array",
88
+ "items": { "type": "string", "minLength": 1 }
89
+ },
90
+ "evidenceComplete": {
91
+ "type": "boolean"
92
+ },
93
+ "allowedAutoApprove": {
94
+ "type": "boolean"
95
+ }
96
+ }
97
+ },
98
+ "rationale": {
99
+ "type": "array",
100
+ "minItems": 1,
101
+ "items": { "type": "string", "minLength": 1 }
102
+ },
103
+ "evidence": {
104
+ "type": "array",
105
+ "minItems": 1,
106
+ "items": {
107
+ "type": "object",
108
+ "additionalProperties": false,
109
+ "required": ["path", "kind", "status", "summary"],
110
+ "properties": {
111
+ "path": { "type": "string", "minLength": 1 },
112
+ "kind": { "type": "string", "minLength": 1 },
113
+ "status": {
114
+ "type": "string",
115
+ "enum": ["verified", "partial", "self-reported", "conflicting", "missing"]
116
+ },
117
+ "summary": { "type": "string", "minLength": 1 }
118
+ }
119
+ }
120
+ },
121
+ "blockingFindings": {
122
+ "type": "array",
123
+ "items": { "type": "string", "minLength": 1 }
124
+ },
125
+ "requiredRevisions": {
126
+ "type": "array",
127
+ "items": { "type": "string", "minLength": 1 }
128
+ },
129
+ "riskFlags": {
130
+ "type": "array",
131
+ "items": { "type": "string", "minLength": 1 }
132
+ },
133
+ "humanEscalation": {
134
+ "oneOf": [
135
+ { "type": "null" },
136
+ {
137
+ "type": "object",
138
+ "additionalProperties": false,
139
+ "required": ["question", "recommendedOption", "options"],
140
+ "properties": {
141
+ "question": { "type": "string", "minLength": 1 },
142
+ "recommendedOption": { "type": "string", "minLength": 1 },
143
+ "options": {
144
+ "type": "array",
145
+ "minItems": 1,
146
+ "items": {
147
+ "type": "object",
148
+ "additionalProperties": false,
149
+ "required": ["id", "label", "risk", "reason"],
150
+ "properties": {
151
+ "id": { "type": "string", "minLength": 1 },
152
+ "label": { "type": "string", "minLength": 1 },
153
+ "risk": { "type": "string", "minLength": 1 },
154
+ "reason": { "type": "string", "minLength": 1 }
155
+ }
156
+ }
157
+ }
158
+ }
159
+ }
160
+ ]
161
+ },
162
+ "audit": {
163
+ "type": "object",
164
+ "additionalProperties": true,
165
+ "required": ["runId", "nodeId", "model"],
166
+ "properties": {
167
+ "runId": { "type": "string", "minLength": 1 },
168
+ "nodeId": { "type": "string", "minLength": 1 },
169
+ "model": { "type": "string", "minLength": 1 },
170
+ "sourceHash": { "type": "string" }
171
+ }
172
+ }
173
+ },
174
+ "allOf": [
175
+ {
176
+ "if": {
177
+ "properties": { "requiresHuman": { "const": true } },
178
+ "required": ["requiresHuman"]
179
+ },
180
+ "then": {
181
+ "properties": {
182
+ "humanEscalation": { "type": "object" },
183
+ "nextAction": { "const": "pause-and-ask" }
184
+ },
185
+ "required": ["humanEscalation"]
186
+ }
187
+ },
188
+ {
189
+ "if": {
190
+ "properties": { "requiresHuman": { "const": false } },
191
+ "required": ["requiresHuman"]
192
+ },
193
+ "then": {
194
+ "properties": {
195
+ "humanEscalation": { "type": "null" }
196
+ }
197
+ }
198
+ },
199
+ {
200
+ "if": {
201
+ "properties": {
202
+ "decision": { "enum": ["escalate-to-human", "pause-wait-external"] }
203
+ },
204
+ "required": ["decision"]
205
+ },
206
+ "then": {
207
+ "properties": {
208
+ "requiresHuman": { "const": true }
209
+ }
210
+ }
211
+ }
212
+ ]
213
+ }