memorix 1.2.2 → 1.2.3

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (139) hide show
  1. package/CHANGELOG.md +7 -0
  2. package/TEAM.md +86 -86
  3. package/dist/cli/index.js +34 -18
  4. package/dist/cli/index.js.map +1 -1
  5. package/dist/index.js +17 -8
  6. package/dist/index.js.map +1 -1
  7. package/dist/maintenance-runner.js.map +1 -1
  8. package/dist/memcode-runtime/CHANGELOG.md +7 -0
  9. package/dist/sdk.js +17 -8
  10. package/dist/sdk.js.map +1 -1
  11. package/docs/DESIGN_DECISIONS.md +357 -357
  12. package/docs/dev-log/progress.txt +18 -8
  13. package/package.json +1 -1
  14. package/plugins/codex/memorix/.codex-plugin/plugin.json +1 -1
  15. package/src/audit/index.ts +156 -156
  16. package/src/cli/commands/audit-list.ts +89 -89
  17. package/src/cli/commands/background.ts +659 -659
  18. package/src/cli/commands/formation.ts +48 -48
  19. package/src/cli/commands/git-hook-install.ts +111 -111
  20. package/src/cli/commands/handoff.ts +54 -54
  21. package/src/cli/commands/hooks-status.ts +63 -63
  22. package/src/cli/commands/ingest-commit.ts +153 -153
  23. package/src/cli/commands/ingest-image.ts +66 -66
  24. package/src/cli/commands/ingest-log.ts +180 -180
  25. package/src/cli/commands/ingest.ts +44 -44
  26. package/src/cli/commands/integrate-shared.ts +15 -15
  27. package/src/cli/commands/lock.ts +82 -82
  28. package/src/cli/commands/message.ts +104 -104
  29. package/src/cli/commands/poll.ts +58 -58
  30. package/src/cli/commands/purge-all-memory.ts +85 -85
  31. package/src/cli/commands/purge-project-memory.ts +83 -83
  32. package/src/cli/commands/reasoning.ts +118 -118
  33. package/src/cli/commands/serve-shared.ts +118 -118
  34. package/src/cli/commands/session.ts +15 -7
  35. package/src/cli/commands/skills.ts +114 -114
  36. package/src/cli/commands/task.ts +167 -167
  37. package/src/cli/commands/transfer.ts +47 -47
  38. package/src/cli/commands/uninstall-project-artifacts.ts +85 -85
  39. package/src/cli/tui/ChatView.tsx +234 -234
  40. package/src/cli/tui/CommandBar.tsx +312 -312
  41. package/src/cli/tui/ContextRail.tsx +118 -118
  42. package/src/cli/tui/HeaderBar.tsx +72 -72
  43. package/src/cli/tui/LogoBanner.tsx +51 -51
  44. package/src/cli/tui/Sidebar.tsx +179 -179
  45. package/src/cli/tui/index.ts +41 -41
  46. package/src/cli/tui/markdown-render.tsx +371 -371
  47. package/src/cli/tui/session-service.ts +3 -2
  48. package/src/cli/tui/use-mouse.ts +157 -157
  49. package/src/cli/tui/useNavigation.ts +56 -56
  50. package/src/cli/update-checker.ts +211 -211
  51. package/src/cli/version.ts +7 -7
  52. package/src/cli/workbench.ts +1 -1
  53. package/src/compact/token-budget.ts +74 -74
  54. package/src/dashboard/project-classification.ts +64 -64
  55. package/src/embedding/fastembed-provider.ts +142 -142
  56. package/src/embedding/transformers-provider.ts +111 -111
  57. package/src/git/extractor.ts +209 -209
  58. package/src/git/hooks-path.ts +85 -85
  59. package/src/hooks/pattern-detector.ts +173 -173
  60. package/src/hooks/significance-filter.ts +250 -250
  61. package/src/llm/memory-manager.ts +328 -328
  62. package/src/llm/provider.ts +885 -885
  63. package/src/llm/quality.ts +248 -248
  64. package/src/memory/attribution-guard.ts +249 -249
  65. package/src/memory/disclosure-policy.ts +135 -135
  66. package/src/memory/entity-extractor.ts +197 -197
  67. package/src/memory/formation/evaluate.ts +217 -217
  68. package/src/memory/formation/extract.ts +361 -361
  69. package/src/memory/formation/index.ts +417 -417
  70. package/src/memory/formation/resolve.ts +344 -344
  71. package/src/memory/formation/types.ts +315 -315
  72. package/src/memory/freshness.ts +122 -122
  73. package/src/memory/graph.ts +197 -197
  74. package/src/memory/refs.ts +94 -94
  75. package/src/memory/secret-filter.ts +79 -79
  76. package/src/memory/session.ts +24 -9
  77. package/src/multimodal/image-loader.ts +143 -143
  78. package/src/orchestrate/adapters/claude-stream.ts +192 -192
  79. package/src/orchestrate/adapters/claude.ts +111 -111
  80. package/src/orchestrate/adapters/codex-stream.ts +134 -134
  81. package/src/orchestrate/adapters/codex.ts +41 -41
  82. package/src/orchestrate/adapters/gemini-stream.ts +166 -166
  83. package/src/orchestrate/adapters/gemini.ts +42 -42
  84. package/src/orchestrate/adapters/index.ts +73 -73
  85. package/src/orchestrate/adapters/opencode-stream.ts +143 -143
  86. package/src/orchestrate/adapters/opencode.ts +47 -47
  87. package/src/orchestrate/adapters/spawn-helper.ts +286 -286
  88. package/src/orchestrate/adapters/types.ts +77 -77
  89. package/src/orchestrate/capability-router.ts +284 -284
  90. package/src/orchestrate/context-compact.ts +188 -188
  91. package/src/orchestrate/cost-tracker.ts +219 -219
  92. package/src/orchestrate/error-recovery.ts +191 -191
  93. package/src/orchestrate/evidence.ts +140 -140
  94. package/src/orchestrate/ledger.ts +110 -110
  95. package/src/orchestrate/memorix-bridge.ts +343 -343
  96. package/src/orchestrate/output-budget.ts +80 -80
  97. package/src/orchestrate/permission.ts +152 -152
  98. package/src/orchestrate/pipeline-trace.ts +131 -131
  99. package/src/orchestrate/prompt-builder.ts +155 -155
  100. package/src/orchestrate/ring-buffer.ts +37 -37
  101. package/src/orchestrate/task-graph.ts +389 -389
  102. package/src/orchestrate/worktree.ts +232 -232
  103. package/src/project/aliases.ts +374 -374
  104. package/src/project/detector.ts +268 -268
  105. package/src/rules/adapters/claude-code.ts +99 -99
  106. package/src/rules/adapters/codex.ts +97 -97
  107. package/src/rules/adapters/copilot.ts +124 -124
  108. package/src/rules/adapters/cursor.ts +114 -114
  109. package/src/rules/adapters/kiro.ts +126 -126
  110. package/src/rules/adapters/trae.ts +56 -56
  111. package/src/rules/adapters/windsurf.ts +83 -83
  112. package/src/rules/syncer.ts +235 -235
  113. package/src/sdk.ts +299 -299
  114. package/src/search/intent-detector.ts +289 -289
  115. package/src/search/query-expansion.ts +52 -52
  116. package/src/server/formation-timeout.ts +27 -27
  117. package/src/server.ts +7 -2
  118. package/src/skills/mini-skills.ts +386 -386
  119. package/src/store/chat-store.ts +119 -119
  120. package/src/store/graph-store.ts +249 -249
  121. package/src/store/mini-skill-store.ts +349 -349
  122. package/src/store/persistence-json.ts +212 -212
  123. package/src/store/persistence.ts +291 -291
  124. package/src/store/project-affinity.ts +195 -195
  125. package/src/team/event-bus.ts +76 -76
  126. package/src/team/file-locks.ts +173 -173
  127. package/src/team/handoff.ts +161 -161
  128. package/src/team/messages.ts +203 -203
  129. package/src/team/poll.ts +132 -132
  130. package/src/team/tasks.ts +211 -211
  131. package/src/workspace/mcp-adapters/codex.ts +191 -191
  132. package/src/workspace/mcp-adapters/copilot.ts +105 -105
  133. package/src/workspace/mcp-adapters/cursor.ts +53 -53
  134. package/src/workspace/mcp-adapters/kiro.ts +64 -64
  135. package/src/workspace/mcp-adapters/opencode.ts +123 -123
  136. package/src/workspace/mcp-adapters/trae.ts +134 -134
  137. package/src/workspace/mcp-adapters/windsurf.ts +91 -91
  138. package/src/workspace/sanitizer.ts +60 -60
  139. package/src/workspace/workflow-sync.ts +131 -131
@@ -1,191 +1,191 @@
1
- /**
2
- * Error Recovery — Phase 7, Step 3: Layered error classification and recovery.
3
- *
4
- * Detects three failure modes from agent output/exit and recommends a
5
- * recovery strategy. The coordinator uses this to decide between:
6
- * - continuation (truncated output)
7
- * - context compaction + re-dispatch (context overflow)
8
- * - exponential backoff + retry (transient errors)
9
- * - normal retry (unknown / unrecoverable)
10
- *
11
- * Design principle: detection is heuristic-based pattern matching on
12
- * stderr/tailOutput. No false positive should crash the pipeline —
13
- * the worst case of a wrong classification is a wasted retry attempt.
14
- */
15
-
16
- // ── Types ──────────────────────────────────────────────────────────
17
-
18
- export type ErrorCategory = 'truncated' | 'context_overflow' | 'transient' | 'unknown';
19
-
20
- export interface RecoveryAction {
21
- category: ErrorCategory;
22
- /** Recommended recovery strategy */
23
- strategy: 'continue' | 'compact_and_retry' | 'backoff_and_retry' | 'normal_retry';
24
- /** Suggested delay before retry in ms (0 = immediate) */
25
- delayMs: number;
26
- /** Human-readable explanation of the detection */
27
- reason: string;
28
- /** Suggested continuation prompt (for 'continue' strategy) */
29
- continuationPrompt?: string;
30
- /** Backoff attempt number (for 'backoff_and_retry') */
31
- backoffAttempt?: number;
32
- }
33
-
34
- export interface ErrorContext {
35
- exitCode: number | null;
36
- killed: boolean;
37
- tailOutput: string;
38
- /** Whether the agent stream had an end_turn signal */
39
- hasEndTurn?: boolean;
40
- }
41
-
42
- // ── Pattern Registry ───────────────────────────────────────────────
43
-
44
- const TRUNCATED_PATTERNS = [
45
- /max.?output/i,
46
- /output.?limit/i,
47
- /max.?tokens/i,
48
- /response.?truncat/i,
49
- /generation.?limit/i,
50
- ];
51
-
52
- const CONTEXT_OVERFLOW_PATTERNS = [
53
- /overlong.?prompt/i,
54
- /context.?length.?exceed/i,
55
- /prompt.?too.?long/i,
56
- /maximum.?context/i,
57
- /token.?limit.?exceed/i,
58
- /input.?too.?long/i,
59
- /request.?too.?large/i,
60
- /max.?input.?tokens/i,
61
- /content.?too.?long/i,
62
- /exceeds.?(?:the\s+)?(?:model'?s?\s+)?(?:maximum|max).?(?:context|token)/i,
63
- ];
64
-
65
- const TRANSIENT_PATTERNS = [
66
- /rate.?limit/i,
67
- /429/,
68
- /too.?many.?requests/i,
69
- /server.?error/i,
70
- /502|503|504/,
71
- /bad.?gateway/i,
72
- /service.?unavailable/i,
73
- /gateway.?timeout/i,
74
- /ECONNRESET/i,
75
- /ECONNREFUSED/i,
76
- /ETIMEDOUT/i,
77
- /socket.?hang.?up/i,
78
- /network.?error/i,
79
- /temporary.?failure/i,
80
- /overloaded/i,
81
- /capacity/i,
82
- ];
83
-
84
- // ── Backoff State ──────────────────────────────────────────────────
85
-
86
- const backoffState = new Map<string, number>(); // taskId → consecutive transient failures
87
-
88
- /** Reset backoff state for a task (call on success or non-transient failure) */
89
- export function resetBackoff(taskId: string): void {
90
- backoffState.delete(taskId);
91
- }
92
-
93
- /** Get current backoff attempt count for a task */
94
- export function getBackoffAttempt(taskId: string): number {
95
- return backoffState.get(taskId) ?? 0;
96
- }
97
-
98
- // ── Core Classification ────────────────────────────────────────────
99
-
100
- /**
101
- * Classify an agent failure and recommend a recovery action.
102
- *
103
- * Priority order: truncated > context_overflow > transient > unknown.
104
- * This ensures the most specific recovery strategy is chosen.
105
- */
106
- export function classifyError(ctx: ErrorContext, taskId?: string): RecoveryAction {
107
- const { exitCode, killed, tailOutput, hasEndTurn } = ctx;
108
- const text = tailOutput.slice(-2000); // Only check last 2KB
109
-
110
- // 1. Truncated output detection
111
- // Agent was cut off mid-response (no end_turn + non-zero exit)
112
- if (!killed && exitCode !== 0 && hasEndTurn === false) {
113
- return {
114
- category: 'truncated',
115
- strategy: 'continue',
116
- delayMs: 0,
117
- reason: 'Agent output appears truncated (no end_turn signal with non-zero exit)',
118
- continuationPrompt: 'Your previous output was truncated due to output limits. Continue from where you stopped. Do not repeat what you already wrote.',
119
- };
120
- }
121
-
122
- // Also detect explicit truncation error messages
123
- if (matchesAny(text, TRUNCATED_PATTERNS)) {
124
- return {
125
- category: 'truncated',
126
- strategy: 'continue',
127
- delayMs: 0,
128
- reason: `Truncation pattern detected in output`,
129
- continuationPrompt: 'Your previous output was truncated due to output limits. Continue from where you stopped. Do not repeat what you already wrote.',
130
- };
131
- }
132
-
133
- // 2. Context overflow detection
134
- if (matchesAny(text, CONTEXT_OVERFLOW_PATTERNS)) {
135
- return {
136
- category: 'context_overflow',
137
- strategy: 'compact_and_retry',
138
- delayMs: 0,
139
- reason: 'Context overflow detected — prompt is too large for the model',
140
- };
141
- }
142
-
143
- // 3. Transient error detection
144
- if (matchesAny(text, TRANSIENT_PATTERNS)) {
145
- const attempt = taskId ? (backoffState.get(taskId) ?? 0) + 1 : 1;
146
- if (taskId) backoffState.set(taskId, attempt);
147
-
148
- // Exponential backoff: 2s, 4s, 8s, capped at 30s
149
- const baseDelay = 2000;
150
- const maxDelay = 30_000;
151
- const delay = Math.min(baseDelay * Math.pow(2, attempt - 1), maxDelay);
152
- // Add jitter: ±25%
153
- const jitter = delay * 0.25 * (Math.random() * 2 - 1);
154
-
155
- return {
156
- category: 'transient',
157
- strategy: 'backoff_and_retry',
158
- delayMs: Math.round(delay + jitter),
159
- reason: `Transient error detected (attempt ${attempt})`,
160
- backoffAttempt: attempt,
161
- };
162
- }
163
-
164
- // 4. Unknown — fall back to normal retry
165
- if (taskId) resetBackoff(taskId); // Non-transient failure resets backoff
166
-
167
- return {
168
- category: 'unknown',
169
- strategy: 'normal_retry',
170
- delayMs: 0,
171
- reason: killed
172
- ? 'Agent killed by timeout'
173
- : `Agent exited with code ${exitCode}`,
174
- };
175
- }
176
-
177
- /**
178
- * Check if the error is recoverable (not a permanent failure).
179
- * Truncated and transient errors are always recoverable.
180
- * Context overflow is recoverable if compaction is available.
181
- * Unknown errors may be recoverable via normal retry.
182
- */
183
- export function isRecoverable(action: RecoveryAction): boolean {
184
- return action.strategy !== 'normal_retry' || action.category === 'unknown';
185
- }
186
-
187
- // ── Helpers ────────────────────────────────────────────────────────
188
-
189
- function matchesAny(text: string, patterns: RegExp[]): boolean {
190
- return patterns.some(p => p.test(text));
191
- }
1
+ /**
2
+ * Error Recovery — Phase 7, Step 3: Layered error classification and recovery.
3
+ *
4
+ * Detects three failure modes from agent output/exit and recommends a
5
+ * recovery strategy. The coordinator uses this to decide between:
6
+ * - continuation (truncated output)
7
+ * - context compaction + re-dispatch (context overflow)
8
+ * - exponential backoff + retry (transient errors)
9
+ * - normal retry (unknown / unrecoverable)
10
+ *
11
+ * Design principle: detection is heuristic-based pattern matching on
12
+ * stderr/tailOutput. No false positive should crash the pipeline —
13
+ * the worst case of a wrong classification is a wasted retry attempt.
14
+ */
15
+
16
+ // ── Types ──────────────────────────────────────────────────────────
17
+
18
+ export type ErrorCategory = 'truncated' | 'context_overflow' | 'transient' | 'unknown';
19
+
20
+ export interface RecoveryAction {
21
+ category: ErrorCategory;
22
+ /** Recommended recovery strategy */
23
+ strategy: 'continue' | 'compact_and_retry' | 'backoff_and_retry' | 'normal_retry';
24
+ /** Suggested delay before retry in ms (0 = immediate) */
25
+ delayMs: number;
26
+ /** Human-readable explanation of the detection */
27
+ reason: string;
28
+ /** Suggested continuation prompt (for 'continue' strategy) */
29
+ continuationPrompt?: string;
30
+ /** Backoff attempt number (for 'backoff_and_retry') */
31
+ backoffAttempt?: number;
32
+ }
33
+
34
+ export interface ErrorContext {
35
+ exitCode: number | null;
36
+ killed: boolean;
37
+ tailOutput: string;
38
+ /** Whether the agent stream had an end_turn signal */
39
+ hasEndTurn?: boolean;
40
+ }
41
+
42
+ // ── Pattern Registry ───────────────────────────────────────────────
43
+
44
+ const TRUNCATED_PATTERNS = [
45
+ /max.?output/i,
46
+ /output.?limit/i,
47
+ /max.?tokens/i,
48
+ /response.?truncat/i,
49
+ /generation.?limit/i,
50
+ ];
51
+
52
+ const CONTEXT_OVERFLOW_PATTERNS = [
53
+ /overlong.?prompt/i,
54
+ /context.?length.?exceed/i,
55
+ /prompt.?too.?long/i,
56
+ /maximum.?context/i,
57
+ /token.?limit.?exceed/i,
58
+ /input.?too.?long/i,
59
+ /request.?too.?large/i,
60
+ /max.?input.?tokens/i,
61
+ /content.?too.?long/i,
62
+ /exceeds.?(?:the\s+)?(?:model'?s?\s+)?(?:maximum|max).?(?:context|token)/i,
63
+ ];
64
+
65
+ const TRANSIENT_PATTERNS = [
66
+ /rate.?limit/i,
67
+ /429/,
68
+ /too.?many.?requests/i,
69
+ /server.?error/i,
70
+ /502|503|504/,
71
+ /bad.?gateway/i,
72
+ /service.?unavailable/i,
73
+ /gateway.?timeout/i,
74
+ /ECONNRESET/i,
75
+ /ECONNREFUSED/i,
76
+ /ETIMEDOUT/i,
77
+ /socket.?hang.?up/i,
78
+ /network.?error/i,
79
+ /temporary.?failure/i,
80
+ /overloaded/i,
81
+ /capacity/i,
82
+ ];
83
+
84
+ // ── Backoff State ──────────────────────────────────────────────────
85
+
86
+ const backoffState = new Map<string, number>(); // taskId → consecutive transient failures
87
+
88
+ /** Reset backoff state for a task (call on success or non-transient failure) */
89
+ export function resetBackoff(taskId: string): void {
90
+ backoffState.delete(taskId);
91
+ }
92
+
93
+ /** Get current backoff attempt count for a task */
94
+ export function getBackoffAttempt(taskId: string): number {
95
+ return backoffState.get(taskId) ?? 0;
96
+ }
97
+
98
+ // ── Core Classification ────────────────────────────────────────────
99
+
100
+ /**
101
+ * Classify an agent failure and recommend a recovery action.
102
+ *
103
+ * Priority order: truncated > context_overflow > transient > unknown.
104
+ * This ensures the most specific recovery strategy is chosen.
105
+ */
106
+ export function classifyError(ctx: ErrorContext, taskId?: string): RecoveryAction {
107
+ const { exitCode, killed, tailOutput, hasEndTurn } = ctx;
108
+ const text = tailOutput.slice(-2000); // Only check last 2KB
109
+
110
+ // 1. Truncated output detection
111
+ // Agent was cut off mid-response (no end_turn + non-zero exit)
112
+ if (!killed && exitCode !== 0 && hasEndTurn === false) {
113
+ return {
114
+ category: 'truncated',
115
+ strategy: 'continue',
116
+ delayMs: 0,
117
+ reason: 'Agent output appears truncated (no end_turn signal with non-zero exit)',
118
+ continuationPrompt: 'Your previous output was truncated due to output limits. Continue from where you stopped. Do not repeat what you already wrote.',
119
+ };
120
+ }
121
+
122
+ // Also detect explicit truncation error messages
123
+ if (matchesAny(text, TRUNCATED_PATTERNS)) {
124
+ return {
125
+ category: 'truncated',
126
+ strategy: 'continue',
127
+ delayMs: 0,
128
+ reason: `Truncation pattern detected in output`,
129
+ continuationPrompt: 'Your previous output was truncated due to output limits. Continue from where you stopped. Do not repeat what you already wrote.',
130
+ };
131
+ }
132
+
133
+ // 2. Context overflow detection
134
+ if (matchesAny(text, CONTEXT_OVERFLOW_PATTERNS)) {
135
+ return {
136
+ category: 'context_overflow',
137
+ strategy: 'compact_and_retry',
138
+ delayMs: 0,
139
+ reason: 'Context overflow detected — prompt is too large for the model',
140
+ };
141
+ }
142
+
143
+ // 3. Transient error detection
144
+ if (matchesAny(text, TRANSIENT_PATTERNS)) {
145
+ const attempt = taskId ? (backoffState.get(taskId) ?? 0) + 1 : 1;
146
+ if (taskId) backoffState.set(taskId, attempt);
147
+
148
+ // Exponential backoff: 2s, 4s, 8s, capped at 30s
149
+ const baseDelay = 2000;
150
+ const maxDelay = 30_000;
151
+ const delay = Math.min(baseDelay * Math.pow(2, attempt - 1), maxDelay);
152
+ // Add jitter: ±25%
153
+ const jitter = delay * 0.25 * (Math.random() * 2 - 1);
154
+
155
+ return {
156
+ category: 'transient',
157
+ strategy: 'backoff_and_retry',
158
+ delayMs: Math.round(delay + jitter),
159
+ reason: `Transient error detected (attempt ${attempt})`,
160
+ backoffAttempt: attempt,
161
+ };
162
+ }
163
+
164
+ // 4. Unknown — fall back to normal retry
165
+ if (taskId) resetBackoff(taskId); // Non-transient failure resets backoff
166
+
167
+ return {
168
+ category: 'unknown',
169
+ strategy: 'normal_retry',
170
+ delayMs: 0,
171
+ reason: killed
172
+ ? 'Agent killed by timeout'
173
+ : `Agent exited with code ${exitCode}`,
174
+ };
175
+ }
176
+
177
+ /**
178
+ * Check if the error is recoverable (not a permanent failure).
179
+ * Truncated and transient errors are always recoverable.
180
+ * Context overflow is recoverable if compaction is available.
181
+ * Unknown errors may be recoverable via normal retry.
182
+ */
183
+ export function isRecoverable(action: RecoveryAction): boolean {
184
+ return action.strategy !== 'normal_retry' || action.category === 'unknown';
185
+ }
186
+
187
+ // ── Helpers ────────────────────────────────────────────────────────
188
+
189
+ function matchesAny(text: string, patterns: RegExp[]): boolean {
190
+ return patterns.some(p => p.test(text));
191
+ }
@@ -1,140 +1,140 @@
1
- /**
2
- * Evidence Directory — Phase 7, Step 9: Per-task evidence collection.
3
- *
4
- * Creates a structured evidence directory for each pipeline run:
5
- * .pipeline/<pipelineId>/evidence/<taskId>/
6
- * prompt.md, output.txt, compile.txt, test.txt, result.json
7
- *
8
- * Design principle: evidence collection is best-effort. Disk write failures
9
- * are logged and ignored — never crash the pipeline for evidence.
10
- */
11
-
12
- import { mkdirSync, writeFileSync } from 'node:fs';
13
- import { join } from 'node:path';
14
- import type { GateResult } from './verify-gate.js';
15
- import type { TokenUsage } from './adapters/types.js';
16
-
17
- // ── Types ──────────────────────────────────────────────────────────
18
-
19
- export interface TaskEvidence {
20
- taskId: string;
21
- taskDescription: string;
22
- agentName: string;
23
- status: 'completed' | 'failed';
24
- durationMs: number;
25
- prompt?: string;
26
- tailOutput?: string;
27
- gateResults?: GateResult[];
28
- tokenUsage?: Record<string, TokenUsage>;
29
- costUSD?: number | null;
30
- fixAttempts?: number;
31
- }
32
-
33
- export interface PipelineSummary {
34
- pipelineId: string;
35
- goal: string;
36
- totalTasks: number;
37
- completed: number;
38
- failed: number;
39
- elapsedMs: number;
40
- tokenUsage?: Record<string, TokenUsage>;
41
- costUSD?: number | null;
42
- tasks: TaskEvidence[];
43
- /** A4: Idle agents and why they weren't dispatched */
44
- idleAgents?: Array<{ name: string; reason: string }>;
45
- /** A2: Routing decisions for explainability */
46
- routingDecisions?: Array<{ role: string; selected: string; reason: string; available: string[] }>;
47
- }
48
-
49
- // ── Core ───────────────────────────────────────────────────────────
50
-
51
- /**
52
- * Write evidence for a single task. Best-effort — never throws.
53
- */
54
- export function writeTaskEvidence(
55
- projectDir: string,
56
- pipelineId: string,
57
- evidence: TaskEvidence,
58
- ): string | null {
59
- try {
60
- const dir = join(projectDir, '.pipeline', pipelineId, 'evidence', evidence.taskId.slice(0, 12));
61
- mkdirSync(dir, { recursive: true });
62
-
63
- if (evidence.prompt) {
64
- safeWrite(join(dir, 'prompt.md'), evidence.prompt);
65
- }
66
- if (evidence.tailOutput) {
67
- safeWrite(join(dir, 'output.txt'), evidence.tailOutput);
68
- }
69
- if (evidence.gateResults) {
70
- for (const g of evidence.gateResults) {
71
- const content = `Command: ${g.command}\nStatus: ${g.passed ? 'PASS' : 'FAIL'}\nDuration: ${(g.durationMs / 1000).toFixed(1)}s\n\n${g.output}`;
72
- safeWrite(join(dir, `${g.gate}.txt`), content);
73
- }
74
- }
75
-
76
- const resultJson = {
77
- taskId: evidence.taskId,
78
- description: evidence.taskDescription,
79
- agent: evidence.agentName,
80
- status: evidence.status,
81
- durationMs: evidence.durationMs,
82
- fixAttempts: evidence.fixAttempts ?? 0,
83
- gates: evidence.gateResults?.map(g => ({ gate: g.gate, passed: g.passed, durationMs: g.durationMs, command: g.command })),
84
- tokenUsage: evidence.tokenUsage,
85
- costUSD: evidence.costUSD,
86
- };
87
- safeWrite(join(dir, 'result.json'), JSON.stringify(resultJson, null, 2));
88
- return dir;
89
- } catch {
90
- return null;
91
- }
92
- }
93
-
94
- /**
95
- * Write pipeline summary markdown. Best-effort — never throws.
96
- */
97
- export function writePipelineSummary(
98
- projectDir: string,
99
- summary: PipelineSummary,
100
- ): string | null {
101
- try {
102
- const dir = join(projectDir, '.pipeline', summary.pipelineId);
103
- mkdirSync(dir, { recursive: true });
104
-
105
- const elapsed = (summary.elapsedMs / 1000).toFixed(0);
106
- const lines = [
107
- `# Pipeline: ${summary.goal.slice(0, 100)}`,
108
- '',
109
- `- **ID**: ${summary.pipelineId}`,
110
- `- **Tasks**: ${summary.completed}/${summary.totalTasks} completed, ${summary.failed} failed`,
111
- `- **Elapsed**: ${elapsed}s`,
112
- ];
113
-
114
- if (summary.costUSD != null) {
115
- lines.push(`- **Cost**: $${summary.costUSD.toFixed(4)}`);
116
- }
117
-
118
- lines.push('', '## Tasks', '');
119
- for (const t of summary.tasks) {
120
- const status = t.status === 'completed' ? 'PASS' : 'FAIL';
121
- const dur = (t.durationMs / 1000).toFixed(1);
122
- const fixes = t.fixAttempts ? ` (${t.fixAttempts} fix attempts)` : '';
123
- lines.push(`- [${status}] ${t.taskDescription.slice(0, 80)} — ${t.agentName}, ${dur}s${fixes}`);
124
- }
125
-
126
- const content = lines.join('\n') + '\n';
127
- safeWrite(join(dir, 'summary.md'), content);
128
- return join(dir, 'summary.md');
129
- } catch {
130
- return null;
131
- }
132
- }
133
-
134
- // ── Helpers ────────────────────────────────────────────────────────
135
-
136
- function safeWrite(path: string, content: string): void {
137
- try {
138
- writeFileSync(path, content, 'utf-8');
139
- } catch { /* best-effort */ }
140
- }
1
+ /**
2
+ * Evidence Directory — Phase 7, Step 9: Per-task evidence collection.
3
+ *
4
+ * Creates a structured evidence directory for each pipeline run:
5
+ * .pipeline/<pipelineId>/evidence/<taskId>/
6
+ * prompt.md, output.txt, compile.txt, test.txt, result.json
7
+ *
8
+ * Design principle: evidence collection is best-effort. Disk write failures
9
+ * are logged and ignored — never crash the pipeline for evidence.
10
+ */
11
+
12
+ import { mkdirSync, writeFileSync } from 'node:fs';
13
+ import { join } from 'node:path';
14
+ import type { GateResult } from './verify-gate.js';
15
+ import type { TokenUsage } from './adapters/types.js';
16
+
17
+ // ── Types ──────────────────────────────────────────────────────────
18
+
19
+ export interface TaskEvidence {
20
+ taskId: string;
21
+ taskDescription: string;
22
+ agentName: string;
23
+ status: 'completed' | 'failed';
24
+ durationMs: number;
25
+ prompt?: string;
26
+ tailOutput?: string;
27
+ gateResults?: GateResult[];
28
+ tokenUsage?: Record<string, TokenUsage>;
29
+ costUSD?: number | null;
30
+ fixAttempts?: number;
31
+ }
32
+
33
+ export interface PipelineSummary {
34
+ pipelineId: string;
35
+ goal: string;
36
+ totalTasks: number;
37
+ completed: number;
38
+ failed: number;
39
+ elapsedMs: number;
40
+ tokenUsage?: Record<string, TokenUsage>;
41
+ costUSD?: number | null;
42
+ tasks: TaskEvidence[];
43
+ /** A4: Idle agents and why they weren't dispatched */
44
+ idleAgents?: Array<{ name: string; reason: string }>;
45
+ /** A2: Routing decisions for explainability */
46
+ routingDecisions?: Array<{ role: string; selected: string; reason: string; available: string[] }>;
47
+ }
48
+
49
+ // ── Core ───────────────────────────────────────────────────────────
50
+
51
+ /**
52
+ * Write evidence for a single task. Best-effort — never throws.
53
+ */
54
+ export function writeTaskEvidence(
55
+ projectDir: string,
56
+ pipelineId: string,
57
+ evidence: TaskEvidence,
58
+ ): string | null {
59
+ try {
60
+ const dir = join(projectDir, '.pipeline', pipelineId, 'evidence', evidence.taskId.slice(0, 12));
61
+ mkdirSync(dir, { recursive: true });
62
+
63
+ if (evidence.prompt) {
64
+ safeWrite(join(dir, 'prompt.md'), evidence.prompt);
65
+ }
66
+ if (evidence.tailOutput) {
67
+ safeWrite(join(dir, 'output.txt'), evidence.tailOutput);
68
+ }
69
+ if (evidence.gateResults) {
70
+ for (const g of evidence.gateResults) {
71
+ const content = `Command: ${g.command}\nStatus: ${g.passed ? 'PASS' : 'FAIL'}\nDuration: ${(g.durationMs / 1000).toFixed(1)}s\n\n${g.output}`;
72
+ safeWrite(join(dir, `${g.gate}.txt`), content);
73
+ }
74
+ }
75
+
76
+ const resultJson = {
77
+ taskId: evidence.taskId,
78
+ description: evidence.taskDescription,
79
+ agent: evidence.agentName,
80
+ status: evidence.status,
81
+ durationMs: evidence.durationMs,
82
+ fixAttempts: evidence.fixAttempts ?? 0,
83
+ gates: evidence.gateResults?.map(g => ({ gate: g.gate, passed: g.passed, durationMs: g.durationMs, command: g.command })),
84
+ tokenUsage: evidence.tokenUsage,
85
+ costUSD: evidence.costUSD,
86
+ };
87
+ safeWrite(join(dir, 'result.json'), JSON.stringify(resultJson, null, 2));
88
+ return dir;
89
+ } catch {
90
+ return null;
91
+ }
92
+ }
93
+
94
+ /**
95
+ * Write pipeline summary markdown. Best-effort — never throws.
96
+ */
97
+ export function writePipelineSummary(
98
+ projectDir: string,
99
+ summary: PipelineSummary,
100
+ ): string | null {
101
+ try {
102
+ const dir = join(projectDir, '.pipeline', summary.pipelineId);
103
+ mkdirSync(dir, { recursive: true });
104
+
105
+ const elapsed = (summary.elapsedMs / 1000).toFixed(0);
106
+ const lines = [
107
+ `# Pipeline: ${summary.goal.slice(0, 100)}`,
108
+ '',
109
+ `- **ID**: ${summary.pipelineId}`,
110
+ `- **Tasks**: ${summary.completed}/${summary.totalTasks} completed, ${summary.failed} failed`,
111
+ `- **Elapsed**: ${elapsed}s`,
112
+ ];
113
+
114
+ if (summary.costUSD != null) {
115
+ lines.push(`- **Cost**: $${summary.costUSD.toFixed(4)}`);
116
+ }
117
+
118
+ lines.push('', '## Tasks', '');
119
+ for (const t of summary.tasks) {
120
+ const status = t.status === 'completed' ? 'PASS' : 'FAIL';
121
+ const dur = (t.durationMs / 1000).toFixed(1);
122
+ const fixes = t.fixAttempts ? ` (${t.fixAttempts} fix attempts)` : '';
123
+ lines.push(`- [${status}] ${t.taskDescription.slice(0, 80)} — ${t.agentName}, ${dur}s${fixes}`);
124
+ }
125
+
126
+ const content = lines.join('\n') + '\n';
127
+ safeWrite(join(dir, 'summary.md'), content);
128
+ return join(dir, 'summary.md');
129
+ } catch {
130
+ return null;
131
+ }
132
+ }
133
+
134
+ // ── Helpers ────────────────────────────────────────────────────────
135
+
136
+ function safeWrite(path: string, content: string): void {
137
+ try {
138
+ writeFileSync(path, content, 'utf-8');
139
+ } catch { /* best-effort */ }
140
+ }