memorix 1.2.1 → 1.2.2

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (199) hide show
  1. package/CHANGELOG.md +16 -0
  2. package/README.md +14 -2
  3. package/README.zh-CN.md +14 -2
  4. package/TEAM.md +86 -86
  5. package/dist/cli/index.js +15407 -13779
  6. package/dist/cli/index.js.map +1 -1
  7. package/dist/index.js +1321 -529
  8. package/dist/index.js.map +1 -1
  9. package/dist/maintenance-runner.d.ts +1 -1
  10. package/dist/maintenance-runner.js +8458 -8087
  11. package/dist/maintenance-runner.js.map +1 -1
  12. package/dist/memcode-runtime/CHANGELOG.md +16 -0
  13. package/dist/sdk.d.ts +7 -2
  14. package/dist/sdk.js +1349 -535
  15. package/dist/sdk.js.map +1 -1
  16. package/dist/types.d.ts +49 -1
  17. package/dist/types.js.map +1 -1
  18. package/docs/1.2.2-MEMORY-CONTROL-PLANE.md +434 -0
  19. package/docs/AGENT_OPERATOR_PLAYBOOK.md +4 -0
  20. package/docs/API_REFERENCE.md +24 -4
  21. package/docs/DESIGN_DECISIONS.md +357 -357
  22. package/docs/README.md +1 -1
  23. package/docs/dev-log/progress.txt +91 -11
  24. package/package.json +1 -1
  25. package/plugins/codex/memorix/.codex-plugin/plugin.json +1 -1
  26. package/src/audit/index.ts +156 -156
  27. package/src/cli/command-guide.ts +192 -0
  28. package/src/cli/commands/audit-list.ts +89 -89
  29. package/src/cli/commands/audit.ts +9 -4
  30. package/src/cli/commands/background.ts +659 -659
  31. package/src/cli/commands/cleanup.ts +5 -1
  32. package/src/cli/commands/codegraph.ts +15 -5
  33. package/src/cli/commands/context.ts +3 -2
  34. package/src/cli/commands/doctor.ts +4 -2
  35. package/src/cli/commands/explain.ts +9 -3
  36. package/src/cli/commands/formation.ts +48 -48
  37. package/src/cli/commands/git-hook-install.ts +111 -111
  38. package/src/cli/commands/handoff.ts +75 -61
  39. package/src/cli/commands/hooks-status.ts +63 -63
  40. package/src/cli/commands/identity.ts +116 -0
  41. package/src/cli/commands/ingest-commit.ts +153 -153
  42. package/src/cli/commands/ingest-image.ts +71 -69
  43. package/src/cli/commands/ingest-log.ts +180 -180
  44. package/src/cli/commands/ingest.ts +44 -44
  45. package/src/cli/commands/integrate-shared.ts +15 -15
  46. package/src/cli/commands/lock.ts +93 -92
  47. package/src/cli/commands/memory.ts +58 -21
  48. package/src/cli/commands/message.ts +123 -118
  49. package/src/cli/commands/operator-shared.ts +98 -3
  50. package/src/cli/commands/poll.ts +74 -64
  51. package/src/cli/commands/purge-all-memory.ts +85 -85
  52. package/src/cli/commands/purge-project-memory.ts +83 -83
  53. package/src/cli/commands/reasoning.ts +135 -121
  54. package/src/cli/commands/retention.ts +9 -4
  55. package/src/cli/commands/serve-http.ts +8 -2
  56. package/src/cli/commands/serve-shared.ts +118 -118
  57. package/src/cli/commands/session.ts +29 -3
  58. package/src/cli/commands/skills.ts +124 -119
  59. package/src/cli/commands/status.ts +4 -3
  60. package/src/cli/commands/task.ts +193 -184
  61. package/src/cli/commands/team.ts +14 -10
  62. package/src/cli/commands/transfer.ts +108 -55
  63. package/src/cli/commands/uninstall-project-artifacts.ts +85 -85
  64. package/src/cli/identity.ts +89 -0
  65. package/src/cli/index.ts +96 -19
  66. package/src/cli/invocation.ts +115 -0
  67. package/src/cli/tui/ChatView.tsx +234 -234
  68. package/src/cli/tui/CommandBar.tsx +312 -312
  69. package/src/cli/tui/ContextRail.tsx +118 -118
  70. package/src/cli/tui/HeaderBar.tsx +72 -72
  71. package/src/cli/tui/LogoBanner.tsx +51 -51
  72. package/src/cli/tui/Sidebar.tsx +179 -179
  73. package/src/cli/tui/chat-service.ts +41 -18
  74. package/src/cli/tui/data.ts +23 -44
  75. package/src/cli/tui/index.ts +41 -41
  76. package/src/cli/tui/markdown-render.tsx +371 -371
  77. package/src/cli/tui/operator-context.ts +60 -0
  78. package/src/cli/tui/use-mouse.ts +157 -157
  79. package/src/cli/tui/useNavigation.ts +56 -56
  80. package/src/cli/tui/views/MemoryView.tsx +10 -8
  81. package/src/cli/update-checker.ts +211 -211
  82. package/src/cli/version.ts +7 -7
  83. package/src/cli/workbench.ts +1 -1
  84. package/src/codegraph/auto-context.ts +31 -2
  85. package/src/codegraph/context-pack.ts +1 -0
  86. package/src/codegraph/project-context.ts +2 -0
  87. package/src/compact/engine.ts +26 -10
  88. package/src/compact/index-format.ts +25 -2
  89. package/src/compact/token-budget.ts +74 -74
  90. package/src/dashboard/project-classification.ts +64 -64
  91. package/src/dashboard/server.ts +46 -9
  92. package/src/embedding/fastembed-provider.ts +142 -142
  93. package/src/embedding/transformers-provider.ts +111 -111
  94. package/src/git/extractor.ts +209 -209
  95. package/src/git/hooks-path.ts +85 -85
  96. package/src/hooks/admission.ts +117 -0
  97. package/src/hooks/handler.ts +98 -91
  98. package/src/hooks/pattern-detector.ts +173 -173
  99. package/src/hooks/significance-filter.ts +250 -250
  100. package/src/knowledge/context-assembly.ts +97 -0
  101. package/src/knowledge/workset.ts +179 -10
  102. package/src/llm/memory-manager.ts +328 -328
  103. package/src/llm/provider.ts +885 -885
  104. package/src/llm/quality.ts +248 -248
  105. package/src/memory/admission.ts +57 -0
  106. package/src/memory/attribution-guard.ts +249 -249
  107. package/src/memory/consolidation.ts +13 -2
  108. package/src/memory/disclosure-policy.ts +140 -135
  109. package/src/memory/entity-extractor.ts +197 -197
  110. package/src/memory/export-import.ts +11 -3
  111. package/src/memory/formation/evaluate.ts +217 -217
  112. package/src/memory/formation/extract.ts +361 -361
  113. package/src/memory/formation/index.ts +417 -417
  114. package/src/memory/formation/resolve.ts +344 -344
  115. package/src/memory/formation/types.ts +315 -315
  116. package/src/memory/freshness.ts +122 -122
  117. package/src/memory/graph-context.ts +8 -2
  118. package/src/memory/graph.ts +197 -197
  119. package/src/memory/observations.ts +162 -4
  120. package/src/memory/quality-audit.ts +2 -0
  121. package/src/memory/refs.ts +94 -94
  122. package/src/memory/retention.ts +22 -2
  123. package/src/memory/secret-filter.ts +79 -79
  124. package/src/memory/session.ts +5 -2
  125. package/src/memory/visibility.ts +80 -0
  126. package/src/multimodal/image-loader.ts +143 -143
  127. package/src/orchestrate/adapters/claude-stream.ts +192 -192
  128. package/src/orchestrate/adapters/claude.ts +111 -111
  129. package/src/orchestrate/adapters/codex-stream.ts +134 -134
  130. package/src/orchestrate/adapters/codex.ts +41 -41
  131. package/src/orchestrate/adapters/gemini-stream.ts +166 -166
  132. package/src/orchestrate/adapters/gemini.ts +42 -42
  133. package/src/orchestrate/adapters/index.ts +73 -73
  134. package/src/orchestrate/adapters/opencode-stream.ts +143 -143
  135. package/src/orchestrate/adapters/opencode.ts +47 -47
  136. package/src/orchestrate/adapters/spawn-helper.ts +286 -286
  137. package/src/orchestrate/adapters/types.ts +77 -77
  138. package/src/orchestrate/capability-router.ts +284 -284
  139. package/src/orchestrate/context-compact.ts +188 -188
  140. package/src/orchestrate/cost-tracker.ts +219 -219
  141. package/src/orchestrate/error-recovery.ts +191 -191
  142. package/src/orchestrate/evidence.ts +140 -140
  143. package/src/orchestrate/ledger.ts +110 -110
  144. package/src/orchestrate/memorix-bridge.ts +378 -340
  145. package/src/orchestrate/output-budget.ts +80 -80
  146. package/src/orchestrate/permission.ts +152 -152
  147. package/src/orchestrate/pipeline-trace.ts +131 -131
  148. package/src/orchestrate/prompt-builder.ts +155 -155
  149. package/src/orchestrate/ring-buffer.ts +37 -37
  150. package/src/orchestrate/task-graph.ts +389 -389
  151. package/src/orchestrate/worktree.ts +232 -232
  152. package/src/project/aliases.ts +374 -374
  153. package/src/project/detector.ts +268 -268
  154. package/src/rules/adapters/claude-code.ts +99 -99
  155. package/src/rules/adapters/codex.ts +97 -97
  156. package/src/rules/adapters/copilot.ts +124 -124
  157. package/src/rules/adapters/cursor.ts +114 -114
  158. package/src/rules/adapters/kiro.ts +126 -126
  159. package/src/rules/adapters/trae.ts +56 -56
  160. package/src/rules/adapters/windsurf.ts +83 -83
  161. package/src/rules/syncer.ts +235 -235
  162. package/src/runtime/control-plane-maintenance.ts +1 -0
  163. package/src/runtime/isolated-maintenance.ts +1 -0
  164. package/src/runtime/lifecycle.ts +18 -0
  165. package/src/runtime/maintenance-jobs.ts +1 -0
  166. package/src/runtime/maintenance-runner.ts +2 -0
  167. package/src/runtime/project-maintenance.ts +89 -0
  168. package/src/sdk.ts +334 -304
  169. package/src/search/intent-detector.ts +289 -289
  170. package/src/search/query-expansion.ts +52 -52
  171. package/src/server/formation-timeout.ts +27 -27
  172. package/src/server.ts +260 -81
  173. package/src/skills/mini-skills.ts +386 -386
  174. package/src/store/chat-store.ts +119 -119
  175. package/src/store/graph-store.ts +249 -249
  176. package/src/store/mini-skill-store.ts +349 -349
  177. package/src/store/orama-store.ts +61 -6
  178. package/src/store/persistence-json.ts +212 -212
  179. package/src/store/persistence.ts +291 -291
  180. package/src/store/project-affinity.ts +195 -195
  181. package/src/store/sqlite-db.ts +23 -1
  182. package/src/store/sqlite-store.ts +12 -2
  183. package/src/team/event-bus.ts +76 -76
  184. package/src/team/file-locks.ts +173 -173
  185. package/src/team/handoff.ts +168 -161
  186. package/src/team/messages.ts +203 -203
  187. package/src/team/poll.ts +132 -132
  188. package/src/team/tasks.ts +211 -211
  189. package/src/types.ts +51 -0
  190. package/src/wiki/generator.ts +2 -0
  191. package/src/workspace/mcp-adapters/codex.ts +191 -191
  192. package/src/workspace/mcp-adapters/copilot.ts +105 -105
  193. package/src/workspace/mcp-adapters/cursor.ts +53 -53
  194. package/src/workspace/mcp-adapters/kiro.ts +64 -64
  195. package/src/workspace/mcp-adapters/opencode.ts +123 -123
  196. package/src/workspace/mcp-adapters/trae.ts +134 -134
  197. package/src/workspace/mcp-adapters/windsurf.ts +91 -91
  198. package/src/workspace/sanitizer.ts +60 -60
  199. package/src/workspace/workflow-sync.ts +131 -131
@@ -1,80 +1,80 @@
1
- /**
2
- * Output Budget — Phase 7, Step 4: Trim large outputs to a token-friendly size.
3
- *
4
- * When gate outputs, agent tail outputs, or other large strings need to be
5
- * embedded in prompts (e.g., fix prompts, ledger entries), this module trims
6
- * them to a configurable budget while preserving the most useful parts
7
- * (beginning + end).
8
- *
9
- * Full output is optionally persisted to disk for debugging.
10
- */
11
-
12
- import { writeFileSync, mkdirSync } from 'node:fs';
13
- import { join, dirname } from 'node:path';
14
-
15
- // ── Constants ──────────────────────────────────────────────────────
16
-
17
- /** Default output budget in bytes */
18
- export const DEFAULT_OUTPUT_BUDGET = 4096;
19
-
20
- // ── Core ───────────────────────────────────────────────────────────
21
-
22
- /**
23
- * Trim text to a byte budget. If the text exceeds the budget, keep
24
- * the first half and last half with an omission marker in between.
25
- *
26
- * Returns the trimmed text. Never throws.
27
- */
28
- export function trimToBudget(text: string, budget: number = DEFAULT_OUTPUT_BUDGET): string {
29
- if (text.length <= budget) return text;
30
-
31
- const half = Math.floor(budget / 2) - 40; // Reserve space for marker
32
- if (half <= 0) return text.slice(0, budget);
33
-
34
- const omitted = text.length - half * 2;
35
- return (
36
- text.slice(0, half) +
37
- `\n\n... (${omitted} bytes omitted — full output saved to disk) ...\n\n` +
38
- text.slice(-half)
39
- );
40
- }
41
-
42
- /**
43
- * Persist full output to a file and return the trimmed version.
44
- *
45
- * If disk write fails, returns the trimmed text without file reference.
46
- * Never throws.
47
- */
48
- export function trimAndPersist(
49
- text: string,
50
- filePath: string,
51
- budget: number = DEFAULT_OUTPUT_BUDGET,
52
- ): { trimmed: string; persisted: boolean; fullPath?: string } {
53
- if (text.length <= budget) {
54
- return { trimmed: text, persisted: false };
55
- }
56
-
57
- // Persist full output to disk (best-effort)
58
- let persisted = false;
59
- try {
60
- mkdirSync(dirname(filePath), { recursive: true });
61
- writeFileSync(filePath, text, 'utf-8');
62
- persisted = true;
63
- } catch {
64
- // Disk write failed — degrade gracefully
65
- }
66
-
67
- const half = Math.floor(budget / 2) - 60; // Reserve space for marker + path
68
- if (half <= 0) {
69
- return { trimmed: text.slice(0, budget), persisted, fullPath: persisted ? filePath : undefined };
70
- }
71
-
72
- const omitted = text.length - half * 2;
73
- const pathNote = persisted ? ` Full output at: ${filePath}` : '';
74
- const trimmed =
75
- text.slice(0, half) +
76
- `\n\n... (${omitted} bytes omitted.${pathNote}) ...\n\n` +
77
- text.slice(-half);
78
-
79
- return { trimmed, persisted, fullPath: persisted ? filePath : undefined };
80
- }
1
+ /**
2
+ * Output Budget — Phase 7, Step 4: Trim large outputs to a token-friendly size.
3
+ *
4
+ * When gate outputs, agent tail outputs, or other large strings need to be
5
+ * embedded in prompts (e.g., fix prompts, ledger entries), this module trims
6
+ * them to a configurable budget while preserving the most useful parts
7
+ * (beginning + end).
8
+ *
9
+ * Full output is optionally persisted to disk for debugging.
10
+ */
11
+
12
+ import { writeFileSync, mkdirSync } from 'node:fs';
13
+ import { join, dirname } from 'node:path';
14
+
15
+ // ── Constants ──────────────────────────────────────────────────────
16
+
17
+ /** Default output budget in bytes */
18
+ export const DEFAULT_OUTPUT_BUDGET = 4096;
19
+
20
+ // ── Core ───────────────────────────────────────────────────────────
21
+
22
+ /**
23
+ * Trim text to a byte budget. If the text exceeds the budget, keep
24
+ * the first half and last half with an omission marker in between.
25
+ *
26
+ * Returns the trimmed text. Never throws.
27
+ */
28
+ export function trimToBudget(text: string, budget: number = DEFAULT_OUTPUT_BUDGET): string {
29
+ if (text.length <= budget) return text;
30
+
31
+ const half = Math.floor(budget / 2) - 40; // Reserve space for marker
32
+ if (half <= 0) return text.slice(0, budget);
33
+
34
+ const omitted = text.length - half * 2;
35
+ return (
36
+ text.slice(0, half) +
37
+ `\n\n... (${omitted} bytes omitted — full output saved to disk) ...\n\n` +
38
+ text.slice(-half)
39
+ );
40
+ }
41
+
42
+ /**
43
+ * Persist full output to a file and return the trimmed version.
44
+ *
45
+ * If disk write fails, returns the trimmed text without file reference.
46
+ * Never throws.
47
+ */
48
+ export function trimAndPersist(
49
+ text: string,
50
+ filePath: string,
51
+ budget: number = DEFAULT_OUTPUT_BUDGET,
52
+ ): { trimmed: string; persisted: boolean; fullPath?: string } {
53
+ if (text.length <= budget) {
54
+ return { trimmed: text, persisted: false };
55
+ }
56
+
57
+ // Persist full output to disk (best-effort)
58
+ let persisted = false;
59
+ try {
60
+ mkdirSync(dirname(filePath), { recursive: true });
61
+ writeFileSync(filePath, text, 'utf-8');
62
+ persisted = true;
63
+ } catch {
64
+ // Disk write failed — degrade gracefully
65
+ }
66
+
67
+ const half = Math.floor(budget / 2) - 60; // Reserve space for marker + path
68
+ if (half <= 0) {
69
+ return { trimmed: text.slice(0, budget), persisted, fullPath: persisted ? filePath : undefined };
70
+ }
71
+
72
+ const omitted = text.length - half * 2;
73
+ const pathNote = persisted ? ` Full output at: ${filePath}` : '';
74
+ const trimmed =
75
+ text.slice(0, half) +
76
+ `\n\n... (${omitted} bytes omitted.${pathNote}) ...\n\n` +
77
+ text.slice(-half);
78
+
79
+ return { trimmed, persisted, fullPath: persisted ? filePath : undefined };
80
+ }
@@ -1,152 +1,152 @@
1
- /**
2
- * Permission — Phase 7, Step 10: Risk classification + monitoring for agent tools.
3
- *
4
- * Classifies tool calls by risk tier and records them for post-run audit.
5
- * Does NOT block or restrict agent tools (Non-Invasiveness principle).
6
- * The primary use case is monitoring and reporting:
7
- * - Which high-risk tools were used and how often
8
- * - Whether agents used filesystem-destructive or network operations
9
- * - Per-task risk summary in the evidence directory
10
- *
11
- * Three tiers:
12
- * - safe: read-only operations (Read, Grep, List, etc.)
13
- * - moderate: write operations (Edit, Write, Execute safe commands)
14
- * - dangerous: destructive/network/exec operations (Delete, Bash, Deploy, etc.)
15
- */
16
-
17
- // ── Types ──────────────────────────────────────────────────────────
18
-
19
- export type RiskTier = 'safe' | 'moderate' | 'dangerous';
20
-
21
- export interface ToolUsageRecord {
22
- tool: string;
23
- tier: RiskTier;
24
- count: number;
25
- firstSeen: number;
26
- lastSeen: number;
27
- }
28
-
29
- export interface TaskRiskProfile {
30
- taskId: string;
31
- /** Highest risk tier used by this task */
32
- maxTier: RiskTier;
33
- /** Total tool calls */
34
- totalCalls: number;
35
- /** Tool usage breakdown */
36
- tools: ToolUsageRecord[];
37
- }
38
-
39
- // ── Classification Registry ────────────────────────────────────────
40
-
41
- const SAFE_TOOLS = new Set([
42
- 'read_file', 'grep_search', 'find_by_name', 'list_dir', 'code_search',
43
- 'read_notebook', 'read_url_content', 'view_content_chunk',
44
- 'memorix_search', 'memorix_detail', 'memorix_search_reasoning',
45
- 'mcp_query', 'search_web',
46
- ]);
47
-
48
- const MODERATE_TOOLS = new Set([
49
- 'edit', 'multi_edit', 'write_to_file', 'edit_notebook',
50
- 'memorix_store', 'memorix_store_reasoning', 'memorix_resolve',
51
- 'memorix_session_start', 'memorix_handoff', 'memorix_poll',
52
- 'git_commit', 'git_add', 'git_checkout',
53
- ]);
54
-
55
- const DANGEROUS_TOOLS = new Set([
56
- 'run_command', 'bash', 'execute_command', 'terminal',
57
- 'delete_file', 'remove_directory',
58
- 'deploy_web_app', 'browser_navigate', 'browser_click',
59
- 'http_request', 'fetch_url',
60
- ]);
61
-
62
- // ── Core ───────────────────────────────────────────────────────────
63
-
64
- /**
65
- * Classify a tool name into a risk tier.
66
- * Unknown tools default to 'moderate' (conservative but not alarming).
67
- */
68
- export function classifyTool(toolName: string): RiskTier {
69
- const name = toolName.toLowerCase().replace(/[^a-z0-9_]/g, '_');
70
- if (SAFE_TOOLS.has(name)) return 'safe';
71
- if (DANGEROUS_TOOLS.has(name)) return 'dangerous';
72
- if (MODERATE_TOOLS.has(name)) return 'moderate';
73
- // Heuristic fallback: tools with 'write', 'delete', 'exec', 'run' in name
74
- if (/delete|remove|exec|run|deploy|bash|terminal/.test(name)) return 'dangerous';
75
- if (/write|edit|create|update|store|commit/.test(name)) return 'moderate';
76
- return 'moderate'; // Unknown → moderate (safe default)
77
- }
78
-
79
- /**
80
- * Task-level tool usage tracker. Create one per dispatched task.
81
- */
82
- export class TaskToolTracker {
83
- private tools = new Map<string, ToolUsageRecord>();
84
- readonly taskId: string;
85
-
86
- constructor(taskId: string) {
87
- this.taskId = taskId;
88
- }
89
-
90
- /** Record a tool call. Call this when an agent uses a tool. */
91
- record(toolName: string): void {
92
- const tier = classifyTool(toolName);
93
- const now = Date.now();
94
- const existing = this.tools.get(toolName);
95
-
96
- if (existing) {
97
- existing.count++;
98
- existing.lastSeen = now;
99
- } else {
100
- this.tools.set(toolName, {
101
- tool: toolName,
102
- tier,
103
- count: 1,
104
- firstSeen: now,
105
- lastSeen: now,
106
- });
107
- }
108
- }
109
-
110
- /** Get the risk profile for this task. */
111
- getProfile(): TaskRiskProfile {
112
- const tools = Array.from(this.tools.values());
113
- const totalCalls = tools.reduce((sum, t) => sum + t.count, 0);
114
-
115
- let maxTier: RiskTier = 'safe';
116
- for (const t of tools) {
117
- if (t.tier === 'dangerous') { maxTier = 'dangerous'; break; }
118
- if (t.tier === 'moderate') maxTier = 'moderate';
119
- }
120
-
121
- return {
122
- taskId: this.taskId,
123
- maxTier,
124
- totalCalls,
125
- tools,
126
- };
127
- }
128
-
129
- /** Format a human-readable risk summary. */
130
- formatSummary(): string {
131
- const profile = this.getProfile();
132
- const lines = [`Risk: ${profile.maxTier} (${profile.totalCalls} tool calls)`];
133
-
134
- const dangerous = profile.tools.filter(t => t.tier === 'dangerous');
135
- if (dangerous.length > 0) {
136
- lines.push(' Dangerous tools:');
137
- for (const t of dangerous) {
138
- lines.push(` - ${t.tool}: ${t.count}x`);
139
- }
140
- }
141
-
142
- return lines.join('\n');
143
- }
144
- }
145
-
146
- /**
147
- * Compare two risk tiers. Returns positive if a > b, negative if a < b, 0 if equal.
148
- */
149
- export function compareTiers(a: RiskTier, b: RiskTier): number {
150
- const order: Record<RiskTier, number> = { safe: 0, moderate: 1, dangerous: 2 };
151
- return order[a] - order[b];
152
- }
1
+ /**
2
+ * Permission — Phase 7, Step 10: Risk classification + monitoring for agent tools.
3
+ *
4
+ * Classifies tool calls by risk tier and records them for post-run audit.
5
+ * Does NOT block or restrict agent tools (Non-Invasiveness principle).
6
+ * The primary use case is monitoring and reporting:
7
+ * - Which high-risk tools were used and how often
8
+ * - Whether agents used filesystem-destructive or network operations
9
+ * - Per-task risk summary in the evidence directory
10
+ *
11
+ * Three tiers:
12
+ * - safe: read-only operations (Read, Grep, List, etc.)
13
+ * - moderate: write operations (Edit, Write, Execute safe commands)
14
+ * - dangerous: destructive/network/exec operations (Delete, Bash, Deploy, etc.)
15
+ */
16
+
17
+ // ── Types ──────────────────────────────────────────────────────────
18
+
19
+ export type RiskTier = 'safe' | 'moderate' | 'dangerous';
20
+
21
+ export interface ToolUsageRecord {
22
+ tool: string;
23
+ tier: RiskTier;
24
+ count: number;
25
+ firstSeen: number;
26
+ lastSeen: number;
27
+ }
28
+
29
+ export interface TaskRiskProfile {
30
+ taskId: string;
31
+ /** Highest risk tier used by this task */
32
+ maxTier: RiskTier;
33
+ /** Total tool calls */
34
+ totalCalls: number;
35
+ /** Tool usage breakdown */
36
+ tools: ToolUsageRecord[];
37
+ }
38
+
39
+ // ── Classification Registry ────────────────────────────────────────
40
+
41
+ const SAFE_TOOLS = new Set([
42
+ 'read_file', 'grep_search', 'find_by_name', 'list_dir', 'code_search',
43
+ 'read_notebook', 'read_url_content', 'view_content_chunk',
44
+ 'memorix_search', 'memorix_detail', 'memorix_search_reasoning',
45
+ 'mcp_query', 'search_web',
46
+ ]);
47
+
48
+ const MODERATE_TOOLS = new Set([
49
+ 'edit', 'multi_edit', 'write_to_file', 'edit_notebook',
50
+ 'memorix_store', 'memorix_store_reasoning', 'memorix_resolve',
51
+ 'memorix_session_start', 'memorix_handoff', 'memorix_poll',
52
+ 'git_commit', 'git_add', 'git_checkout',
53
+ ]);
54
+
55
+ const DANGEROUS_TOOLS = new Set([
56
+ 'run_command', 'bash', 'execute_command', 'terminal',
57
+ 'delete_file', 'remove_directory',
58
+ 'deploy_web_app', 'browser_navigate', 'browser_click',
59
+ 'http_request', 'fetch_url',
60
+ ]);
61
+
62
+ // ── Core ───────────────────────────────────────────────────────────
63
+
64
+ /**
65
+ * Classify a tool name into a risk tier.
66
+ * Unknown tools default to 'moderate' (conservative but not alarming).
67
+ */
68
+ export function classifyTool(toolName: string): RiskTier {
69
+ const name = toolName.toLowerCase().replace(/[^a-z0-9_]/g, '_');
70
+ if (SAFE_TOOLS.has(name)) return 'safe';
71
+ if (DANGEROUS_TOOLS.has(name)) return 'dangerous';
72
+ if (MODERATE_TOOLS.has(name)) return 'moderate';
73
+ // Heuristic fallback: tools with 'write', 'delete', 'exec', 'run' in name
74
+ if (/delete|remove|exec|run|deploy|bash|terminal/.test(name)) return 'dangerous';
75
+ if (/write|edit|create|update|store|commit/.test(name)) return 'moderate';
76
+ return 'moderate'; // Unknown → moderate (safe default)
77
+ }
78
+
79
+ /**
80
+ * Task-level tool usage tracker. Create one per dispatched task.
81
+ */
82
+ export class TaskToolTracker {
83
+ private tools = new Map<string, ToolUsageRecord>();
84
+ readonly taskId: string;
85
+
86
+ constructor(taskId: string) {
87
+ this.taskId = taskId;
88
+ }
89
+
90
+ /** Record a tool call. Call this when an agent uses a tool. */
91
+ record(toolName: string): void {
92
+ const tier = classifyTool(toolName);
93
+ const now = Date.now();
94
+ const existing = this.tools.get(toolName);
95
+
96
+ if (existing) {
97
+ existing.count++;
98
+ existing.lastSeen = now;
99
+ } else {
100
+ this.tools.set(toolName, {
101
+ tool: toolName,
102
+ tier,
103
+ count: 1,
104
+ firstSeen: now,
105
+ lastSeen: now,
106
+ });
107
+ }
108
+ }
109
+
110
+ /** Get the risk profile for this task. */
111
+ getProfile(): TaskRiskProfile {
112
+ const tools = Array.from(this.tools.values());
113
+ const totalCalls = tools.reduce((sum, t) => sum + t.count, 0);
114
+
115
+ let maxTier: RiskTier = 'safe';
116
+ for (const t of tools) {
117
+ if (t.tier === 'dangerous') { maxTier = 'dangerous'; break; }
118
+ if (t.tier === 'moderate') maxTier = 'moderate';
119
+ }
120
+
121
+ return {
122
+ taskId: this.taskId,
123
+ maxTier,
124
+ totalCalls,
125
+ tools,
126
+ };
127
+ }
128
+
129
+ /** Format a human-readable risk summary. */
130
+ formatSummary(): string {
131
+ const profile = this.getProfile();
132
+ const lines = [`Risk: ${profile.maxTier} (${profile.totalCalls} tool calls)`];
133
+
134
+ const dangerous = profile.tools.filter(t => t.tier === 'dangerous');
135
+ if (dangerous.length > 0) {
136
+ lines.push(' Dangerous tools:');
137
+ for (const t of dangerous) {
138
+ lines.push(` - ${t.tool}: ${t.count}x`);
139
+ }
140
+ }
141
+
142
+ return lines.join('\n');
143
+ }
144
+ }
145
+
146
+ /**
147
+ * Compare two risk tiers. Returns positive if a > b, negative if a < b, 0 if equal.
148
+ */
149
+ export function compareTiers(a: RiskTier, b: RiskTier): number {
150
+ const order: Record<RiskTier, number> = { safe: 0, moderate: 1, dangerous: 2 };
151
+ return order[a] - order[b];
152
+ }