memorix 1.2.2 → 1.2.4

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (163) hide show
  1. package/CHANGELOG.md +27 -0
  2. package/README.md +3 -3
  3. package/README.zh-CN.md +3 -3
  4. package/TEAM.md +86 -86
  5. package/dist/cli/index.js +5199 -4726
  6. package/dist/cli/index.js.map +1 -1
  7. package/dist/index.js +428 -49
  8. package/dist/index.js.map +1 -1
  9. package/dist/maintenance-runner.js +97 -18
  10. package/dist/maintenance-runner.js.map +1 -1
  11. package/dist/memcode-runtime/CHANGELOG.md +27 -0
  12. package/dist/sdk.js +428 -49
  13. package/dist/sdk.js.map +1 -1
  14. package/docs/1.2.4-PERSISTENT-MEMORY-DELIVERY.md +86 -0
  15. package/docs/AGENT_OPERATOR_PLAYBOOK.md +13 -1
  16. package/docs/API_REFERENCE.md +13 -3
  17. package/docs/DESIGN_DECISIONS.md +357 -357
  18. package/docs/dev-log/progress.txt +60 -9
  19. package/package.json +1 -1
  20. package/plugins/codex/memorix/.codex-plugin/plugin.json +1 -1
  21. package/src/audit/index.ts +156 -156
  22. package/src/cli/capability-map.ts +1 -1
  23. package/src/cli/command-guide.ts +4 -1
  24. package/src/cli/commands/agent-integrations.ts +5 -1
  25. package/src/cli/commands/audit-list.ts +89 -89
  26. package/src/cli/commands/background.ts +659 -659
  27. package/src/cli/commands/codegraph.ts +1 -1
  28. package/src/cli/commands/context.ts +9 -1
  29. package/src/cli/commands/formation.ts +48 -48
  30. package/src/cli/commands/git-hook-install.ts +111 -111
  31. package/src/cli/commands/handoff.ts +54 -54
  32. package/src/cli/commands/hooks-status.ts +63 -63
  33. package/src/cli/commands/ingest-commit.ts +153 -153
  34. package/src/cli/commands/ingest-image.ts +66 -66
  35. package/src/cli/commands/ingest-log.ts +180 -180
  36. package/src/cli/commands/ingest.ts +44 -44
  37. package/src/cli/commands/integrate-shared.ts +15 -15
  38. package/src/cli/commands/lock.ts +82 -82
  39. package/src/cli/commands/message.ts +104 -104
  40. package/src/cli/commands/poll.ts +58 -58
  41. package/src/cli/commands/purge-all-memory.ts +85 -85
  42. package/src/cli/commands/purge-project-memory.ts +83 -83
  43. package/src/cli/commands/reasoning.ts +118 -118
  44. package/src/cli/commands/resume.ts +31 -0
  45. package/src/cli/commands/serve-shared.ts +118 -118
  46. package/src/cli/commands/session.ts +15 -7
  47. package/src/cli/commands/skills.ts +114 -114
  48. package/src/cli/commands/task.ts +167 -167
  49. package/src/cli/commands/transfer.ts +47 -47
  50. package/src/cli/commands/uninstall-project-artifacts.ts +85 -85
  51. package/src/cli/index.ts +3 -1
  52. package/src/cli/tui/ChatView.tsx +234 -234
  53. package/src/cli/tui/CommandBar.tsx +312 -312
  54. package/src/cli/tui/ContextRail.tsx +118 -118
  55. package/src/cli/tui/HeaderBar.tsx +72 -72
  56. package/src/cli/tui/LogoBanner.tsx +51 -51
  57. package/src/cli/tui/Sidebar.tsx +179 -179
  58. package/src/cli/tui/index.ts +41 -41
  59. package/src/cli/tui/markdown-render.tsx +371 -371
  60. package/src/cli/tui/session-service.ts +3 -2
  61. package/src/cli/tui/use-mouse.ts +157 -157
  62. package/src/cli/tui/useNavigation.ts +56 -56
  63. package/src/cli/update-checker.ts +211 -211
  64. package/src/cli/version.ts +7 -7
  65. package/src/cli/workbench.ts +1 -1
  66. package/src/codegraph/auto-context.ts +54 -1
  67. package/src/codegraph/task-lens.ts +29 -0
  68. package/src/compact/token-budget.ts +89 -74
  69. package/src/config/toml-loader.ts +9 -5
  70. package/src/dashboard/project-classification.ts +64 -64
  71. package/src/embedding/fastembed-provider.ts +142 -142
  72. package/src/embedding/transformers-provider.ts +111 -111
  73. package/src/git/extractor.ts +209 -209
  74. package/src/git/hooks-path.ts +85 -85
  75. package/src/hooks/handler.ts +127 -66
  76. package/src/hooks/installers/index.ts +5 -4
  77. package/src/hooks/official-skills.ts +6 -4
  78. package/src/hooks/pattern-detector.ts +173 -173
  79. package/src/hooks/rules/memorix-agent-rules.md +9 -7
  80. package/src/hooks/significance-filter.ts +250 -250
  81. package/src/knowledge/context-assembly.ts +4 -1
  82. package/src/knowledge/workset.ts +89 -1
  83. package/src/llm/memory-manager.ts +328 -328
  84. package/src/llm/provider.ts +885 -885
  85. package/src/llm/quality.ts +248 -248
  86. package/src/memory/attribution-guard.ts +249 -249
  87. package/src/memory/disclosure-policy.ts +135 -135
  88. package/src/memory/entity-extractor.ts +197 -197
  89. package/src/memory/formation/evaluate.ts +217 -217
  90. package/src/memory/formation/extract.ts +361 -361
  91. package/src/memory/formation/index.ts +417 -417
  92. package/src/memory/formation/resolve.ts +344 -344
  93. package/src/memory/formation/types.ts +315 -315
  94. package/src/memory/freshness.ts +122 -122
  95. package/src/memory/graph.ts +197 -197
  96. package/src/memory/refs.ts +94 -94
  97. package/src/memory/secret-filter.ts +79 -79
  98. package/src/memory/session.ts +158 -9
  99. package/src/multimodal/image-loader.ts +143 -143
  100. package/src/orchestrate/adapters/claude-stream.ts +192 -192
  101. package/src/orchestrate/adapters/claude.ts +111 -111
  102. package/src/orchestrate/adapters/codex-stream.ts +134 -134
  103. package/src/orchestrate/adapters/codex.ts +41 -41
  104. package/src/orchestrate/adapters/gemini-stream.ts +166 -166
  105. package/src/orchestrate/adapters/gemini.ts +42 -42
  106. package/src/orchestrate/adapters/index.ts +73 -73
  107. package/src/orchestrate/adapters/opencode-stream.ts +143 -143
  108. package/src/orchestrate/adapters/opencode.ts +47 -47
  109. package/src/orchestrate/adapters/spawn-helper.ts +286 -286
  110. package/src/orchestrate/adapters/types.ts +77 -77
  111. package/src/orchestrate/capability-router.ts +284 -284
  112. package/src/orchestrate/context-compact.ts +188 -188
  113. package/src/orchestrate/cost-tracker.ts +219 -219
  114. package/src/orchestrate/error-recovery.ts +191 -191
  115. package/src/orchestrate/evidence.ts +140 -140
  116. package/src/orchestrate/ledger.ts +110 -110
  117. package/src/orchestrate/memorix-bridge.ts +343 -343
  118. package/src/orchestrate/output-budget.ts +80 -80
  119. package/src/orchestrate/permission.ts +152 -152
  120. package/src/orchestrate/pipeline-trace.ts +131 -131
  121. package/src/orchestrate/prompt-builder.ts +155 -155
  122. package/src/orchestrate/ring-buffer.ts +37 -37
  123. package/src/orchestrate/task-graph.ts +389 -389
  124. package/src/orchestrate/worktree.ts +232 -232
  125. package/src/project/aliases.ts +374 -374
  126. package/src/project/detector.ts +268 -268
  127. package/src/rules/adapters/claude-code.ts +99 -99
  128. package/src/rules/adapters/codex.ts +97 -97
  129. package/src/rules/adapters/copilot.ts +124 -124
  130. package/src/rules/adapters/cursor.ts +114 -114
  131. package/src/rules/adapters/kiro.ts +126 -126
  132. package/src/rules/adapters/trae.ts +56 -56
  133. package/src/rules/adapters/windsurf.ts +83 -83
  134. package/src/rules/syncer.ts +235 -235
  135. package/src/sdk.ts +299 -299
  136. package/src/search/intent-detector.ts +289 -289
  137. package/src/search/query-expansion.ts +52 -52
  138. package/src/server/formation-timeout.ts +27 -27
  139. package/src/server.ts +144 -10
  140. package/src/skills/mini-skills.ts +386 -386
  141. package/src/store/bun-sqlite-compat.ts +118 -15
  142. package/src/store/chat-store.ts +119 -119
  143. package/src/store/graph-store.ts +249 -249
  144. package/src/store/mini-skill-store.ts +349 -349
  145. package/src/store/persistence-json.ts +212 -212
  146. package/src/store/persistence.ts +291 -291
  147. package/src/store/project-affinity.ts +195 -195
  148. package/src/store/sqlite-db.ts +3 -3
  149. package/src/team/event-bus.ts +76 -76
  150. package/src/team/file-locks.ts +173 -173
  151. package/src/team/handoff.ts +161 -161
  152. package/src/team/messages.ts +203 -203
  153. package/src/team/poll.ts +132 -132
  154. package/src/team/tasks.ts +211 -211
  155. package/src/workspace/mcp-adapters/codex.ts +191 -191
  156. package/src/workspace/mcp-adapters/copilot.ts +105 -105
  157. package/src/workspace/mcp-adapters/cursor.ts +53 -53
  158. package/src/workspace/mcp-adapters/kiro.ts +64 -64
  159. package/src/workspace/mcp-adapters/opencode.ts +123 -123
  160. package/src/workspace/mcp-adapters/trae.ts +134 -134
  161. package/src/workspace/mcp-adapters/windsurf.ts +91 -91
  162. package/src/workspace/sanitizer.ts +60 -60
  163. package/src/workspace/workflow-sync.ts +131 -131
@@ -1,80 +1,80 @@
1
- /**
2
- * Output Budget — Phase 7, Step 4: Trim large outputs to a token-friendly size.
3
- *
4
- * When gate outputs, agent tail outputs, or other large strings need to be
5
- * embedded in prompts (e.g., fix prompts, ledger entries), this module trims
6
- * them to a configurable budget while preserving the most useful parts
7
- * (beginning + end).
8
- *
9
- * Full output is optionally persisted to disk for debugging.
10
- */
11
-
12
- import { writeFileSync, mkdirSync } from 'node:fs';
13
- import { join, dirname } from 'node:path';
14
-
15
- // ── Constants ──────────────────────────────────────────────────────
16
-
17
- /** Default output budget in bytes */
18
- export const DEFAULT_OUTPUT_BUDGET = 4096;
19
-
20
- // ── Core ───────────────────────────────────────────────────────────
21
-
22
- /**
23
- * Trim text to a byte budget. If the text exceeds the budget, keep
24
- * the first half and last half with an omission marker in between.
25
- *
26
- * Returns the trimmed text. Never throws.
27
- */
28
- export function trimToBudget(text: string, budget: number = DEFAULT_OUTPUT_BUDGET): string {
29
- if (text.length <= budget) return text;
30
-
31
- const half = Math.floor(budget / 2) - 40; // Reserve space for marker
32
- if (half <= 0) return text.slice(0, budget);
33
-
34
- const omitted = text.length - half * 2;
35
- return (
36
- text.slice(0, half) +
37
- `\n\n... (${omitted} bytes omitted — full output saved to disk) ...\n\n` +
38
- text.slice(-half)
39
- );
40
- }
41
-
42
- /**
43
- * Persist full output to a file and return the trimmed version.
44
- *
45
- * If disk write fails, returns the trimmed text without file reference.
46
- * Never throws.
47
- */
48
- export function trimAndPersist(
49
- text: string,
50
- filePath: string,
51
- budget: number = DEFAULT_OUTPUT_BUDGET,
52
- ): { trimmed: string; persisted: boolean; fullPath?: string } {
53
- if (text.length <= budget) {
54
- return { trimmed: text, persisted: false };
55
- }
56
-
57
- // Persist full output to disk (best-effort)
58
- let persisted = false;
59
- try {
60
- mkdirSync(dirname(filePath), { recursive: true });
61
- writeFileSync(filePath, text, 'utf-8');
62
- persisted = true;
63
- } catch {
64
- // Disk write failed — degrade gracefully
65
- }
66
-
67
- const half = Math.floor(budget / 2) - 60; // Reserve space for marker + path
68
- if (half <= 0) {
69
- return { trimmed: text.slice(0, budget), persisted, fullPath: persisted ? filePath : undefined };
70
- }
71
-
72
- const omitted = text.length - half * 2;
73
- const pathNote = persisted ? ` Full output at: ${filePath}` : '';
74
- const trimmed =
75
- text.slice(0, half) +
76
- `\n\n... (${omitted} bytes omitted.${pathNote}) ...\n\n` +
77
- text.slice(-half);
78
-
79
- return { trimmed, persisted, fullPath: persisted ? filePath : undefined };
80
- }
1
+ /**
2
+ * Output Budget — Phase 7, Step 4: Trim large outputs to a token-friendly size.
3
+ *
4
+ * When gate outputs, agent tail outputs, or other large strings need to be
5
+ * embedded in prompts (e.g., fix prompts, ledger entries), this module trims
6
+ * them to a configurable budget while preserving the most useful parts
7
+ * (beginning + end).
8
+ *
9
+ * Full output is optionally persisted to disk for debugging.
10
+ */
11
+
12
+ import { writeFileSync, mkdirSync } from 'node:fs';
13
+ import { join, dirname } from 'node:path';
14
+
15
+ // ── Constants ──────────────────────────────────────────────────────
16
+
17
+ /** Default output budget in bytes */
18
+ export const DEFAULT_OUTPUT_BUDGET = 4096;
19
+
20
+ // ── Core ───────────────────────────────────────────────────────────
21
+
22
+ /**
23
+ * Trim text to a byte budget. If the text exceeds the budget, keep
24
+ * the first half and last half with an omission marker in between.
25
+ *
26
+ * Returns the trimmed text. Never throws.
27
+ */
28
+ export function trimToBudget(text: string, budget: number = DEFAULT_OUTPUT_BUDGET): string {
29
+ if (text.length <= budget) return text;
30
+
31
+ const half = Math.floor(budget / 2) - 40; // Reserve space for marker
32
+ if (half <= 0) return text.slice(0, budget);
33
+
34
+ const omitted = text.length - half * 2;
35
+ return (
36
+ text.slice(0, half) +
37
+ `\n\n... (${omitted} bytes omitted — full output saved to disk) ...\n\n` +
38
+ text.slice(-half)
39
+ );
40
+ }
41
+
42
+ /**
43
+ * Persist full output to a file and return the trimmed version.
44
+ *
45
+ * If disk write fails, returns the trimmed text without file reference.
46
+ * Never throws.
47
+ */
48
+ export function trimAndPersist(
49
+ text: string,
50
+ filePath: string,
51
+ budget: number = DEFAULT_OUTPUT_BUDGET,
52
+ ): { trimmed: string; persisted: boolean; fullPath?: string } {
53
+ if (text.length <= budget) {
54
+ return { trimmed: text, persisted: false };
55
+ }
56
+
57
+ // Persist full output to disk (best-effort)
58
+ let persisted = false;
59
+ try {
60
+ mkdirSync(dirname(filePath), { recursive: true });
61
+ writeFileSync(filePath, text, 'utf-8');
62
+ persisted = true;
63
+ } catch {
64
+ // Disk write failed — degrade gracefully
65
+ }
66
+
67
+ const half = Math.floor(budget / 2) - 60; // Reserve space for marker + path
68
+ if (half <= 0) {
69
+ return { trimmed: text.slice(0, budget), persisted, fullPath: persisted ? filePath : undefined };
70
+ }
71
+
72
+ const omitted = text.length - half * 2;
73
+ const pathNote = persisted ? ` Full output at: ${filePath}` : '';
74
+ const trimmed =
75
+ text.slice(0, half) +
76
+ `\n\n... (${omitted} bytes omitted.${pathNote}) ...\n\n` +
77
+ text.slice(-half);
78
+
79
+ return { trimmed, persisted, fullPath: persisted ? filePath : undefined };
80
+ }
@@ -1,152 +1,152 @@
1
- /**
2
- * Permission — Phase 7, Step 10: Risk classification + monitoring for agent tools.
3
- *
4
- * Classifies tool calls by risk tier and records them for post-run audit.
5
- * Does NOT block or restrict agent tools (Non-Invasiveness principle).
6
- * The primary use case is monitoring and reporting:
7
- * - Which high-risk tools were used and how often
8
- * - Whether agents used filesystem-destructive or network operations
9
- * - Per-task risk summary in the evidence directory
10
- *
11
- * Three tiers:
12
- * - safe: read-only operations (Read, Grep, List, etc.)
13
- * - moderate: write operations (Edit, Write, Execute safe commands)
14
- * - dangerous: destructive/network/exec operations (Delete, Bash, Deploy, etc.)
15
- */
16
-
17
- // ── Types ──────────────────────────────────────────────────────────
18
-
19
- export type RiskTier = 'safe' | 'moderate' | 'dangerous';
20
-
21
- export interface ToolUsageRecord {
22
- tool: string;
23
- tier: RiskTier;
24
- count: number;
25
- firstSeen: number;
26
- lastSeen: number;
27
- }
28
-
29
- export interface TaskRiskProfile {
30
- taskId: string;
31
- /** Highest risk tier used by this task */
32
- maxTier: RiskTier;
33
- /** Total tool calls */
34
- totalCalls: number;
35
- /** Tool usage breakdown */
36
- tools: ToolUsageRecord[];
37
- }
38
-
39
- // ── Classification Registry ────────────────────────────────────────
40
-
41
- const SAFE_TOOLS = new Set([
42
- 'read_file', 'grep_search', 'find_by_name', 'list_dir', 'code_search',
43
- 'read_notebook', 'read_url_content', 'view_content_chunk',
44
- 'memorix_search', 'memorix_detail', 'memorix_search_reasoning',
45
- 'mcp_query', 'search_web',
46
- ]);
47
-
48
- const MODERATE_TOOLS = new Set([
49
- 'edit', 'multi_edit', 'write_to_file', 'edit_notebook',
50
- 'memorix_store', 'memorix_store_reasoning', 'memorix_resolve',
51
- 'memorix_session_start', 'memorix_handoff', 'memorix_poll',
52
- 'git_commit', 'git_add', 'git_checkout',
53
- ]);
54
-
55
- const DANGEROUS_TOOLS = new Set([
56
- 'run_command', 'bash', 'execute_command', 'terminal',
57
- 'delete_file', 'remove_directory',
58
- 'deploy_web_app', 'browser_navigate', 'browser_click',
59
- 'http_request', 'fetch_url',
60
- ]);
61
-
62
- // ── Core ───────────────────────────────────────────────────────────
63
-
64
- /**
65
- * Classify a tool name into a risk tier.
66
- * Unknown tools default to 'moderate' (conservative but not alarming).
67
- */
68
- export function classifyTool(toolName: string): RiskTier {
69
- const name = toolName.toLowerCase().replace(/[^a-z0-9_]/g, '_');
70
- if (SAFE_TOOLS.has(name)) return 'safe';
71
- if (DANGEROUS_TOOLS.has(name)) return 'dangerous';
72
- if (MODERATE_TOOLS.has(name)) return 'moderate';
73
- // Heuristic fallback: tools with 'write', 'delete', 'exec', 'run' in name
74
- if (/delete|remove|exec|run|deploy|bash|terminal/.test(name)) return 'dangerous';
75
- if (/write|edit|create|update|store|commit/.test(name)) return 'moderate';
76
- return 'moderate'; // Unknown → moderate (safe default)
77
- }
78
-
79
- /**
80
- * Task-level tool usage tracker. Create one per dispatched task.
81
- */
82
- export class TaskToolTracker {
83
- private tools = new Map<string, ToolUsageRecord>();
84
- readonly taskId: string;
85
-
86
- constructor(taskId: string) {
87
- this.taskId = taskId;
88
- }
89
-
90
- /** Record a tool call. Call this when an agent uses a tool. */
91
- record(toolName: string): void {
92
- const tier = classifyTool(toolName);
93
- const now = Date.now();
94
- const existing = this.tools.get(toolName);
95
-
96
- if (existing) {
97
- existing.count++;
98
- existing.lastSeen = now;
99
- } else {
100
- this.tools.set(toolName, {
101
- tool: toolName,
102
- tier,
103
- count: 1,
104
- firstSeen: now,
105
- lastSeen: now,
106
- });
107
- }
108
- }
109
-
110
- /** Get the risk profile for this task. */
111
- getProfile(): TaskRiskProfile {
112
- const tools = Array.from(this.tools.values());
113
- const totalCalls = tools.reduce((sum, t) => sum + t.count, 0);
114
-
115
- let maxTier: RiskTier = 'safe';
116
- for (const t of tools) {
117
- if (t.tier === 'dangerous') { maxTier = 'dangerous'; break; }
118
- if (t.tier === 'moderate') maxTier = 'moderate';
119
- }
120
-
121
- return {
122
- taskId: this.taskId,
123
- maxTier,
124
- totalCalls,
125
- tools,
126
- };
127
- }
128
-
129
- /** Format a human-readable risk summary. */
130
- formatSummary(): string {
131
- const profile = this.getProfile();
132
- const lines = [`Risk: ${profile.maxTier} (${profile.totalCalls} tool calls)`];
133
-
134
- const dangerous = profile.tools.filter(t => t.tier === 'dangerous');
135
- if (dangerous.length > 0) {
136
- lines.push(' Dangerous tools:');
137
- for (const t of dangerous) {
138
- lines.push(` - ${t.tool}: ${t.count}x`);
139
- }
140
- }
141
-
142
- return lines.join('\n');
143
- }
144
- }
145
-
146
- /**
147
- * Compare two risk tiers. Returns positive if a > b, negative if a < b, 0 if equal.
148
- */
149
- export function compareTiers(a: RiskTier, b: RiskTier): number {
150
- const order: Record<RiskTier, number> = { safe: 0, moderate: 1, dangerous: 2 };
151
- return order[a] - order[b];
152
- }
1
+ /**
2
+ * Permission — Phase 7, Step 10: Risk classification + monitoring for agent tools.
3
+ *
4
+ * Classifies tool calls by risk tier and records them for post-run audit.
5
+ * Does NOT block or restrict agent tools (Non-Invasiveness principle).
6
+ * The primary use case is monitoring and reporting:
7
+ * - Which high-risk tools were used and how often
8
+ * - Whether agents used filesystem-destructive or network operations
9
+ * - Per-task risk summary in the evidence directory
10
+ *
11
+ * Three tiers:
12
+ * - safe: read-only operations (Read, Grep, List, etc.)
13
+ * - moderate: write operations (Edit, Write, Execute safe commands)
14
+ * - dangerous: destructive/network/exec operations (Delete, Bash, Deploy, etc.)
15
+ */
16
+
17
+ // ── Types ──────────────────────────────────────────────────────────
18
+
19
+ export type RiskTier = 'safe' | 'moderate' | 'dangerous';
20
+
21
+ export interface ToolUsageRecord {
22
+ tool: string;
23
+ tier: RiskTier;
24
+ count: number;
25
+ firstSeen: number;
26
+ lastSeen: number;
27
+ }
28
+
29
+ export interface TaskRiskProfile {
30
+ taskId: string;
31
+ /** Highest risk tier used by this task */
32
+ maxTier: RiskTier;
33
+ /** Total tool calls */
34
+ totalCalls: number;
35
+ /** Tool usage breakdown */
36
+ tools: ToolUsageRecord[];
37
+ }
38
+
39
+ // ── Classification Registry ────────────────────────────────────────
40
+
41
+ const SAFE_TOOLS = new Set([
42
+ 'read_file', 'grep_search', 'find_by_name', 'list_dir', 'code_search',
43
+ 'read_notebook', 'read_url_content', 'view_content_chunk',
44
+ 'memorix_search', 'memorix_detail', 'memorix_search_reasoning',
45
+ 'mcp_query', 'search_web',
46
+ ]);
47
+
48
+ const MODERATE_TOOLS = new Set([
49
+ 'edit', 'multi_edit', 'write_to_file', 'edit_notebook',
50
+ 'memorix_store', 'memorix_store_reasoning', 'memorix_resolve',
51
+ 'memorix_session_start', 'memorix_handoff', 'memorix_poll',
52
+ 'git_commit', 'git_add', 'git_checkout',
53
+ ]);
54
+
55
+ const DANGEROUS_TOOLS = new Set([
56
+ 'run_command', 'bash', 'execute_command', 'terminal',
57
+ 'delete_file', 'remove_directory',
58
+ 'deploy_web_app', 'browser_navigate', 'browser_click',
59
+ 'http_request', 'fetch_url',
60
+ ]);
61
+
62
+ // ── Core ───────────────────────────────────────────────────────────
63
+
64
+ /**
65
+ * Classify a tool name into a risk tier.
66
+ * Unknown tools default to 'moderate' (conservative but not alarming).
67
+ */
68
+ export function classifyTool(toolName: string): RiskTier {
69
+ const name = toolName.toLowerCase().replace(/[^a-z0-9_]/g, '_');
70
+ if (SAFE_TOOLS.has(name)) return 'safe';
71
+ if (DANGEROUS_TOOLS.has(name)) return 'dangerous';
72
+ if (MODERATE_TOOLS.has(name)) return 'moderate';
73
+ // Heuristic fallback: tools with 'write', 'delete', 'exec', 'run' in name
74
+ if (/delete|remove|exec|run|deploy|bash|terminal/.test(name)) return 'dangerous';
75
+ if (/write|edit|create|update|store|commit/.test(name)) return 'moderate';
76
+ return 'moderate'; // Unknown → moderate (safe default)
77
+ }
78
+
79
+ /**
80
+ * Task-level tool usage tracker. Create one per dispatched task.
81
+ */
82
+ export class TaskToolTracker {
83
+ private tools = new Map<string, ToolUsageRecord>();
84
+ readonly taskId: string;
85
+
86
+ constructor(taskId: string) {
87
+ this.taskId = taskId;
88
+ }
89
+
90
+ /** Record a tool call. Call this when an agent uses a tool. */
91
+ record(toolName: string): void {
92
+ const tier = classifyTool(toolName);
93
+ const now = Date.now();
94
+ const existing = this.tools.get(toolName);
95
+
96
+ if (existing) {
97
+ existing.count++;
98
+ existing.lastSeen = now;
99
+ } else {
100
+ this.tools.set(toolName, {
101
+ tool: toolName,
102
+ tier,
103
+ count: 1,
104
+ firstSeen: now,
105
+ lastSeen: now,
106
+ });
107
+ }
108
+ }
109
+
110
+ /** Get the risk profile for this task. */
111
+ getProfile(): TaskRiskProfile {
112
+ const tools = Array.from(this.tools.values());
113
+ const totalCalls = tools.reduce((sum, t) => sum + t.count, 0);
114
+
115
+ let maxTier: RiskTier = 'safe';
116
+ for (const t of tools) {
117
+ if (t.tier === 'dangerous') { maxTier = 'dangerous'; break; }
118
+ if (t.tier === 'moderate') maxTier = 'moderate';
119
+ }
120
+
121
+ return {
122
+ taskId: this.taskId,
123
+ maxTier,
124
+ totalCalls,
125
+ tools,
126
+ };
127
+ }
128
+
129
+ /** Format a human-readable risk summary. */
130
+ formatSummary(): string {
131
+ const profile = this.getProfile();
132
+ const lines = [`Risk: ${profile.maxTier} (${profile.totalCalls} tool calls)`];
133
+
134
+ const dangerous = profile.tools.filter(t => t.tier === 'dangerous');
135
+ if (dangerous.length > 0) {
136
+ lines.push(' Dangerous tools:');
137
+ for (const t of dangerous) {
138
+ lines.push(` - ${t.tool}: ${t.count}x`);
139
+ }
140
+ }
141
+
142
+ return lines.join('\n');
143
+ }
144
+ }
145
+
146
+ /**
147
+ * Compare two risk tiers. Returns positive if a > b, negative if a < b, 0 if equal.
148
+ */
149
+ export function compareTiers(a: RiskTier, b: RiskTier): number {
150
+ const order: Record<RiskTier, number> = { safe: 0, moderate: 1, dangerous: 2 };
151
+ return order[a] - order[b];
152
+ }