memorix 1.2.2 → 1.2.3

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (139) hide show
  1. package/CHANGELOG.md +7 -0
  2. package/TEAM.md +86 -86
  3. package/dist/cli/index.js +34 -18
  4. package/dist/cli/index.js.map +1 -1
  5. package/dist/index.js +17 -8
  6. package/dist/index.js.map +1 -1
  7. package/dist/maintenance-runner.js.map +1 -1
  8. package/dist/memcode-runtime/CHANGELOG.md +7 -0
  9. package/dist/sdk.js +17 -8
  10. package/dist/sdk.js.map +1 -1
  11. package/docs/DESIGN_DECISIONS.md +357 -357
  12. package/docs/dev-log/progress.txt +18 -8
  13. package/package.json +1 -1
  14. package/plugins/codex/memorix/.codex-plugin/plugin.json +1 -1
  15. package/src/audit/index.ts +156 -156
  16. package/src/cli/commands/audit-list.ts +89 -89
  17. package/src/cli/commands/background.ts +659 -659
  18. package/src/cli/commands/formation.ts +48 -48
  19. package/src/cli/commands/git-hook-install.ts +111 -111
  20. package/src/cli/commands/handoff.ts +54 -54
  21. package/src/cli/commands/hooks-status.ts +63 -63
  22. package/src/cli/commands/ingest-commit.ts +153 -153
  23. package/src/cli/commands/ingest-image.ts +66 -66
  24. package/src/cli/commands/ingest-log.ts +180 -180
  25. package/src/cli/commands/ingest.ts +44 -44
  26. package/src/cli/commands/integrate-shared.ts +15 -15
  27. package/src/cli/commands/lock.ts +82 -82
  28. package/src/cli/commands/message.ts +104 -104
  29. package/src/cli/commands/poll.ts +58 -58
  30. package/src/cli/commands/purge-all-memory.ts +85 -85
  31. package/src/cli/commands/purge-project-memory.ts +83 -83
  32. package/src/cli/commands/reasoning.ts +118 -118
  33. package/src/cli/commands/serve-shared.ts +118 -118
  34. package/src/cli/commands/session.ts +15 -7
  35. package/src/cli/commands/skills.ts +114 -114
  36. package/src/cli/commands/task.ts +167 -167
  37. package/src/cli/commands/transfer.ts +47 -47
  38. package/src/cli/commands/uninstall-project-artifacts.ts +85 -85
  39. package/src/cli/tui/ChatView.tsx +234 -234
  40. package/src/cli/tui/CommandBar.tsx +312 -312
  41. package/src/cli/tui/ContextRail.tsx +118 -118
  42. package/src/cli/tui/HeaderBar.tsx +72 -72
  43. package/src/cli/tui/LogoBanner.tsx +51 -51
  44. package/src/cli/tui/Sidebar.tsx +179 -179
  45. package/src/cli/tui/index.ts +41 -41
  46. package/src/cli/tui/markdown-render.tsx +371 -371
  47. package/src/cli/tui/session-service.ts +3 -2
  48. package/src/cli/tui/use-mouse.ts +157 -157
  49. package/src/cli/tui/useNavigation.ts +56 -56
  50. package/src/cli/update-checker.ts +211 -211
  51. package/src/cli/version.ts +7 -7
  52. package/src/cli/workbench.ts +1 -1
  53. package/src/compact/token-budget.ts +74 -74
  54. package/src/dashboard/project-classification.ts +64 -64
  55. package/src/embedding/fastembed-provider.ts +142 -142
  56. package/src/embedding/transformers-provider.ts +111 -111
  57. package/src/git/extractor.ts +209 -209
  58. package/src/git/hooks-path.ts +85 -85
  59. package/src/hooks/pattern-detector.ts +173 -173
  60. package/src/hooks/significance-filter.ts +250 -250
  61. package/src/llm/memory-manager.ts +328 -328
  62. package/src/llm/provider.ts +885 -885
  63. package/src/llm/quality.ts +248 -248
  64. package/src/memory/attribution-guard.ts +249 -249
  65. package/src/memory/disclosure-policy.ts +135 -135
  66. package/src/memory/entity-extractor.ts +197 -197
  67. package/src/memory/formation/evaluate.ts +217 -217
  68. package/src/memory/formation/extract.ts +361 -361
  69. package/src/memory/formation/index.ts +417 -417
  70. package/src/memory/formation/resolve.ts +344 -344
  71. package/src/memory/formation/types.ts +315 -315
  72. package/src/memory/freshness.ts +122 -122
  73. package/src/memory/graph.ts +197 -197
  74. package/src/memory/refs.ts +94 -94
  75. package/src/memory/secret-filter.ts +79 -79
  76. package/src/memory/session.ts +24 -9
  77. package/src/multimodal/image-loader.ts +143 -143
  78. package/src/orchestrate/adapters/claude-stream.ts +192 -192
  79. package/src/orchestrate/adapters/claude.ts +111 -111
  80. package/src/orchestrate/adapters/codex-stream.ts +134 -134
  81. package/src/orchestrate/adapters/codex.ts +41 -41
  82. package/src/orchestrate/adapters/gemini-stream.ts +166 -166
  83. package/src/orchestrate/adapters/gemini.ts +42 -42
  84. package/src/orchestrate/adapters/index.ts +73 -73
  85. package/src/orchestrate/adapters/opencode-stream.ts +143 -143
  86. package/src/orchestrate/adapters/opencode.ts +47 -47
  87. package/src/orchestrate/adapters/spawn-helper.ts +286 -286
  88. package/src/orchestrate/adapters/types.ts +77 -77
  89. package/src/orchestrate/capability-router.ts +284 -284
  90. package/src/orchestrate/context-compact.ts +188 -188
  91. package/src/orchestrate/cost-tracker.ts +219 -219
  92. package/src/orchestrate/error-recovery.ts +191 -191
  93. package/src/orchestrate/evidence.ts +140 -140
  94. package/src/orchestrate/ledger.ts +110 -110
  95. package/src/orchestrate/memorix-bridge.ts +343 -343
  96. package/src/orchestrate/output-budget.ts +80 -80
  97. package/src/orchestrate/permission.ts +152 -152
  98. package/src/orchestrate/pipeline-trace.ts +131 -131
  99. package/src/orchestrate/prompt-builder.ts +155 -155
  100. package/src/orchestrate/ring-buffer.ts +37 -37
  101. package/src/orchestrate/task-graph.ts +389 -389
  102. package/src/orchestrate/worktree.ts +232 -232
  103. package/src/project/aliases.ts +374 -374
  104. package/src/project/detector.ts +268 -268
  105. package/src/rules/adapters/claude-code.ts +99 -99
  106. package/src/rules/adapters/codex.ts +97 -97
  107. package/src/rules/adapters/copilot.ts +124 -124
  108. package/src/rules/adapters/cursor.ts +114 -114
  109. package/src/rules/adapters/kiro.ts +126 -126
  110. package/src/rules/adapters/trae.ts +56 -56
  111. package/src/rules/adapters/windsurf.ts +83 -83
  112. package/src/rules/syncer.ts +235 -235
  113. package/src/sdk.ts +299 -299
  114. package/src/search/intent-detector.ts +289 -289
  115. package/src/search/query-expansion.ts +52 -52
  116. package/src/server/formation-timeout.ts +27 -27
  117. package/src/server.ts +7 -2
  118. package/src/skills/mini-skills.ts +386 -386
  119. package/src/store/chat-store.ts +119 -119
  120. package/src/store/graph-store.ts +249 -249
  121. package/src/store/mini-skill-store.ts +349 -349
  122. package/src/store/persistence-json.ts +212 -212
  123. package/src/store/persistence.ts +291 -291
  124. package/src/store/project-affinity.ts +195 -195
  125. package/src/team/event-bus.ts +76 -76
  126. package/src/team/file-locks.ts +173 -173
  127. package/src/team/handoff.ts +161 -161
  128. package/src/team/messages.ts +203 -203
  129. package/src/team/poll.ts +132 -132
  130. package/src/team/tasks.ts +211 -211
  131. package/src/workspace/mcp-adapters/codex.ts +191 -191
  132. package/src/workspace/mcp-adapters/copilot.ts +105 -105
  133. package/src/workspace/mcp-adapters/cursor.ts +53 -53
  134. package/src/workspace/mcp-adapters/kiro.ts +64 -64
  135. package/src/workspace/mcp-adapters/opencode.ts +123 -123
  136. package/src/workspace/mcp-adapters/trae.ts +134 -134
  137. package/src/workspace/mcp-adapters/windsurf.ts +91 -91
  138. package/src/workspace/sanitizer.ts +60 -60
  139. package/src/workspace/workflow-sync.ts +131 -131
@@ -1,80 +1,80 @@
1
- /**
2
- * Output Budget — Phase 7, Step 4: Trim large outputs to a token-friendly size.
3
- *
4
- * When gate outputs, agent tail outputs, or other large strings need to be
5
- * embedded in prompts (e.g., fix prompts, ledger entries), this module trims
6
- * them to a configurable budget while preserving the most useful parts
7
- * (beginning + end).
8
- *
9
- * Full output is optionally persisted to disk for debugging.
10
- */
11
-
12
- import { writeFileSync, mkdirSync } from 'node:fs';
13
- import { join, dirname } from 'node:path';
14
-
15
- // ── Constants ──────────────────────────────────────────────────────
16
-
17
- /** Default output budget in bytes */
18
- export const DEFAULT_OUTPUT_BUDGET = 4096;
19
-
20
- // ── Core ───────────────────────────────────────────────────────────
21
-
22
- /**
23
- * Trim text to a byte budget. If the text exceeds the budget, keep
24
- * the first half and last half with an omission marker in between.
25
- *
26
- * Returns the trimmed text. Never throws.
27
- */
28
- export function trimToBudget(text: string, budget: number = DEFAULT_OUTPUT_BUDGET): string {
29
- if (text.length <= budget) return text;
30
-
31
- const half = Math.floor(budget / 2) - 40; // Reserve space for marker
32
- if (half <= 0) return text.slice(0, budget);
33
-
34
- const omitted = text.length - half * 2;
35
- return (
36
- text.slice(0, half) +
37
- `\n\n... (${omitted} bytes omitted — full output saved to disk) ...\n\n` +
38
- text.slice(-half)
39
- );
40
- }
41
-
42
- /**
43
- * Persist full output to a file and return the trimmed version.
44
- *
45
- * If disk write fails, returns the trimmed text without file reference.
46
- * Never throws.
47
- */
48
- export function trimAndPersist(
49
- text: string,
50
- filePath: string,
51
- budget: number = DEFAULT_OUTPUT_BUDGET,
52
- ): { trimmed: string; persisted: boolean; fullPath?: string } {
53
- if (text.length <= budget) {
54
- return { trimmed: text, persisted: false };
55
- }
56
-
57
- // Persist full output to disk (best-effort)
58
- let persisted = false;
59
- try {
60
- mkdirSync(dirname(filePath), { recursive: true });
61
- writeFileSync(filePath, text, 'utf-8');
62
- persisted = true;
63
- } catch {
64
- // Disk write failed — degrade gracefully
65
- }
66
-
67
- const half = Math.floor(budget / 2) - 60; // Reserve space for marker + path
68
- if (half <= 0) {
69
- return { trimmed: text.slice(0, budget), persisted, fullPath: persisted ? filePath : undefined };
70
- }
71
-
72
- const omitted = text.length - half * 2;
73
- const pathNote = persisted ? ` Full output at: ${filePath}` : '';
74
- const trimmed =
75
- text.slice(0, half) +
76
- `\n\n... (${omitted} bytes omitted.${pathNote}) ...\n\n` +
77
- text.slice(-half);
78
-
79
- return { trimmed, persisted, fullPath: persisted ? filePath : undefined };
80
- }
1
+ /**
2
+ * Output Budget — Phase 7, Step 4: Trim large outputs to a token-friendly size.
3
+ *
4
+ * When gate outputs, agent tail outputs, or other large strings need to be
5
+ * embedded in prompts (e.g., fix prompts, ledger entries), this module trims
6
+ * them to a configurable budget while preserving the most useful parts
7
+ * (beginning + end).
8
+ *
9
+ * Full output is optionally persisted to disk for debugging.
10
+ */
11
+
12
+ import { writeFileSync, mkdirSync } from 'node:fs';
13
+ import { join, dirname } from 'node:path';
14
+
15
+ // ── Constants ──────────────────────────────────────────────────────
16
+
17
+ /** Default output budget in bytes */
18
+ export const DEFAULT_OUTPUT_BUDGET = 4096;
19
+
20
+ // ── Core ───────────────────────────────────────────────────────────
21
+
22
+ /**
23
+ * Trim text to a byte budget. If the text exceeds the budget, keep
24
+ * the first half and last half with an omission marker in between.
25
+ *
26
+ * Returns the trimmed text. Never throws.
27
+ */
28
+ export function trimToBudget(text: string, budget: number = DEFAULT_OUTPUT_BUDGET): string {
29
+ if (text.length <= budget) return text;
30
+
31
+ const half = Math.floor(budget / 2) - 40; // Reserve space for marker
32
+ if (half <= 0) return text.slice(0, budget);
33
+
34
+ const omitted = text.length - half * 2;
35
+ return (
36
+ text.slice(0, half) +
37
+ `\n\n... (${omitted} bytes omitted — full output saved to disk) ...\n\n` +
38
+ text.slice(-half)
39
+ );
40
+ }
41
+
42
+ /**
43
+ * Persist full output to a file and return the trimmed version.
44
+ *
45
+ * If disk write fails, returns the trimmed text without file reference.
46
+ * Never throws.
47
+ */
48
+ export function trimAndPersist(
49
+ text: string,
50
+ filePath: string,
51
+ budget: number = DEFAULT_OUTPUT_BUDGET,
52
+ ): { trimmed: string; persisted: boolean; fullPath?: string } {
53
+ if (text.length <= budget) {
54
+ return { trimmed: text, persisted: false };
55
+ }
56
+
57
+ // Persist full output to disk (best-effort)
58
+ let persisted = false;
59
+ try {
60
+ mkdirSync(dirname(filePath), { recursive: true });
61
+ writeFileSync(filePath, text, 'utf-8');
62
+ persisted = true;
63
+ } catch {
64
+ // Disk write failed — degrade gracefully
65
+ }
66
+
67
+ const half = Math.floor(budget / 2) - 60; // Reserve space for marker + path
68
+ if (half <= 0) {
69
+ return { trimmed: text.slice(0, budget), persisted, fullPath: persisted ? filePath : undefined };
70
+ }
71
+
72
+ const omitted = text.length - half * 2;
73
+ const pathNote = persisted ? ` Full output at: ${filePath}` : '';
74
+ const trimmed =
75
+ text.slice(0, half) +
76
+ `\n\n... (${omitted} bytes omitted.${pathNote}) ...\n\n` +
77
+ text.slice(-half);
78
+
79
+ return { trimmed, persisted, fullPath: persisted ? filePath : undefined };
80
+ }
@@ -1,152 +1,152 @@
1
- /**
2
- * Permission — Phase 7, Step 10: Risk classification + monitoring for agent tools.
3
- *
4
- * Classifies tool calls by risk tier and records them for post-run audit.
5
- * Does NOT block or restrict agent tools (Non-Invasiveness principle).
6
- * The primary use case is monitoring and reporting:
7
- * - Which high-risk tools were used and how often
8
- * - Whether agents used filesystem-destructive or network operations
9
- * - Per-task risk summary in the evidence directory
10
- *
11
- * Three tiers:
12
- * - safe: read-only operations (Read, Grep, List, etc.)
13
- * - moderate: write operations (Edit, Write, Execute safe commands)
14
- * - dangerous: destructive/network/exec operations (Delete, Bash, Deploy, etc.)
15
- */
16
-
17
- // ── Types ──────────────────────────────────────────────────────────
18
-
19
- export type RiskTier = 'safe' | 'moderate' | 'dangerous';
20
-
21
- export interface ToolUsageRecord {
22
- tool: string;
23
- tier: RiskTier;
24
- count: number;
25
- firstSeen: number;
26
- lastSeen: number;
27
- }
28
-
29
- export interface TaskRiskProfile {
30
- taskId: string;
31
- /** Highest risk tier used by this task */
32
- maxTier: RiskTier;
33
- /** Total tool calls */
34
- totalCalls: number;
35
- /** Tool usage breakdown */
36
- tools: ToolUsageRecord[];
37
- }
38
-
39
- // ── Classification Registry ────────────────────────────────────────
40
-
41
- const SAFE_TOOLS = new Set([
42
- 'read_file', 'grep_search', 'find_by_name', 'list_dir', 'code_search',
43
- 'read_notebook', 'read_url_content', 'view_content_chunk',
44
- 'memorix_search', 'memorix_detail', 'memorix_search_reasoning',
45
- 'mcp_query', 'search_web',
46
- ]);
47
-
48
- const MODERATE_TOOLS = new Set([
49
- 'edit', 'multi_edit', 'write_to_file', 'edit_notebook',
50
- 'memorix_store', 'memorix_store_reasoning', 'memorix_resolve',
51
- 'memorix_session_start', 'memorix_handoff', 'memorix_poll',
52
- 'git_commit', 'git_add', 'git_checkout',
53
- ]);
54
-
55
- const DANGEROUS_TOOLS = new Set([
56
- 'run_command', 'bash', 'execute_command', 'terminal',
57
- 'delete_file', 'remove_directory',
58
- 'deploy_web_app', 'browser_navigate', 'browser_click',
59
- 'http_request', 'fetch_url',
60
- ]);
61
-
62
- // ── Core ───────────────────────────────────────────────────────────
63
-
64
- /**
65
- * Classify a tool name into a risk tier.
66
- * Unknown tools default to 'moderate' (conservative but not alarming).
67
- */
68
- export function classifyTool(toolName: string): RiskTier {
69
- const name = toolName.toLowerCase().replace(/[^a-z0-9_]/g, '_');
70
- if (SAFE_TOOLS.has(name)) return 'safe';
71
- if (DANGEROUS_TOOLS.has(name)) return 'dangerous';
72
- if (MODERATE_TOOLS.has(name)) return 'moderate';
73
- // Heuristic fallback: tools with 'write', 'delete', 'exec', 'run' in name
74
- if (/delete|remove|exec|run|deploy|bash|terminal/.test(name)) return 'dangerous';
75
- if (/write|edit|create|update|store|commit/.test(name)) return 'moderate';
76
- return 'moderate'; // Unknown → moderate (safe default)
77
- }
78
-
79
- /**
80
- * Task-level tool usage tracker. Create one per dispatched task.
81
- */
82
- export class TaskToolTracker {
83
- private tools = new Map<string, ToolUsageRecord>();
84
- readonly taskId: string;
85
-
86
- constructor(taskId: string) {
87
- this.taskId = taskId;
88
- }
89
-
90
- /** Record a tool call. Call this when an agent uses a tool. */
91
- record(toolName: string): void {
92
- const tier = classifyTool(toolName);
93
- const now = Date.now();
94
- const existing = this.tools.get(toolName);
95
-
96
- if (existing) {
97
- existing.count++;
98
- existing.lastSeen = now;
99
- } else {
100
- this.tools.set(toolName, {
101
- tool: toolName,
102
- tier,
103
- count: 1,
104
- firstSeen: now,
105
- lastSeen: now,
106
- });
107
- }
108
- }
109
-
110
- /** Get the risk profile for this task. */
111
- getProfile(): TaskRiskProfile {
112
- const tools = Array.from(this.tools.values());
113
- const totalCalls = tools.reduce((sum, t) => sum + t.count, 0);
114
-
115
- let maxTier: RiskTier = 'safe';
116
- for (const t of tools) {
117
- if (t.tier === 'dangerous') { maxTier = 'dangerous'; break; }
118
- if (t.tier === 'moderate') maxTier = 'moderate';
119
- }
120
-
121
- return {
122
- taskId: this.taskId,
123
- maxTier,
124
- totalCalls,
125
- tools,
126
- };
127
- }
128
-
129
- /** Format a human-readable risk summary. */
130
- formatSummary(): string {
131
- const profile = this.getProfile();
132
- const lines = [`Risk: ${profile.maxTier} (${profile.totalCalls} tool calls)`];
133
-
134
- const dangerous = profile.tools.filter(t => t.tier === 'dangerous');
135
- if (dangerous.length > 0) {
136
- lines.push(' Dangerous tools:');
137
- for (const t of dangerous) {
138
- lines.push(` - ${t.tool}: ${t.count}x`);
139
- }
140
- }
141
-
142
- return lines.join('\n');
143
- }
144
- }
145
-
146
- /**
147
- * Compare two risk tiers. Returns positive if a > b, negative if a < b, 0 if equal.
148
- */
149
- export function compareTiers(a: RiskTier, b: RiskTier): number {
150
- const order: Record<RiskTier, number> = { safe: 0, moderate: 1, dangerous: 2 };
151
- return order[a] - order[b];
152
- }
1
+ /**
2
+ * Permission — Phase 7, Step 10: Risk classification + monitoring for agent tools.
3
+ *
4
+ * Classifies tool calls by risk tier and records them for post-run audit.
5
+ * Does NOT block or restrict agent tools (Non-Invasiveness principle).
6
+ * The primary use case is monitoring and reporting:
7
+ * - Which high-risk tools were used and how often
8
+ * - Whether agents used filesystem-destructive or network operations
9
+ * - Per-task risk summary in the evidence directory
10
+ *
11
+ * Three tiers:
12
+ * - safe: read-only operations (Read, Grep, List, etc.)
13
+ * - moderate: write operations (Edit, Write, Execute safe commands)
14
+ * - dangerous: destructive/network/exec operations (Delete, Bash, Deploy, etc.)
15
+ */
16
+
17
+ // ── Types ──────────────────────────────────────────────────────────
18
+
19
+ export type RiskTier = 'safe' | 'moderate' | 'dangerous';
20
+
21
+ export interface ToolUsageRecord {
22
+ tool: string;
23
+ tier: RiskTier;
24
+ count: number;
25
+ firstSeen: number;
26
+ lastSeen: number;
27
+ }
28
+
29
+ export interface TaskRiskProfile {
30
+ taskId: string;
31
+ /** Highest risk tier used by this task */
32
+ maxTier: RiskTier;
33
+ /** Total tool calls */
34
+ totalCalls: number;
35
+ /** Tool usage breakdown */
36
+ tools: ToolUsageRecord[];
37
+ }
38
+
39
+ // ── Classification Registry ────────────────────────────────────────
40
+
41
+ const SAFE_TOOLS = new Set([
42
+ 'read_file', 'grep_search', 'find_by_name', 'list_dir', 'code_search',
43
+ 'read_notebook', 'read_url_content', 'view_content_chunk',
44
+ 'memorix_search', 'memorix_detail', 'memorix_search_reasoning',
45
+ 'mcp_query', 'search_web',
46
+ ]);
47
+
48
+ const MODERATE_TOOLS = new Set([
49
+ 'edit', 'multi_edit', 'write_to_file', 'edit_notebook',
50
+ 'memorix_store', 'memorix_store_reasoning', 'memorix_resolve',
51
+ 'memorix_session_start', 'memorix_handoff', 'memorix_poll',
52
+ 'git_commit', 'git_add', 'git_checkout',
53
+ ]);
54
+
55
+ const DANGEROUS_TOOLS = new Set([
56
+ 'run_command', 'bash', 'execute_command', 'terminal',
57
+ 'delete_file', 'remove_directory',
58
+ 'deploy_web_app', 'browser_navigate', 'browser_click',
59
+ 'http_request', 'fetch_url',
60
+ ]);
61
+
62
+ // ── Core ───────────────────────────────────────────────────────────
63
+
64
+ /**
65
+ * Classify a tool name into a risk tier.
66
+ * Unknown tools default to 'moderate' (conservative but not alarming).
67
+ */
68
+ export function classifyTool(toolName: string): RiskTier {
69
+ const name = toolName.toLowerCase().replace(/[^a-z0-9_]/g, '_');
70
+ if (SAFE_TOOLS.has(name)) return 'safe';
71
+ if (DANGEROUS_TOOLS.has(name)) return 'dangerous';
72
+ if (MODERATE_TOOLS.has(name)) return 'moderate';
73
+ // Heuristic fallback: tools with 'write', 'delete', 'exec', 'run' in name
74
+ if (/delete|remove|exec|run|deploy|bash|terminal/.test(name)) return 'dangerous';
75
+ if (/write|edit|create|update|store|commit/.test(name)) return 'moderate';
76
+ return 'moderate'; // Unknown → moderate (safe default)
77
+ }
78
+
79
+ /**
80
+ * Task-level tool usage tracker. Create one per dispatched task.
81
+ */
82
+ export class TaskToolTracker {
83
+ private tools = new Map<string, ToolUsageRecord>();
84
+ readonly taskId: string;
85
+
86
+ constructor(taskId: string) {
87
+ this.taskId = taskId;
88
+ }
89
+
90
+ /** Record a tool call. Call this when an agent uses a tool. */
91
+ record(toolName: string): void {
92
+ const tier = classifyTool(toolName);
93
+ const now = Date.now();
94
+ const existing = this.tools.get(toolName);
95
+
96
+ if (existing) {
97
+ existing.count++;
98
+ existing.lastSeen = now;
99
+ } else {
100
+ this.tools.set(toolName, {
101
+ tool: toolName,
102
+ tier,
103
+ count: 1,
104
+ firstSeen: now,
105
+ lastSeen: now,
106
+ });
107
+ }
108
+ }
109
+
110
+ /** Get the risk profile for this task. */
111
+ getProfile(): TaskRiskProfile {
112
+ const tools = Array.from(this.tools.values());
113
+ const totalCalls = tools.reduce((sum, t) => sum + t.count, 0);
114
+
115
+ let maxTier: RiskTier = 'safe';
116
+ for (const t of tools) {
117
+ if (t.tier === 'dangerous') { maxTier = 'dangerous'; break; }
118
+ if (t.tier === 'moderate') maxTier = 'moderate';
119
+ }
120
+
121
+ return {
122
+ taskId: this.taskId,
123
+ maxTier,
124
+ totalCalls,
125
+ tools,
126
+ };
127
+ }
128
+
129
+ /** Format a human-readable risk summary. */
130
+ formatSummary(): string {
131
+ const profile = this.getProfile();
132
+ const lines = [`Risk: ${profile.maxTier} (${profile.totalCalls} tool calls)`];
133
+
134
+ const dangerous = profile.tools.filter(t => t.tier === 'dangerous');
135
+ if (dangerous.length > 0) {
136
+ lines.push(' Dangerous tools:');
137
+ for (const t of dangerous) {
138
+ lines.push(` - ${t.tool}: ${t.count}x`);
139
+ }
140
+ }
141
+
142
+ return lines.join('\n');
143
+ }
144
+ }
145
+
146
+ /**
147
+ * Compare two risk tiers. Returns positive if a > b, negative if a < b, 0 if equal.
148
+ */
149
+ export function compareTiers(a: RiskTier, b: RiskTier): number {
150
+ const order: Record<RiskTier, number> = { safe: 0, moderate: 1, dangerous: 2 };
151
+ return order[a] - order[b];
152
+ }