memorix 1.2.1 → 1.2.2

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (199) hide show
  1. package/CHANGELOG.md +16 -0
  2. package/README.md +14 -2
  3. package/README.zh-CN.md +14 -2
  4. package/TEAM.md +86 -86
  5. package/dist/cli/index.js +15407 -13779
  6. package/dist/cli/index.js.map +1 -1
  7. package/dist/index.js +1321 -529
  8. package/dist/index.js.map +1 -1
  9. package/dist/maintenance-runner.d.ts +1 -1
  10. package/dist/maintenance-runner.js +8458 -8087
  11. package/dist/maintenance-runner.js.map +1 -1
  12. package/dist/memcode-runtime/CHANGELOG.md +16 -0
  13. package/dist/sdk.d.ts +7 -2
  14. package/dist/sdk.js +1349 -535
  15. package/dist/sdk.js.map +1 -1
  16. package/dist/types.d.ts +49 -1
  17. package/dist/types.js.map +1 -1
  18. package/docs/1.2.2-MEMORY-CONTROL-PLANE.md +434 -0
  19. package/docs/AGENT_OPERATOR_PLAYBOOK.md +4 -0
  20. package/docs/API_REFERENCE.md +24 -4
  21. package/docs/DESIGN_DECISIONS.md +357 -357
  22. package/docs/README.md +1 -1
  23. package/docs/dev-log/progress.txt +91 -11
  24. package/package.json +1 -1
  25. package/plugins/codex/memorix/.codex-plugin/plugin.json +1 -1
  26. package/src/audit/index.ts +156 -156
  27. package/src/cli/command-guide.ts +192 -0
  28. package/src/cli/commands/audit-list.ts +89 -89
  29. package/src/cli/commands/audit.ts +9 -4
  30. package/src/cli/commands/background.ts +659 -659
  31. package/src/cli/commands/cleanup.ts +5 -1
  32. package/src/cli/commands/codegraph.ts +15 -5
  33. package/src/cli/commands/context.ts +3 -2
  34. package/src/cli/commands/doctor.ts +4 -2
  35. package/src/cli/commands/explain.ts +9 -3
  36. package/src/cli/commands/formation.ts +48 -48
  37. package/src/cli/commands/git-hook-install.ts +111 -111
  38. package/src/cli/commands/handoff.ts +75 -61
  39. package/src/cli/commands/hooks-status.ts +63 -63
  40. package/src/cli/commands/identity.ts +116 -0
  41. package/src/cli/commands/ingest-commit.ts +153 -153
  42. package/src/cli/commands/ingest-image.ts +71 -69
  43. package/src/cli/commands/ingest-log.ts +180 -180
  44. package/src/cli/commands/ingest.ts +44 -44
  45. package/src/cli/commands/integrate-shared.ts +15 -15
  46. package/src/cli/commands/lock.ts +93 -92
  47. package/src/cli/commands/memory.ts +58 -21
  48. package/src/cli/commands/message.ts +123 -118
  49. package/src/cli/commands/operator-shared.ts +98 -3
  50. package/src/cli/commands/poll.ts +74 -64
  51. package/src/cli/commands/purge-all-memory.ts +85 -85
  52. package/src/cli/commands/purge-project-memory.ts +83 -83
  53. package/src/cli/commands/reasoning.ts +135 -121
  54. package/src/cli/commands/retention.ts +9 -4
  55. package/src/cli/commands/serve-http.ts +8 -2
  56. package/src/cli/commands/serve-shared.ts +118 -118
  57. package/src/cli/commands/session.ts +29 -3
  58. package/src/cli/commands/skills.ts +124 -119
  59. package/src/cli/commands/status.ts +4 -3
  60. package/src/cli/commands/task.ts +193 -184
  61. package/src/cli/commands/team.ts +14 -10
  62. package/src/cli/commands/transfer.ts +108 -55
  63. package/src/cli/commands/uninstall-project-artifacts.ts +85 -85
  64. package/src/cli/identity.ts +89 -0
  65. package/src/cli/index.ts +96 -19
  66. package/src/cli/invocation.ts +115 -0
  67. package/src/cli/tui/ChatView.tsx +234 -234
  68. package/src/cli/tui/CommandBar.tsx +312 -312
  69. package/src/cli/tui/ContextRail.tsx +118 -118
  70. package/src/cli/tui/HeaderBar.tsx +72 -72
  71. package/src/cli/tui/LogoBanner.tsx +51 -51
  72. package/src/cli/tui/Sidebar.tsx +179 -179
  73. package/src/cli/tui/chat-service.ts +41 -18
  74. package/src/cli/tui/data.ts +23 -44
  75. package/src/cli/tui/index.ts +41 -41
  76. package/src/cli/tui/markdown-render.tsx +371 -371
  77. package/src/cli/tui/operator-context.ts +60 -0
  78. package/src/cli/tui/use-mouse.ts +157 -157
  79. package/src/cli/tui/useNavigation.ts +56 -56
  80. package/src/cli/tui/views/MemoryView.tsx +10 -8
  81. package/src/cli/update-checker.ts +211 -211
  82. package/src/cli/version.ts +7 -7
  83. package/src/cli/workbench.ts +1 -1
  84. package/src/codegraph/auto-context.ts +31 -2
  85. package/src/codegraph/context-pack.ts +1 -0
  86. package/src/codegraph/project-context.ts +2 -0
  87. package/src/compact/engine.ts +26 -10
  88. package/src/compact/index-format.ts +25 -2
  89. package/src/compact/token-budget.ts +74 -74
  90. package/src/dashboard/project-classification.ts +64 -64
  91. package/src/dashboard/server.ts +46 -9
  92. package/src/embedding/fastembed-provider.ts +142 -142
  93. package/src/embedding/transformers-provider.ts +111 -111
  94. package/src/git/extractor.ts +209 -209
  95. package/src/git/hooks-path.ts +85 -85
  96. package/src/hooks/admission.ts +117 -0
  97. package/src/hooks/handler.ts +98 -91
  98. package/src/hooks/pattern-detector.ts +173 -173
  99. package/src/hooks/significance-filter.ts +250 -250
  100. package/src/knowledge/context-assembly.ts +97 -0
  101. package/src/knowledge/workset.ts +179 -10
  102. package/src/llm/memory-manager.ts +328 -328
  103. package/src/llm/provider.ts +885 -885
  104. package/src/llm/quality.ts +248 -248
  105. package/src/memory/admission.ts +57 -0
  106. package/src/memory/attribution-guard.ts +249 -249
  107. package/src/memory/consolidation.ts +13 -2
  108. package/src/memory/disclosure-policy.ts +140 -135
  109. package/src/memory/entity-extractor.ts +197 -197
  110. package/src/memory/export-import.ts +11 -3
  111. package/src/memory/formation/evaluate.ts +217 -217
  112. package/src/memory/formation/extract.ts +361 -361
  113. package/src/memory/formation/index.ts +417 -417
  114. package/src/memory/formation/resolve.ts +344 -344
  115. package/src/memory/formation/types.ts +315 -315
  116. package/src/memory/freshness.ts +122 -122
  117. package/src/memory/graph-context.ts +8 -2
  118. package/src/memory/graph.ts +197 -197
  119. package/src/memory/observations.ts +162 -4
  120. package/src/memory/quality-audit.ts +2 -0
  121. package/src/memory/refs.ts +94 -94
  122. package/src/memory/retention.ts +22 -2
  123. package/src/memory/secret-filter.ts +79 -79
  124. package/src/memory/session.ts +5 -2
  125. package/src/memory/visibility.ts +80 -0
  126. package/src/multimodal/image-loader.ts +143 -143
  127. package/src/orchestrate/adapters/claude-stream.ts +192 -192
  128. package/src/orchestrate/adapters/claude.ts +111 -111
  129. package/src/orchestrate/adapters/codex-stream.ts +134 -134
  130. package/src/orchestrate/adapters/codex.ts +41 -41
  131. package/src/orchestrate/adapters/gemini-stream.ts +166 -166
  132. package/src/orchestrate/adapters/gemini.ts +42 -42
  133. package/src/orchestrate/adapters/index.ts +73 -73
  134. package/src/orchestrate/adapters/opencode-stream.ts +143 -143
  135. package/src/orchestrate/adapters/opencode.ts +47 -47
  136. package/src/orchestrate/adapters/spawn-helper.ts +286 -286
  137. package/src/orchestrate/adapters/types.ts +77 -77
  138. package/src/orchestrate/capability-router.ts +284 -284
  139. package/src/orchestrate/context-compact.ts +188 -188
  140. package/src/orchestrate/cost-tracker.ts +219 -219
  141. package/src/orchestrate/error-recovery.ts +191 -191
  142. package/src/orchestrate/evidence.ts +140 -140
  143. package/src/orchestrate/ledger.ts +110 -110
  144. package/src/orchestrate/memorix-bridge.ts +378 -340
  145. package/src/orchestrate/output-budget.ts +80 -80
  146. package/src/orchestrate/permission.ts +152 -152
  147. package/src/orchestrate/pipeline-trace.ts +131 -131
  148. package/src/orchestrate/prompt-builder.ts +155 -155
  149. package/src/orchestrate/ring-buffer.ts +37 -37
  150. package/src/orchestrate/task-graph.ts +389 -389
  151. package/src/orchestrate/worktree.ts +232 -232
  152. package/src/project/aliases.ts +374 -374
  153. package/src/project/detector.ts +268 -268
  154. package/src/rules/adapters/claude-code.ts +99 -99
  155. package/src/rules/adapters/codex.ts +97 -97
  156. package/src/rules/adapters/copilot.ts +124 -124
  157. package/src/rules/adapters/cursor.ts +114 -114
  158. package/src/rules/adapters/kiro.ts +126 -126
  159. package/src/rules/adapters/trae.ts +56 -56
  160. package/src/rules/adapters/windsurf.ts +83 -83
  161. package/src/rules/syncer.ts +235 -235
  162. package/src/runtime/control-plane-maintenance.ts +1 -0
  163. package/src/runtime/isolated-maintenance.ts +1 -0
  164. package/src/runtime/lifecycle.ts +18 -0
  165. package/src/runtime/maintenance-jobs.ts +1 -0
  166. package/src/runtime/maintenance-runner.ts +2 -0
  167. package/src/runtime/project-maintenance.ts +89 -0
  168. package/src/sdk.ts +334 -304
  169. package/src/search/intent-detector.ts +289 -289
  170. package/src/search/query-expansion.ts +52 -52
  171. package/src/server/formation-timeout.ts +27 -27
  172. package/src/server.ts +260 -81
  173. package/src/skills/mini-skills.ts +386 -386
  174. package/src/store/chat-store.ts +119 -119
  175. package/src/store/graph-store.ts +249 -249
  176. package/src/store/mini-skill-store.ts +349 -349
  177. package/src/store/orama-store.ts +61 -6
  178. package/src/store/persistence-json.ts +212 -212
  179. package/src/store/persistence.ts +291 -291
  180. package/src/store/project-affinity.ts +195 -195
  181. package/src/store/sqlite-db.ts +23 -1
  182. package/src/store/sqlite-store.ts +12 -2
  183. package/src/team/event-bus.ts +76 -76
  184. package/src/team/file-locks.ts +173 -173
  185. package/src/team/handoff.ts +168 -161
  186. package/src/team/messages.ts +203 -203
  187. package/src/team/poll.ts +132 -132
  188. package/src/team/tasks.ts +211 -211
  189. package/src/types.ts +51 -0
  190. package/src/wiki/generator.ts +2 -0
  191. package/src/workspace/mcp-adapters/codex.ts +191 -191
  192. package/src/workspace/mcp-adapters/copilot.ts +105 -105
  193. package/src/workspace/mcp-adapters/cursor.ts +53 -53
  194. package/src/workspace/mcp-adapters/kiro.ts +64 -64
  195. package/src/workspace/mcp-adapters/opencode.ts +123 -123
  196. package/src/workspace/mcp-adapters/trae.ts +134 -134
  197. package/src/workspace/mcp-adapters/windsurf.ts +91 -91
  198. package/src/workspace/sanitizer.ts +60 -60
  199. package/src/workspace/workflow-sync.ts +131 -131
@@ -1,219 +1,219 @@
1
- /**
2
- * Cost Tracker — Phase 7, Step 8: Token → USD conversion and budget enforcement.
3
- *
4
- * Converts accumulated token usage (per model) into USD using a configurable
5
- * price table. Supports budget limits that abort the pipeline when exceeded.
6
- *
7
- * Design principle: if the price table doesn't have a model, report tokens only
8
- * (skip USD calculation). Never crash because of a missing price entry.
9
- */
10
-
11
- import type { TokenUsage } from './adapters/types.js';
12
-
13
- // ── Types ──────────────────────────────────────────────────────────
14
-
15
- export interface ModelPrice {
16
- /** Cost per 1M input tokens in USD */
17
- inputPer1M: number;
18
- /** Cost per 1M output tokens in USD */
19
- outputPer1M: number;
20
- /** Cost per 1M cache read tokens in USD (default: 0) */
21
- cacheReadPer1M?: number;
22
- /** Cost per 1M cache write tokens in USD (default: 0) */
23
- cacheWritePer1M?: number;
24
- }
25
-
26
- export interface CostSummary {
27
- /** Total cost in USD (null if no price data available) */
28
- totalUSD: number | null;
29
- /** Per-model breakdown */
30
- models: {
31
- model: string;
32
- inputTokens: number;
33
- outputTokens: number;
34
- cacheReadTokens: number;
35
- cacheWriteTokens: number;
36
- costUSD: number | null;
37
- }[];
38
- /** Whether the budget has been exceeded */
39
- budgetExceeded: boolean;
40
- /** Configured budget (null if no budget) */
41
- budgetUSD: number | null;
42
- }
43
-
44
- // ── Default Price Table ────────────────────────────────────────────
45
-
46
- const DEFAULT_PRICES: Record<string, ModelPrice> = {
47
- // Claude models — verified 2026-04-13 via Anthropic docs
48
- 'claude-opus-4-6': { inputPer1M: 5, outputPer1M: 25, cacheReadPer1M: 0.50, cacheWritePer1M: 6.25 },
49
- 'claude-sonnet-4-6': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
50
- 'claude-haiku-4-5': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
51
- 'claude-haiku-4-5-20251001': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
52
- // Legacy Claude models (still in use by some adapters)
53
- 'claude-sonnet-4-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
54
- 'claude-sonnet-4-0-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
55
- 'claude-opus-4-20250514': { inputPer1M: 15, outputPer1M: 75, cacheReadPer1M: 1.50, cacheWritePer1M: 18.75 },
56
- 'claude-3-5-sonnet': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
57
- 'claude-3-5-haiku': { inputPer1M: 0.80, outputPer1M: 4, cacheReadPer1M: 0.08, cacheWritePer1M: 1 },
58
-
59
- // OpenAI / Codex models — verified 2026-04-13
60
- 'gpt-5': { inputPer1M: 1.25, outputPer1M: 10 },
61
- 'gpt-5-mini': { inputPer1M: 0.25, outputPer1M: 2 },
62
- 'gpt-5-nano': { inputPer1M: 0.05, outputPer1M: 0.40 },
63
- 'gpt-5.2': { inputPer1M: 1.75, outputPer1M: 14 },
64
- 'o3': { inputPer1M: 2, outputPer1M: 8 },
65
- 'o3-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
66
- 'o4-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
67
- 'gpt-4.1': { inputPer1M: 2, outputPer1M: 8, cacheReadPer1M: 0.50, cacheWritePer1M: 2 },
68
- 'gpt-4.1-mini': { inputPer1M: 0.40, outputPer1M: 1.60, cacheReadPer1M: 0.10, cacheWritePer1M: 0.40 },
69
- 'gpt-4.1-nano': { inputPer1M: 0.10, outputPer1M: 0.40 },
70
- 'gpt-4o': { inputPer1M: 2.50, outputPer1M: 10, cacheReadPer1M: 1.25 },
71
- 'gpt-4o-mini': { inputPer1M: 0.15, outputPer1M: 0.60, cacheReadPer1M: 0.075 },
72
-
73
- // Google Gemini models — verified 2026-04-13 via Google AI pricing
74
- 'gemini-2.5-flash': { inputPer1M: 0.30, outputPer1M: 2.50 },
75
- 'gemini-2.5-flash-preview-05-20': { inputPer1M: 0.30, outputPer1M: 2.50 },
76
- 'gemini-2.5-flash-lite': { inputPer1M: 0.10, outputPer1M: 0.40 },
77
- 'gemini-2.5-pro': { inputPer1M: 1.25, outputPer1M: 10 },
78
- 'gemini-2.0-flash': { inputPer1M: 0.10, outputPer1M: 0.40 },
79
- 'gemini-2.0-flash-lite': { inputPer1M: 0.075, outputPer1M: 0.30 },
80
- // Gemini 3.x generation
81
- 'gemini-3.1-pro': { inputPer1M: 2, outputPer1M: 12 },
82
- 'gemini-3-flash': { inputPer1M: 0.50, outputPer1M: 3 },
83
-
84
- // OpenCode / open models — verified 2026-04-13
85
- 'deepseek-coder': { inputPer1M: 0.28, outputPer1M: 0.42 },
86
- 'deepseek-chat': { inputPer1M: 0.28, outputPer1M: 0.42 },
87
- };
88
-
89
- // ── Core ───────────────────────────────────────────────────────────
90
-
91
- /**
92
- * Calculate cost for a single model's token usage.
93
- * Returns null if the model is not in the price table.
94
- */
95
- export function calculateModelCost(
96
- model: string,
97
- usage: TokenUsage,
98
- customPrices?: Record<string, ModelPrice>,
99
- ): number | null {
100
- const prices = { ...DEFAULT_PRICES, ...customPrices };
101
- const price = findPrice(model, prices);
102
- if (!price) return null;
103
-
104
- const inputCost = (usage.inputTokens / 1_000_000) * price.inputPer1M;
105
- const outputCost = (usage.outputTokens / 1_000_000) * price.outputPer1M;
106
- const cacheReadCost = (usage.cacheReadTokens / 1_000_000) * (price.cacheReadPer1M ?? 0);
107
- const cacheWriteCost = (usage.cacheWriteTokens / 1_000_000) * (price.cacheWritePer1M ?? 0);
108
-
109
- return inputCost + outputCost + cacheReadCost + cacheWriteCost;
110
- }
111
-
112
- /**
113
- * Calculate total pipeline cost from accumulated token usage.
114
- */
115
- export function calculatePipelineCost(
116
- tokenUsage: Record<string, TokenUsage>,
117
- budgetUSD?: number,
118
- customPrices?: Record<string, ModelPrice>,
119
- ): CostSummary {
120
- const models: CostSummary['models'] = [];
121
- let totalUSD: number | null = 0;
122
- let hasAnyPrice = false;
123
-
124
- for (const [model, usage] of Object.entries(tokenUsage)) {
125
- const cost = calculateModelCost(model, usage, customPrices);
126
- if (cost !== null) {
127
- hasAnyPrice = true;
128
- totalUSD = (totalUSD ?? 0) + cost;
129
- }
130
-
131
- models.push({
132
- model,
133
- inputTokens: usage.inputTokens,
134
- outputTokens: usage.outputTokens,
135
- cacheReadTokens: usage.cacheReadTokens,
136
- cacheWriteTokens: usage.cacheWriteTokens,
137
- costUSD: cost !== null ? Math.round(cost * 10000) / 10000 : null, // Round to 4 decimals
138
- });
139
- }
140
-
141
- if (!hasAnyPrice) totalUSD = null;
142
- else totalUSD = Math.round(totalUSD! * 10000) / 10000;
143
-
144
- return {
145
- totalUSD,
146
- models,
147
- budgetExceeded: budgetUSD != null && totalUSD != null && totalUSD > budgetUSD,
148
- budgetUSD: budgetUSD ?? null,
149
- };
150
- }
151
-
152
- /**
153
- * Check if the current cost exceeds the budget.
154
- * Returns false if no budget is set or cost cannot be calculated.
155
- */
156
- export function isBudgetExceeded(
157
- tokenUsage: Record<string, TokenUsage>,
158
- budgetUSD?: number,
159
- customPrices?: Record<string, ModelPrice>,
160
- ): boolean {
161
- if (budgetUSD == null) return false;
162
- const summary = calculatePipelineCost(tokenUsage, budgetUSD, customPrices);
163
- return summary.budgetExceeded;
164
- }
165
-
166
- /**
167
- * Format cost summary for CLI display.
168
- */
169
- export function formatCostSummary(summary: CostSummary): string {
170
- const lines: string[] = [];
171
-
172
- for (const m of summary.models) {
173
- const tokens = `in=${fmtNum(m.inputTokens)} out=${fmtNum(m.outputTokens)}`;
174
- const cache = m.cacheReadTokens > 0 || m.cacheWriteTokens > 0
175
- ? ` cache_r=${fmtNum(m.cacheReadTokens)} cache_w=${fmtNum(m.cacheWriteTokens)}`
176
- : '';
177
- const cost = m.costUSD !== null ? ` ($${m.costUSD.toFixed(4)})` : '';
178
- lines.push(` ${m.model}: ${tokens}${cache}${cost}`);
179
- }
180
-
181
- if (summary.totalUSD !== null) {
182
- lines.push(` Total: $${summary.totalUSD.toFixed(4)}`);
183
- }
184
-
185
- if (summary.budgetUSD !== null) {
186
- const status = summary.budgetExceeded ? ' [WARN] EXCEEDED' : '';
187
- lines.push(` Budget: $${summary.budgetUSD.toFixed(2)}${status}`);
188
- }
189
-
190
- return lines.join('\n');
191
- }
192
-
193
- // ── Helpers ────────────────────────────────────────────────────────
194
-
195
- /**
196
- * Find price entry by model name. Supports fuzzy matching:
197
- * "claude-sonnet-4-20250514" matches "claude-sonnet-4-20250514" exactly,
198
- * but also tries prefix matching for versioned models.
199
- */
200
- function findPrice(model: string, prices: Record<string, ModelPrice>): ModelPrice | null {
201
- // Exact match
202
- if (prices[model]) return prices[model];
203
-
204
- // Prefix match: "claude-sonnet-4-20250514" → try "claude-sonnet-4"
205
- const normalized = model.toLowerCase();
206
- for (const [key, price] of Object.entries(prices)) {
207
- if (normalized.startsWith(key.toLowerCase()) || key.toLowerCase().startsWith(normalized)) {
208
- return price;
209
- }
210
- }
211
-
212
- return null;
213
- }
214
-
215
- function fmtNum(n: number): string {
216
- if (n >= 1_000_000) return `${(n / 1_000_000).toFixed(1)}M`;
217
- if (n >= 1_000) return `${(n / 1_000).toFixed(1)}K`;
218
- return String(n);
219
- }
1
+ /**
2
+ * Cost Tracker — Phase 7, Step 8: Token → USD conversion and budget enforcement.
3
+ *
4
+ * Converts accumulated token usage (per model) into USD using a configurable
5
+ * price table. Supports budget limits that abort the pipeline when exceeded.
6
+ *
7
+ * Design principle: if the price table doesn't have a model, report tokens only
8
+ * (skip USD calculation). Never crash because of a missing price entry.
9
+ */
10
+
11
+ import type { TokenUsage } from './adapters/types.js';
12
+
13
+ // ── Types ──────────────────────────────────────────────────────────
14
+
15
+ export interface ModelPrice {
16
+ /** Cost per 1M input tokens in USD */
17
+ inputPer1M: number;
18
+ /** Cost per 1M output tokens in USD */
19
+ outputPer1M: number;
20
+ /** Cost per 1M cache read tokens in USD (default: 0) */
21
+ cacheReadPer1M?: number;
22
+ /** Cost per 1M cache write tokens in USD (default: 0) */
23
+ cacheWritePer1M?: number;
24
+ }
25
+
26
+ export interface CostSummary {
27
+ /** Total cost in USD (null if no price data available) */
28
+ totalUSD: number | null;
29
+ /** Per-model breakdown */
30
+ models: {
31
+ model: string;
32
+ inputTokens: number;
33
+ outputTokens: number;
34
+ cacheReadTokens: number;
35
+ cacheWriteTokens: number;
36
+ costUSD: number | null;
37
+ }[];
38
+ /** Whether the budget has been exceeded */
39
+ budgetExceeded: boolean;
40
+ /** Configured budget (null if no budget) */
41
+ budgetUSD: number | null;
42
+ }
43
+
44
+ // ── Default Price Table ────────────────────────────────────────────
45
+
46
+ const DEFAULT_PRICES: Record<string, ModelPrice> = {
47
+ // Claude models — verified 2026-04-13 via Anthropic docs
48
+ 'claude-opus-4-6': { inputPer1M: 5, outputPer1M: 25, cacheReadPer1M: 0.50, cacheWritePer1M: 6.25 },
49
+ 'claude-sonnet-4-6': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
50
+ 'claude-haiku-4-5': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
51
+ 'claude-haiku-4-5-20251001': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
52
+ // Legacy Claude models (still in use by some adapters)
53
+ 'claude-sonnet-4-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
54
+ 'claude-sonnet-4-0-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
55
+ 'claude-opus-4-20250514': { inputPer1M: 15, outputPer1M: 75, cacheReadPer1M: 1.50, cacheWritePer1M: 18.75 },
56
+ 'claude-3-5-sonnet': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
57
+ 'claude-3-5-haiku': { inputPer1M: 0.80, outputPer1M: 4, cacheReadPer1M: 0.08, cacheWritePer1M: 1 },
58
+
59
+ // OpenAI / Codex models — verified 2026-04-13
60
+ 'gpt-5': { inputPer1M: 1.25, outputPer1M: 10 },
61
+ 'gpt-5-mini': { inputPer1M: 0.25, outputPer1M: 2 },
62
+ 'gpt-5-nano': { inputPer1M: 0.05, outputPer1M: 0.40 },
63
+ 'gpt-5.2': { inputPer1M: 1.75, outputPer1M: 14 },
64
+ 'o3': { inputPer1M: 2, outputPer1M: 8 },
65
+ 'o3-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
66
+ 'o4-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
67
+ 'gpt-4.1': { inputPer1M: 2, outputPer1M: 8, cacheReadPer1M: 0.50, cacheWritePer1M: 2 },
68
+ 'gpt-4.1-mini': { inputPer1M: 0.40, outputPer1M: 1.60, cacheReadPer1M: 0.10, cacheWritePer1M: 0.40 },
69
+ 'gpt-4.1-nano': { inputPer1M: 0.10, outputPer1M: 0.40 },
70
+ 'gpt-4o': { inputPer1M: 2.50, outputPer1M: 10, cacheReadPer1M: 1.25 },
71
+ 'gpt-4o-mini': { inputPer1M: 0.15, outputPer1M: 0.60, cacheReadPer1M: 0.075 },
72
+
73
+ // Google Gemini models — verified 2026-04-13 via Google AI pricing
74
+ 'gemini-2.5-flash': { inputPer1M: 0.30, outputPer1M: 2.50 },
75
+ 'gemini-2.5-flash-preview-05-20': { inputPer1M: 0.30, outputPer1M: 2.50 },
76
+ 'gemini-2.5-flash-lite': { inputPer1M: 0.10, outputPer1M: 0.40 },
77
+ 'gemini-2.5-pro': { inputPer1M: 1.25, outputPer1M: 10 },
78
+ 'gemini-2.0-flash': { inputPer1M: 0.10, outputPer1M: 0.40 },
79
+ 'gemini-2.0-flash-lite': { inputPer1M: 0.075, outputPer1M: 0.30 },
80
+ // Gemini 3.x generation
81
+ 'gemini-3.1-pro': { inputPer1M: 2, outputPer1M: 12 },
82
+ 'gemini-3-flash': { inputPer1M: 0.50, outputPer1M: 3 },
83
+
84
+ // OpenCode / open models — verified 2026-04-13
85
+ 'deepseek-coder': { inputPer1M: 0.28, outputPer1M: 0.42 },
86
+ 'deepseek-chat': { inputPer1M: 0.28, outputPer1M: 0.42 },
87
+ };
88
+
89
+ // ── Core ───────────────────────────────────────────────────────────
90
+
91
+ /**
92
+ * Calculate cost for a single model's token usage.
93
+ * Returns null if the model is not in the price table.
94
+ */
95
+ export function calculateModelCost(
96
+ model: string,
97
+ usage: TokenUsage,
98
+ customPrices?: Record<string, ModelPrice>,
99
+ ): number | null {
100
+ const prices = { ...DEFAULT_PRICES, ...customPrices };
101
+ const price = findPrice(model, prices);
102
+ if (!price) return null;
103
+
104
+ const inputCost = (usage.inputTokens / 1_000_000) * price.inputPer1M;
105
+ const outputCost = (usage.outputTokens / 1_000_000) * price.outputPer1M;
106
+ const cacheReadCost = (usage.cacheReadTokens / 1_000_000) * (price.cacheReadPer1M ?? 0);
107
+ const cacheWriteCost = (usage.cacheWriteTokens / 1_000_000) * (price.cacheWritePer1M ?? 0);
108
+
109
+ return inputCost + outputCost + cacheReadCost + cacheWriteCost;
110
+ }
111
+
112
+ /**
113
+ * Calculate total pipeline cost from accumulated token usage.
114
+ */
115
+ export function calculatePipelineCost(
116
+ tokenUsage: Record<string, TokenUsage>,
117
+ budgetUSD?: number,
118
+ customPrices?: Record<string, ModelPrice>,
119
+ ): CostSummary {
120
+ const models: CostSummary['models'] = [];
121
+ let totalUSD: number | null = 0;
122
+ let hasAnyPrice = false;
123
+
124
+ for (const [model, usage] of Object.entries(tokenUsage)) {
125
+ const cost = calculateModelCost(model, usage, customPrices);
126
+ if (cost !== null) {
127
+ hasAnyPrice = true;
128
+ totalUSD = (totalUSD ?? 0) + cost;
129
+ }
130
+
131
+ models.push({
132
+ model,
133
+ inputTokens: usage.inputTokens,
134
+ outputTokens: usage.outputTokens,
135
+ cacheReadTokens: usage.cacheReadTokens,
136
+ cacheWriteTokens: usage.cacheWriteTokens,
137
+ costUSD: cost !== null ? Math.round(cost * 10000) / 10000 : null, // Round to 4 decimals
138
+ });
139
+ }
140
+
141
+ if (!hasAnyPrice) totalUSD = null;
142
+ else totalUSD = Math.round(totalUSD! * 10000) / 10000;
143
+
144
+ return {
145
+ totalUSD,
146
+ models,
147
+ budgetExceeded: budgetUSD != null && totalUSD != null && totalUSD > budgetUSD,
148
+ budgetUSD: budgetUSD ?? null,
149
+ };
150
+ }
151
+
152
+ /**
153
+ * Check if the current cost exceeds the budget.
154
+ * Returns false if no budget is set or cost cannot be calculated.
155
+ */
156
+ export function isBudgetExceeded(
157
+ tokenUsage: Record<string, TokenUsage>,
158
+ budgetUSD?: number,
159
+ customPrices?: Record<string, ModelPrice>,
160
+ ): boolean {
161
+ if (budgetUSD == null) return false;
162
+ const summary = calculatePipelineCost(tokenUsage, budgetUSD, customPrices);
163
+ return summary.budgetExceeded;
164
+ }
165
+
166
+ /**
167
+ * Format cost summary for CLI display.
168
+ */
169
+ export function formatCostSummary(summary: CostSummary): string {
170
+ const lines: string[] = [];
171
+
172
+ for (const m of summary.models) {
173
+ const tokens = `in=${fmtNum(m.inputTokens)} out=${fmtNum(m.outputTokens)}`;
174
+ const cache = m.cacheReadTokens > 0 || m.cacheWriteTokens > 0
175
+ ? ` cache_r=${fmtNum(m.cacheReadTokens)} cache_w=${fmtNum(m.cacheWriteTokens)}`
176
+ : '';
177
+ const cost = m.costUSD !== null ? ` ($${m.costUSD.toFixed(4)})` : '';
178
+ lines.push(` ${m.model}: ${tokens}${cache}${cost}`);
179
+ }
180
+
181
+ if (summary.totalUSD !== null) {
182
+ lines.push(` Total: $${summary.totalUSD.toFixed(4)}`);
183
+ }
184
+
185
+ if (summary.budgetUSD !== null) {
186
+ const status = summary.budgetExceeded ? ' [WARN] EXCEEDED' : '';
187
+ lines.push(` Budget: $${summary.budgetUSD.toFixed(2)}${status}`);
188
+ }
189
+
190
+ return lines.join('\n');
191
+ }
192
+
193
+ // ── Helpers ────────────────────────────────────────────────────────
194
+
195
+ /**
196
+ * Find price entry by model name. Supports fuzzy matching:
197
+ * "claude-sonnet-4-20250514" matches "claude-sonnet-4-20250514" exactly,
198
+ * but also tries prefix matching for versioned models.
199
+ */
200
+ function findPrice(model: string, prices: Record<string, ModelPrice>): ModelPrice | null {
201
+ // Exact match
202
+ if (prices[model]) return prices[model];
203
+
204
+ // Prefix match: "claude-sonnet-4-20250514" → try "claude-sonnet-4"
205
+ const normalized = model.toLowerCase();
206
+ for (const [key, price] of Object.entries(prices)) {
207
+ if (normalized.startsWith(key.toLowerCase()) || key.toLowerCase().startsWith(normalized)) {
208
+ return price;
209
+ }
210
+ }
211
+
212
+ return null;
213
+ }
214
+
215
+ function fmtNum(n: number): string {
216
+ if (n >= 1_000_000) return `${(n / 1_000_000).toFixed(1)}M`;
217
+ if (n >= 1_000) return `${(n / 1_000).toFixed(1)}K`;
218
+ return String(n);
219
+ }