memorix 1.2.2 → 1.2.4

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (163) hide show
  1. package/CHANGELOG.md +27 -0
  2. package/README.md +3 -3
  3. package/README.zh-CN.md +3 -3
  4. package/TEAM.md +86 -86
  5. package/dist/cli/index.js +5199 -4726
  6. package/dist/cli/index.js.map +1 -1
  7. package/dist/index.js +428 -49
  8. package/dist/index.js.map +1 -1
  9. package/dist/maintenance-runner.js +97 -18
  10. package/dist/maintenance-runner.js.map +1 -1
  11. package/dist/memcode-runtime/CHANGELOG.md +27 -0
  12. package/dist/sdk.js +428 -49
  13. package/dist/sdk.js.map +1 -1
  14. package/docs/1.2.4-PERSISTENT-MEMORY-DELIVERY.md +86 -0
  15. package/docs/AGENT_OPERATOR_PLAYBOOK.md +13 -1
  16. package/docs/API_REFERENCE.md +13 -3
  17. package/docs/DESIGN_DECISIONS.md +357 -357
  18. package/docs/dev-log/progress.txt +60 -9
  19. package/package.json +1 -1
  20. package/plugins/codex/memorix/.codex-plugin/plugin.json +1 -1
  21. package/src/audit/index.ts +156 -156
  22. package/src/cli/capability-map.ts +1 -1
  23. package/src/cli/command-guide.ts +4 -1
  24. package/src/cli/commands/agent-integrations.ts +5 -1
  25. package/src/cli/commands/audit-list.ts +89 -89
  26. package/src/cli/commands/background.ts +659 -659
  27. package/src/cli/commands/codegraph.ts +1 -1
  28. package/src/cli/commands/context.ts +9 -1
  29. package/src/cli/commands/formation.ts +48 -48
  30. package/src/cli/commands/git-hook-install.ts +111 -111
  31. package/src/cli/commands/handoff.ts +54 -54
  32. package/src/cli/commands/hooks-status.ts +63 -63
  33. package/src/cli/commands/ingest-commit.ts +153 -153
  34. package/src/cli/commands/ingest-image.ts +66 -66
  35. package/src/cli/commands/ingest-log.ts +180 -180
  36. package/src/cli/commands/ingest.ts +44 -44
  37. package/src/cli/commands/integrate-shared.ts +15 -15
  38. package/src/cli/commands/lock.ts +82 -82
  39. package/src/cli/commands/message.ts +104 -104
  40. package/src/cli/commands/poll.ts +58 -58
  41. package/src/cli/commands/purge-all-memory.ts +85 -85
  42. package/src/cli/commands/purge-project-memory.ts +83 -83
  43. package/src/cli/commands/reasoning.ts +118 -118
  44. package/src/cli/commands/resume.ts +31 -0
  45. package/src/cli/commands/serve-shared.ts +118 -118
  46. package/src/cli/commands/session.ts +15 -7
  47. package/src/cli/commands/skills.ts +114 -114
  48. package/src/cli/commands/task.ts +167 -167
  49. package/src/cli/commands/transfer.ts +47 -47
  50. package/src/cli/commands/uninstall-project-artifacts.ts +85 -85
  51. package/src/cli/index.ts +3 -1
  52. package/src/cli/tui/ChatView.tsx +234 -234
  53. package/src/cli/tui/CommandBar.tsx +312 -312
  54. package/src/cli/tui/ContextRail.tsx +118 -118
  55. package/src/cli/tui/HeaderBar.tsx +72 -72
  56. package/src/cli/tui/LogoBanner.tsx +51 -51
  57. package/src/cli/tui/Sidebar.tsx +179 -179
  58. package/src/cli/tui/index.ts +41 -41
  59. package/src/cli/tui/markdown-render.tsx +371 -371
  60. package/src/cli/tui/session-service.ts +3 -2
  61. package/src/cli/tui/use-mouse.ts +157 -157
  62. package/src/cli/tui/useNavigation.ts +56 -56
  63. package/src/cli/update-checker.ts +211 -211
  64. package/src/cli/version.ts +7 -7
  65. package/src/cli/workbench.ts +1 -1
  66. package/src/codegraph/auto-context.ts +54 -1
  67. package/src/codegraph/task-lens.ts +29 -0
  68. package/src/compact/token-budget.ts +89 -74
  69. package/src/config/toml-loader.ts +9 -5
  70. package/src/dashboard/project-classification.ts +64 -64
  71. package/src/embedding/fastembed-provider.ts +142 -142
  72. package/src/embedding/transformers-provider.ts +111 -111
  73. package/src/git/extractor.ts +209 -209
  74. package/src/git/hooks-path.ts +85 -85
  75. package/src/hooks/handler.ts +127 -66
  76. package/src/hooks/installers/index.ts +5 -4
  77. package/src/hooks/official-skills.ts +6 -4
  78. package/src/hooks/pattern-detector.ts +173 -173
  79. package/src/hooks/rules/memorix-agent-rules.md +9 -7
  80. package/src/hooks/significance-filter.ts +250 -250
  81. package/src/knowledge/context-assembly.ts +4 -1
  82. package/src/knowledge/workset.ts +89 -1
  83. package/src/llm/memory-manager.ts +328 -328
  84. package/src/llm/provider.ts +885 -885
  85. package/src/llm/quality.ts +248 -248
  86. package/src/memory/attribution-guard.ts +249 -249
  87. package/src/memory/disclosure-policy.ts +135 -135
  88. package/src/memory/entity-extractor.ts +197 -197
  89. package/src/memory/formation/evaluate.ts +217 -217
  90. package/src/memory/formation/extract.ts +361 -361
  91. package/src/memory/formation/index.ts +417 -417
  92. package/src/memory/formation/resolve.ts +344 -344
  93. package/src/memory/formation/types.ts +315 -315
  94. package/src/memory/freshness.ts +122 -122
  95. package/src/memory/graph.ts +197 -197
  96. package/src/memory/refs.ts +94 -94
  97. package/src/memory/secret-filter.ts +79 -79
  98. package/src/memory/session.ts +158 -9
  99. package/src/multimodal/image-loader.ts +143 -143
  100. package/src/orchestrate/adapters/claude-stream.ts +192 -192
  101. package/src/orchestrate/adapters/claude.ts +111 -111
  102. package/src/orchestrate/adapters/codex-stream.ts +134 -134
  103. package/src/orchestrate/adapters/codex.ts +41 -41
  104. package/src/orchestrate/adapters/gemini-stream.ts +166 -166
  105. package/src/orchestrate/adapters/gemini.ts +42 -42
  106. package/src/orchestrate/adapters/index.ts +73 -73
  107. package/src/orchestrate/adapters/opencode-stream.ts +143 -143
  108. package/src/orchestrate/adapters/opencode.ts +47 -47
  109. package/src/orchestrate/adapters/spawn-helper.ts +286 -286
  110. package/src/orchestrate/adapters/types.ts +77 -77
  111. package/src/orchestrate/capability-router.ts +284 -284
  112. package/src/orchestrate/context-compact.ts +188 -188
  113. package/src/orchestrate/cost-tracker.ts +219 -219
  114. package/src/orchestrate/error-recovery.ts +191 -191
  115. package/src/orchestrate/evidence.ts +140 -140
  116. package/src/orchestrate/ledger.ts +110 -110
  117. package/src/orchestrate/memorix-bridge.ts +343 -343
  118. package/src/orchestrate/output-budget.ts +80 -80
  119. package/src/orchestrate/permission.ts +152 -152
  120. package/src/orchestrate/pipeline-trace.ts +131 -131
  121. package/src/orchestrate/prompt-builder.ts +155 -155
  122. package/src/orchestrate/ring-buffer.ts +37 -37
  123. package/src/orchestrate/task-graph.ts +389 -389
  124. package/src/orchestrate/worktree.ts +232 -232
  125. package/src/project/aliases.ts +374 -374
  126. package/src/project/detector.ts +268 -268
  127. package/src/rules/adapters/claude-code.ts +99 -99
  128. package/src/rules/adapters/codex.ts +97 -97
  129. package/src/rules/adapters/copilot.ts +124 -124
  130. package/src/rules/adapters/cursor.ts +114 -114
  131. package/src/rules/adapters/kiro.ts +126 -126
  132. package/src/rules/adapters/trae.ts +56 -56
  133. package/src/rules/adapters/windsurf.ts +83 -83
  134. package/src/rules/syncer.ts +235 -235
  135. package/src/sdk.ts +299 -299
  136. package/src/search/intent-detector.ts +289 -289
  137. package/src/search/query-expansion.ts +52 -52
  138. package/src/server/formation-timeout.ts +27 -27
  139. package/src/server.ts +144 -10
  140. package/src/skills/mini-skills.ts +386 -386
  141. package/src/store/bun-sqlite-compat.ts +118 -15
  142. package/src/store/chat-store.ts +119 -119
  143. package/src/store/graph-store.ts +249 -249
  144. package/src/store/mini-skill-store.ts +349 -349
  145. package/src/store/persistence-json.ts +212 -212
  146. package/src/store/persistence.ts +291 -291
  147. package/src/store/project-affinity.ts +195 -195
  148. package/src/store/sqlite-db.ts +3 -3
  149. package/src/team/event-bus.ts +76 -76
  150. package/src/team/file-locks.ts +173 -173
  151. package/src/team/handoff.ts +161 -161
  152. package/src/team/messages.ts +203 -203
  153. package/src/team/poll.ts +132 -132
  154. package/src/team/tasks.ts +211 -211
  155. package/src/workspace/mcp-adapters/codex.ts +191 -191
  156. package/src/workspace/mcp-adapters/copilot.ts +105 -105
  157. package/src/workspace/mcp-adapters/cursor.ts +53 -53
  158. package/src/workspace/mcp-adapters/kiro.ts +64 -64
  159. package/src/workspace/mcp-adapters/opencode.ts +123 -123
  160. package/src/workspace/mcp-adapters/trae.ts +134 -134
  161. package/src/workspace/mcp-adapters/windsurf.ts +91 -91
  162. package/src/workspace/sanitizer.ts +60 -60
  163. package/src/workspace/workflow-sync.ts +131 -131
@@ -1,219 +1,219 @@
1
- /**
2
- * Cost Tracker — Phase 7, Step 8: Token → USD conversion and budget enforcement.
3
- *
4
- * Converts accumulated token usage (per model) into USD using a configurable
5
- * price table. Supports budget limits that abort the pipeline when exceeded.
6
- *
7
- * Design principle: if the price table doesn't have a model, report tokens only
8
- * (skip USD calculation). Never crash because of a missing price entry.
9
- */
10
-
11
- import type { TokenUsage } from './adapters/types.js';
12
-
13
- // ── Types ──────────────────────────────────────────────────────────
14
-
15
- export interface ModelPrice {
16
- /** Cost per 1M input tokens in USD */
17
- inputPer1M: number;
18
- /** Cost per 1M output tokens in USD */
19
- outputPer1M: number;
20
- /** Cost per 1M cache read tokens in USD (default: 0) */
21
- cacheReadPer1M?: number;
22
- /** Cost per 1M cache write tokens in USD (default: 0) */
23
- cacheWritePer1M?: number;
24
- }
25
-
26
- export interface CostSummary {
27
- /** Total cost in USD (null if no price data available) */
28
- totalUSD: number | null;
29
- /** Per-model breakdown */
30
- models: {
31
- model: string;
32
- inputTokens: number;
33
- outputTokens: number;
34
- cacheReadTokens: number;
35
- cacheWriteTokens: number;
36
- costUSD: number | null;
37
- }[];
38
- /** Whether the budget has been exceeded */
39
- budgetExceeded: boolean;
40
- /** Configured budget (null if no budget) */
41
- budgetUSD: number | null;
42
- }
43
-
44
- // ── Default Price Table ────────────────────────────────────────────
45
-
46
- const DEFAULT_PRICES: Record<string, ModelPrice> = {
47
- // Claude models — verified 2026-04-13 via Anthropic docs
48
- 'claude-opus-4-6': { inputPer1M: 5, outputPer1M: 25, cacheReadPer1M: 0.50, cacheWritePer1M: 6.25 },
49
- 'claude-sonnet-4-6': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
50
- 'claude-haiku-4-5': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
51
- 'claude-haiku-4-5-20251001': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
52
- // Legacy Claude models (still in use by some adapters)
53
- 'claude-sonnet-4-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
54
- 'claude-sonnet-4-0-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
55
- 'claude-opus-4-20250514': { inputPer1M: 15, outputPer1M: 75, cacheReadPer1M: 1.50, cacheWritePer1M: 18.75 },
56
- 'claude-3-5-sonnet': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
57
- 'claude-3-5-haiku': { inputPer1M: 0.80, outputPer1M: 4, cacheReadPer1M: 0.08, cacheWritePer1M: 1 },
58
-
59
- // OpenAI / Codex models — verified 2026-04-13
60
- 'gpt-5': { inputPer1M: 1.25, outputPer1M: 10 },
61
- 'gpt-5-mini': { inputPer1M: 0.25, outputPer1M: 2 },
62
- 'gpt-5-nano': { inputPer1M: 0.05, outputPer1M: 0.40 },
63
- 'gpt-5.2': { inputPer1M: 1.75, outputPer1M: 14 },
64
- 'o3': { inputPer1M: 2, outputPer1M: 8 },
65
- 'o3-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
66
- 'o4-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
67
- 'gpt-4.1': { inputPer1M: 2, outputPer1M: 8, cacheReadPer1M: 0.50, cacheWritePer1M: 2 },
68
- 'gpt-4.1-mini': { inputPer1M: 0.40, outputPer1M: 1.60, cacheReadPer1M: 0.10, cacheWritePer1M: 0.40 },
69
- 'gpt-4.1-nano': { inputPer1M: 0.10, outputPer1M: 0.40 },
70
- 'gpt-4o': { inputPer1M: 2.50, outputPer1M: 10, cacheReadPer1M: 1.25 },
71
- 'gpt-4o-mini': { inputPer1M: 0.15, outputPer1M: 0.60, cacheReadPer1M: 0.075 },
72
-
73
- // Google Gemini models — verified 2026-04-13 via Google AI pricing
74
- 'gemini-2.5-flash': { inputPer1M: 0.30, outputPer1M: 2.50 },
75
- 'gemini-2.5-flash-preview-05-20': { inputPer1M: 0.30, outputPer1M: 2.50 },
76
- 'gemini-2.5-flash-lite': { inputPer1M: 0.10, outputPer1M: 0.40 },
77
- 'gemini-2.5-pro': { inputPer1M: 1.25, outputPer1M: 10 },
78
- 'gemini-2.0-flash': { inputPer1M: 0.10, outputPer1M: 0.40 },
79
- 'gemini-2.0-flash-lite': { inputPer1M: 0.075, outputPer1M: 0.30 },
80
- // Gemini 3.x generation
81
- 'gemini-3.1-pro': { inputPer1M: 2, outputPer1M: 12 },
82
- 'gemini-3-flash': { inputPer1M: 0.50, outputPer1M: 3 },
83
-
84
- // OpenCode / open models — verified 2026-04-13
85
- 'deepseek-coder': { inputPer1M: 0.28, outputPer1M: 0.42 },
86
- 'deepseek-chat': { inputPer1M: 0.28, outputPer1M: 0.42 },
87
- };
88
-
89
- // ── Core ───────────────────────────────────────────────────────────
90
-
91
- /**
92
- * Calculate cost for a single model's token usage.
93
- * Returns null if the model is not in the price table.
94
- */
95
- export function calculateModelCost(
96
- model: string,
97
- usage: TokenUsage,
98
- customPrices?: Record<string, ModelPrice>,
99
- ): number | null {
100
- const prices = { ...DEFAULT_PRICES, ...customPrices };
101
- const price = findPrice(model, prices);
102
- if (!price) return null;
103
-
104
- const inputCost = (usage.inputTokens / 1_000_000) * price.inputPer1M;
105
- const outputCost = (usage.outputTokens / 1_000_000) * price.outputPer1M;
106
- const cacheReadCost = (usage.cacheReadTokens / 1_000_000) * (price.cacheReadPer1M ?? 0);
107
- const cacheWriteCost = (usage.cacheWriteTokens / 1_000_000) * (price.cacheWritePer1M ?? 0);
108
-
109
- return inputCost + outputCost + cacheReadCost + cacheWriteCost;
110
- }
111
-
112
- /**
113
- * Calculate total pipeline cost from accumulated token usage.
114
- */
115
- export function calculatePipelineCost(
116
- tokenUsage: Record<string, TokenUsage>,
117
- budgetUSD?: number,
118
- customPrices?: Record<string, ModelPrice>,
119
- ): CostSummary {
120
- const models: CostSummary['models'] = [];
121
- let totalUSD: number | null = 0;
122
- let hasAnyPrice = false;
123
-
124
- for (const [model, usage] of Object.entries(tokenUsage)) {
125
- const cost = calculateModelCost(model, usage, customPrices);
126
- if (cost !== null) {
127
- hasAnyPrice = true;
128
- totalUSD = (totalUSD ?? 0) + cost;
129
- }
130
-
131
- models.push({
132
- model,
133
- inputTokens: usage.inputTokens,
134
- outputTokens: usage.outputTokens,
135
- cacheReadTokens: usage.cacheReadTokens,
136
- cacheWriteTokens: usage.cacheWriteTokens,
137
- costUSD: cost !== null ? Math.round(cost * 10000) / 10000 : null, // Round to 4 decimals
138
- });
139
- }
140
-
141
- if (!hasAnyPrice) totalUSD = null;
142
- else totalUSD = Math.round(totalUSD! * 10000) / 10000;
143
-
144
- return {
145
- totalUSD,
146
- models,
147
- budgetExceeded: budgetUSD != null && totalUSD != null && totalUSD > budgetUSD,
148
- budgetUSD: budgetUSD ?? null,
149
- };
150
- }
151
-
152
- /**
153
- * Check if the current cost exceeds the budget.
154
- * Returns false if no budget is set or cost cannot be calculated.
155
- */
156
- export function isBudgetExceeded(
157
- tokenUsage: Record<string, TokenUsage>,
158
- budgetUSD?: number,
159
- customPrices?: Record<string, ModelPrice>,
160
- ): boolean {
161
- if (budgetUSD == null) return false;
162
- const summary = calculatePipelineCost(tokenUsage, budgetUSD, customPrices);
163
- return summary.budgetExceeded;
164
- }
165
-
166
- /**
167
- * Format cost summary for CLI display.
168
- */
169
- export function formatCostSummary(summary: CostSummary): string {
170
- const lines: string[] = [];
171
-
172
- for (const m of summary.models) {
173
- const tokens = `in=${fmtNum(m.inputTokens)} out=${fmtNum(m.outputTokens)}`;
174
- const cache = m.cacheReadTokens > 0 || m.cacheWriteTokens > 0
175
- ? ` cache_r=${fmtNum(m.cacheReadTokens)} cache_w=${fmtNum(m.cacheWriteTokens)}`
176
- : '';
177
- const cost = m.costUSD !== null ? ` ($${m.costUSD.toFixed(4)})` : '';
178
- lines.push(` ${m.model}: ${tokens}${cache}${cost}`);
179
- }
180
-
181
- if (summary.totalUSD !== null) {
182
- lines.push(` Total: $${summary.totalUSD.toFixed(4)}`);
183
- }
184
-
185
- if (summary.budgetUSD !== null) {
186
- const status = summary.budgetExceeded ? ' [WARN] EXCEEDED' : '';
187
- lines.push(` Budget: $${summary.budgetUSD.toFixed(2)}${status}`);
188
- }
189
-
190
- return lines.join('\n');
191
- }
192
-
193
- // ── Helpers ────────────────────────────────────────────────────────
194
-
195
- /**
196
- * Find price entry by model name. Supports fuzzy matching:
197
- * "claude-sonnet-4-20250514" matches "claude-sonnet-4-20250514" exactly,
198
- * but also tries prefix matching for versioned models.
199
- */
200
- function findPrice(model: string, prices: Record<string, ModelPrice>): ModelPrice | null {
201
- // Exact match
202
- if (prices[model]) return prices[model];
203
-
204
- // Prefix match: "claude-sonnet-4-20250514" → try "claude-sonnet-4"
205
- const normalized = model.toLowerCase();
206
- for (const [key, price] of Object.entries(prices)) {
207
- if (normalized.startsWith(key.toLowerCase()) || key.toLowerCase().startsWith(normalized)) {
208
- return price;
209
- }
210
- }
211
-
212
- return null;
213
- }
214
-
215
- function fmtNum(n: number): string {
216
- if (n >= 1_000_000) return `${(n / 1_000_000).toFixed(1)}M`;
217
- if (n >= 1_000) return `${(n / 1_000).toFixed(1)}K`;
218
- return String(n);
219
- }
1
+ /**
2
+ * Cost Tracker — Phase 7, Step 8: Token → USD conversion and budget enforcement.
3
+ *
4
+ * Converts accumulated token usage (per model) into USD using a configurable
5
+ * price table. Supports budget limits that abort the pipeline when exceeded.
6
+ *
7
+ * Design principle: if the price table doesn't have a model, report tokens only
8
+ * (skip USD calculation). Never crash because of a missing price entry.
9
+ */
10
+
11
+ import type { TokenUsage } from './adapters/types.js';
12
+
13
+ // ── Types ──────────────────────────────────────────────────────────
14
+
15
+ export interface ModelPrice {
16
+ /** Cost per 1M input tokens in USD */
17
+ inputPer1M: number;
18
+ /** Cost per 1M output tokens in USD */
19
+ outputPer1M: number;
20
+ /** Cost per 1M cache read tokens in USD (default: 0) */
21
+ cacheReadPer1M?: number;
22
+ /** Cost per 1M cache write tokens in USD (default: 0) */
23
+ cacheWritePer1M?: number;
24
+ }
25
+
26
+ export interface CostSummary {
27
+ /** Total cost in USD (null if no price data available) */
28
+ totalUSD: number | null;
29
+ /** Per-model breakdown */
30
+ models: {
31
+ model: string;
32
+ inputTokens: number;
33
+ outputTokens: number;
34
+ cacheReadTokens: number;
35
+ cacheWriteTokens: number;
36
+ costUSD: number | null;
37
+ }[];
38
+ /** Whether the budget has been exceeded */
39
+ budgetExceeded: boolean;
40
+ /** Configured budget (null if no budget) */
41
+ budgetUSD: number | null;
42
+ }
43
+
44
+ // ── Default Price Table ────────────────────────────────────────────
45
+
46
+ const DEFAULT_PRICES: Record<string, ModelPrice> = {
47
+ // Claude models — verified 2026-04-13 via Anthropic docs
48
+ 'claude-opus-4-6': { inputPer1M: 5, outputPer1M: 25, cacheReadPer1M: 0.50, cacheWritePer1M: 6.25 },
49
+ 'claude-sonnet-4-6': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
50
+ 'claude-haiku-4-5': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
51
+ 'claude-haiku-4-5-20251001': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
52
+ // Legacy Claude models (still in use by some adapters)
53
+ 'claude-sonnet-4-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
54
+ 'claude-sonnet-4-0-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
55
+ 'claude-opus-4-20250514': { inputPer1M: 15, outputPer1M: 75, cacheReadPer1M: 1.50, cacheWritePer1M: 18.75 },
56
+ 'claude-3-5-sonnet': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
57
+ 'claude-3-5-haiku': { inputPer1M: 0.80, outputPer1M: 4, cacheReadPer1M: 0.08, cacheWritePer1M: 1 },
58
+
59
+ // OpenAI / Codex models — verified 2026-04-13
60
+ 'gpt-5': { inputPer1M: 1.25, outputPer1M: 10 },
61
+ 'gpt-5-mini': { inputPer1M: 0.25, outputPer1M: 2 },
62
+ 'gpt-5-nano': { inputPer1M: 0.05, outputPer1M: 0.40 },
63
+ 'gpt-5.2': { inputPer1M: 1.75, outputPer1M: 14 },
64
+ 'o3': { inputPer1M: 2, outputPer1M: 8 },
65
+ 'o3-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
66
+ 'o4-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
67
+ 'gpt-4.1': { inputPer1M: 2, outputPer1M: 8, cacheReadPer1M: 0.50, cacheWritePer1M: 2 },
68
+ 'gpt-4.1-mini': { inputPer1M: 0.40, outputPer1M: 1.60, cacheReadPer1M: 0.10, cacheWritePer1M: 0.40 },
69
+ 'gpt-4.1-nano': { inputPer1M: 0.10, outputPer1M: 0.40 },
70
+ 'gpt-4o': { inputPer1M: 2.50, outputPer1M: 10, cacheReadPer1M: 1.25 },
71
+ 'gpt-4o-mini': { inputPer1M: 0.15, outputPer1M: 0.60, cacheReadPer1M: 0.075 },
72
+
73
+ // Google Gemini models — verified 2026-04-13 via Google AI pricing
74
+ 'gemini-2.5-flash': { inputPer1M: 0.30, outputPer1M: 2.50 },
75
+ 'gemini-2.5-flash-preview-05-20': { inputPer1M: 0.30, outputPer1M: 2.50 },
76
+ 'gemini-2.5-flash-lite': { inputPer1M: 0.10, outputPer1M: 0.40 },
77
+ 'gemini-2.5-pro': { inputPer1M: 1.25, outputPer1M: 10 },
78
+ 'gemini-2.0-flash': { inputPer1M: 0.10, outputPer1M: 0.40 },
79
+ 'gemini-2.0-flash-lite': { inputPer1M: 0.075, outputPer1M: 0.30 },
80
+ // Gemini 3.x generation
81
+ 'gemini-3.1-pro': { inputPer1M: 2, outputPer1M: 12 },
82
+ 'gemini-3-flash': { inputPer1M: 0.50, outputPer1M: 3 },
83
+
84
+ // OpenCode / open models — verified 2026-04-13
85
+ 'deepseek-coder': { inputPer1M: 0.28, outputPer1M: 0.42 },
86
+ 'deepseek-chat': { inputPer1M: 0.28, outputPer1M: 0.42 },
87
+ };
88
+
89
+ // ── Core ───────────────────────────────────────────────────────────
90
+
91
+ /**
92
+ * Calculate cost for a single model's token usage.
93
+ * Returns null if the model is not in the price table.
94
+ */
95
+ export function calculateModelCost(
96
+ model: string,
97
+ usage: TokenUsage,
98
+ customPrices?: Record<string, ModelPrice>,
99
+ ): number | null {
100
+ const prices = { ...DEFAULT_PRICES, ...customPrices };
101
+ const price = findPrice(model, prices);
102
+ if (!price) return null;
103
+
104
+ const inputCost = (usage.inputTokens / 1_000_000) * price.inputPer1M;
105
+ const outputCost = (usage.outputTokens / 1_000_000) * price.outputPer1M;
106
+ const cacheReadCost = (usage.cacheReadTokens / 1_000_000) * (price.cacheReadPer1M ?? 0);
107
+ const cacheWriteCost = (usage.cacheWriteTokens / 1_000_000) * (price.cacheWritePer1M ?? 0);
108
+
109
+ return inputCost + outputCost + cacheReadCost + cacheWriteCost;
110
+ }
111
+
112
+ /**
113
+ * Calculate total pipeline cost from accumulated token usage.
114
+ */
115
+ export function calculatePipelineCost(
116
+ tokenUsage: Record<string, TokenUsage>,
117
+ budgetUSD?: number,
118
+ customPrices?: Record<string, ModelPrice>,
119
+ ): CostSummary {
120
+ const models: CostSummary['models'] = [];
121
+ let totalUSD: number | null = 0;
122
+ let hasAnyPrice = false;
123
+
124
+ for (const [model, usage] of Object.entries(tokenUsage)) {
125
+ const cost = calculateModelCost(model, usage, customPrices);
126
+ if (cost !== null) {
127
+ hasAnyPrice = true;
128
+ totalUSD = (totalUSD ?? 0) + cost;
129
+ }
130
+
131
+ models.push({
132
+ model,
133
+ inputTokens: usage.inputTokens,
134
+ outputTokens: usage.outputTokens,
135
+ cacheReadTokens: usage.cacheReadTokens,
136
+ cacheWriteTokens: usage.cacheWriteTokens,
137
+ costUSD: cost !== null ? Math.round(cost * 10000) / 10000 : null, // Round to 4 decimals
138
+ });
139
+ }
140
+
141
+ if (!hasAnyPrice) totalUSD = null;
142
+ else totalUSD = Math.round(totalUSD! * 10000) / 10000;
143
+
144
+ return {
145
+ totalUSD,
146
+ models,
147
+ budgetExceeded: budgetUSD != null && totalUSD != null && totalUSD > budgetUSD,
148
+ budgetUSD: budgetUSD ?? null,
149
+ };
150
+ }
151
+
152
+ /**
153
+ * Check if the current cost exceeds the budget.
154
+ * Returns false if no budget is set or cost cannot be calculated.
155
+ */
156
+ export function isBudgetExceeded(
157
+ tokenUsage: Record<string, TokenUsage>,
158
+ budgetUSD?: number,
159
+ customPrices?: Record<string, ModelPrice>,
160
+ ): boolean {
161
+ if (budgetUSD == null) return false;
162
+ const summary = calculatePipelineCost(tokenUsage, budgetUSD, customPrices);
163
+ return summary.budgetExceeded;
164
+ }
165
+
166
+ /**
167
+ * Format cost summary for CLI display.
168
+ */
169
+ export function formatCostSummary(summary: CostSummary): string {
170
+ const lines: string[] = [];
171
+
172
+ for (const m of summary.models) {
173
+ const tokens = `in=${fmtNum(m.inputTokens)} out=${fmtNum(m.outputTokens)}`;
174
+ const cache = m.cacheReadTokens > 0 || m.cacheWriteTokens > 0
175
+ ? ` cache_r=${fmtNum(m.cacheReadTokens)} cache_w=${fmtNum(m.cacheWriteTokens)}`
176
+ : '';
177
+ const cost = m.costUSD !== null ? ` ($${m.costUSD.toFixed(4)})` : '';
178
+ lines.push(` ${m.model}: ${tokens}${cache}${cost}`);
179
+ }
180
+
181
+ if (summary.totalUSD !== null) {
182
+ lines.push(` Total: $${summary.totalUSD.toFixed(4)}`);
183
+ }
184
+
185
+ if (summary.budgetUSD !== null) {
186
+ const status = summary.budgetExceeded ? ' [WARN] EXCEEDED' : '';
187
+ lines.push(` Budget: $${summary.budgetUSD.toFixed(2)}${status}`);
188
+ }
189
+
190
+ return lines.join('\n');
191
+ }
192
+
193
+ // ── Helpers ────────────────────────────────────────────────────────
194
+
195
+ /**
196
+ * Find price entry by model name. Supports fuzzy matching:
197
+ * "claude-sonnet-4-20250514" matches "claude-sonnet-4-20250514" exactly,
198
+ * but also tries prefix matching for versioned models.
199
+ */
200
+ function findPrice(model: string, prices: Record<string, ModelPrice>): ModelPrice | null {
201
+ // Exact match
202
+ if (prices[model]) return prices[model];
203
+
204
+ // Prefix match: "claude-sonnet-4-20250514" → try "claude-sonnet-4"
205
+ const normalized = model.toLowerCase();
206
+ for (const [key, price] of Object.entries(prices)) {
207
+ if (normalized.startsWith(key.toLowerCase()) || key.toLowerCase().startsWith(normalized)) {
208
+ return price;
209
+ }
210
+ }
211
+
212
+ return null;
213
+ }
214
+
215
+ function fmtNum(n: number): string {
216
+ if (n >= 1_000_000) return `${(n / 1_000_000).toFixed(1)}M`;
217
+ if (n >= 1_000) return `${(n / 1_000).toFixed(1)}K`;
218
+ return String(n);
219
+ }