memorix 1.2.2 → 1.2.3

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (139) hide show
  1. package/CHANGELOG.md +7 -0
  2. package/TEAM.md +86 -86
  3. package/dist/cli/index.js +34 -18
  4. package/dist/cli/index.js.map +1 -1
  5. package/dist/index.js +17 -8
  6. package/dist/index.js.map +1 -1
  7. package/dist/maintenance-runner.js.map +1 -1
  8. package/dist/memcode-runtime/CHANGELOG.md +7 -0
  9. package/dist/sdk.js +17 -8
  10. package/dist/sdk.js.map +1 -1
  11. package/docs/DESIGN_DECISIONS.md +357 -357
  12. package/docs/dev-log/progress.txt +18 -8
  13. package/package.json +1 -1
  14. package/plugins/codex/memorix/.codex-plugin/plugin.json +1 -1
  15. package/src/audit/index.ts +156 -156
  16. package/src/cli/commands/audit-list.ts +89 -89
  17. package/src/cli/commands/background.ts +659 -659
  18. package/src/cli/commands/formation.ts +48 -48
  19. package/src/cli/commands/git-hook-install.ts +111 -111
  20. package/src/cli/commands/handoff.ts +54 -54
  21. package/src/cli/commands/hooks-status.ts +63 -63
  22. package/src/cli/commands/ingest-commit.ts +153 -153
  23. package/src/cli/commands/ingest-image.ts +66 -66
  24. package/src/cli/commands/ingest-log.ts +180 -180
  25. package/src/cli/commands/ingest.ts +44 -44
  26. package/src/cli/commands/integrate-shared.ts +15 -15
  27. package/src/cli/commands/lock.ts +82 -82
  28. package/src/cli/commands/message.ts +104 -104
  29. package/src/cli/commands/poll.ts +58 -58
  30. package/src/cli/commands/purge-all-memory.ts +85 -85
  31. package/src/cli/commands/purge-project-memory.ts +83 -83
  32. package/src/cli/commands/reasoning.ts +118 -118
  33. package/src/cli/commands/serve-shared.ts +118 -118
  34. package/src/cli/commands/session.ts +15 -7
  35. package/src/cli/commands/skills.ts +114 -114
  36. package/src/cli/commands/task.ts +167 -167
  37. package/src/cli/commands/transfer.ts +47 -47
  38. package/src/cli/commands/uninstall-project-artifacts.ts +85 -85
  39. package/src/cli/tui/ChatView.tsx +234 -234
  40. package/src/cli/tui/CommandBar.tsx +312 -312
  41. package/src/cli/tui/ContextRail.tsx +118 -118
  42. package/src/cli/tui/HeaderBar.tsx +72 -72
  43. package/src/cli/tui/LogoBanner.tsx +51 -51
  44. package/src/cli/tui/Sidebar.tsx +179 -179
  45. package/src/cli/tui/index.ts +41 -41
  46. package/src/cli/tui/markdown-render.tsx +371 -371
  47. package/src/cli/tui/session-service.ts +3 -2
  48. package/src/cli/tui/use-mouse.ts +157 -157
  49. package/src/cli/tui/useNavigation.ts +56 -56
  50. package/src/cli/update-checker.ts +211 -211
  51. package/src/cli/version.ts +7 -7
  52. package/src/cli/workbench.ts +1 -1
  53. package/src/compact/token-budget.ts +74 -74
  54. package/src/dashboard/project-classification.ts +64 -64
  55. package/src/embedding/fastembed-provider.ts +142 -142
  56. package/src/embedding/transformers-provider.ts +111 -111
  57. package/src/git/extractor.ts +209 -209
  58. package/src/git/hooks-path.ts +85 -85
  59. package/src/hooks/pattern-detector.ts +173 -173
  60. package/src/hooks/significance-filter.ts +250 -250
  61. package/src/llm/memory-manager.ts +328 -328
  62. package/src/llm/provider.ts +885 -885
  63. package/src/llm/quality.ts +248 -248
  64. package/src/memory/attribution-guard.ts +249 -249
  65. package/src/memory/disclosure-policy.ts +135 -135
  66. package/src/memory/entity-extractor.ts +197 -197
  67. package/src/memory/formation/evaluate.ts +217 -217
  68. package/src/memory/formation/extract.ts +361 -361
  69. package/src/memory/formation/index.ts +417 -417
  70. package/src/memory/formation/resolve.ts +344 -344
  71. package/src/memory/formation/types.ts +315 -315
  72. package/src/memory/freshness.ts +122 -122
  73. package/src/memory/graph.ts +197 -197
  74. package/src/memory/refs.ts +94 -94
  75. package/src/memory/secret-filter.ts +79 -79
  76. package/src/memory/session.ts +24 -9
  77. package/src/multimodal/image-loader.ts +143 -143
  78. package/src/orchestrate/adapters/claude-stream.ts +192 -192
  79. package/src/orchestrate/adapters/claude.ts +111 -111
  80. package/src/orchestrate/adapters/codex-stream.ts +134 -134
  81. package/src/orchestrate/adapters/codex.ts +41 -41
  82. package/src/orchestrate/adapters/gemini-stream.ts +166 -166
  83. package/src/orchestrate/adapters/gemini.ts +42 -42
  84. package/src/orchestrate/adapters/index.ts +73 -73
  85. package/src/orchestrate/adapters/opencode-stream.ts +143 -143
  86. package/src/orchestrate/adapters/opencode.ts +47 -47
  87. package/src/orchestrate/adapters/spawn-helper.ts +286 -286
  88. package/src/orchestrate/adapters/types.ts +77 -77
  89. package/src/orchestrate/capability-router.ts +284 -284
  90. package/src/orchestrate/context-compact.ts +188 -188
  91. package/src/orchestrate/cost-tracker.ts +219 -219
  92. package/src/orchestrate/error-recovery.ts +191 -191
  93. package/src/orchestrate/evidence.ts +140 -140
  94. package/src/orchestrate/ledger.ts +110 -110
  95. package/src/orchestrate/memorix-bridge.ts +343 -343
  96. package/src/orchestrate/output-budget.ts +80 -80
  97. package/src/orchestrate/permission.ts +152 -152
  98. package/src/orchestrate/pipeline-trace.ts +131 -131
  99. package/src/orchestrate/prompt-builder.ts +155 -155
  100. package/src/orchestrate/ring-buffer.ts +37 -37
  101. package/src/orchestrate/task-graph.ts +389 -389
  102. package/src/orchestrate/worktree.ts +232 -232
  103. package/src/project/aliases.ts +374 -374
  104. package/src/project/detector.ts +268 -268
  105. package/src/rules/adapters/claude-code.ts +99 -99
  106. package/src/rules/adapters/codex.ts +97 -97
  107. package/src/rules/adapters/copilot.ts +124 -124
  108. package/src/rules/adapters/cursor.ts +114 -114
  109. package/src/rules/adapters/kiro.ts +126 -126
  110. package/src/rules/adapters/trae.ts +56 -56
  111. package/src/rules/adapters/windsurf.ts +83 -83
  112. package/src/rules/syncer.ts +235 -235
  113. package/src/sdk.ts +299 -299
  114. package/src/search/intent-detector.ts +289 -289
  115. package/src/search/query-expansion.ts +52 -52
  116. package/src/server/formation-timeout.ts +27 -27
  117. package/src/server.ts +7 -2
  118. package/src/skills/mini-skills.ts +386 -386
  119. package/src/store/chat-store.ts +119 -119
  120. package/src/store/graph-store.ts +249 -249
  121. package/src/store/mini-skill-store.ts +349 -349
  122. package/src/store/persistence-json.ts +212 -212
  123. package/src/store/persistence.ts +291 -291
  124. package/src/store/project-affinity.ts +195 -195
  125. package/src/team/event-bus.ts +76 -76
  126. package/src/team/file-locks.ts +173 -173
  127. package/src/team/handoff.ts +161 -161
  128. package/src/team/messages.ts +203 -203
  129. package/src/team/poll.ts +132 -132
  130. package/src/team/tasks.ts +211 -211
  131. package/src/workspace/mcp-adapters/codex.ts +191 -191
  132. package/src/workspace/mcp-adapters/copilot.ts +105 -105
  133. package/src/workspace/mcp-adapters/cursor.ts +53 -53
  134. package/src/workspace/mcp-adapters/kiro.ts +64 -64
  135. package/src/workspace/mcp-adapters/opencode.ts +123 -123
  136. package/src/workspace/mcp-adapters/trae.ts +134 -134
  137. package/src/workspace/mcp-adapters/windsurf.ts +91 -91
  138. package/src/workspace/sanitizer.ts +60 -60
  139. package/src/workspace/workflow-sync.ts +131 -131
@@ -1,219 +1,219 @@
1
- /**
2
- * Cost Tracker — Phase 7, Step 8: Token → USD conversion and budget enforcement.
3
- *
4
- * Converts accumulated token usage (per model) into USD using a configurable
5
- * price table. Supports budget limits that abort the pipeline when exceeded.
6
- *
7
- * Design principle: if the price table doesn't have a model, report tokens only
8
- * (skip USD calculation). Never crash because of a missing price entry.
9
- */
10
-
11
- import type { TokenUsage } from './adapters/types.js';
12
-
13
- // ── Types ──────────────────────────────────────────────────────────
14
-
15
- export interface ModelPrice {
16
- /** Cost per 1M input tokens in USD */
17
- inputPer1M: number;
18
- /** Cost per 1M output tokens in USD */
19
- outputPer1M: number;
20
- /** Cost per 1M cache read tokens in USD (default: 0) */
21
- cacheReadPer1M?: number;
22
- /** Cost per 1M cache write tokens in USD (default: 0) */
23
- cacheWritePer1M?: number;
24
- }
25
-
26
- export interface CostSummary {
27
- /** Total cost in USD (null if no price data available) */
28
- totalUSD: number | null;
29
- /** Per-model breakdown */
30
- models: {
31
- model: string;
32
- inputTokens: number;
33
- outputTokens: number;
34
- cacheReadTokens: number;
35
- cacheWriteTokens: number;
36
- costUSD: number | null;
37
- }[];
38
- /** Whether the budget has been exceeded */
39
- budgetExceeded: boolean;
40
- /** Configured budget (null if no budget) */
41
- budgetUSD: number | null;
42
- }
43
-
44
- // ── Default Price Table ────────────────────────────────────────────
45
-
46
- const DEFAULT_PRICES: Record<string, ModelPrice> = {
47
- // Claude models — verified 2026-04-13 via Anthropic docs
48
- 'claude-opus-4-6': { inputPer1M: 5, outputPer1M: 25, cacheReadPer1M: 0.50, cacheWritePer1M: 6.25 },
49
- 'claude-sonnet-4-6': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
50
- 'claude-haiku-4-5': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
51
- 'claude-haiku-4-5-20251001': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
52
- // Legacy Claude models (still in use by some adapters)
53
- 'claude-sonnet-4-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
54
- 'claude-sonnet-4-0-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
55
- 'claude-opus-4-20250514': { inputPer1M: 15, outputPer1M: 75, cacheReadPer1M: 1.50, cacheWritePer1M: 18.75 },
56
- 'claude-3-5-sonnet': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
57
- 'claude-3-5-haiku': { inputPer1M: 0.80, outputPer1M: 4, cacheReadPer1M: 0.08, cacheWritePer1M: 1 },
58
-
59
- // OpenAI / Codex models — verified 2026-04-13
60
- 'gpt-5': { inputPer1M: 1.25, outputPer1M: 10 },
61
- 'gpt-5-mini': { inputPer1M: 0.25, outputPer1M: 2 },
62
- 'gpt-5-nano': { inputPer1M: 0.05, outputPer1M: 0.40 },
63
- 'gpt-5.2': { inputPer1M: 1.75, outputPer1M: 14 },
64
- 'o3': { inputPer1M: 2, outputPer1M: 8 },
65
- 'o3-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
66
- 'o4-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
67
- 'gpt-4.1': { inputPer1M: 2, outputPer1M: 8, cacheReadPer1M: 0.50, cacheWritePer1M: 2 },
68
- 'gpt-4.1-mini': { inputPer1M: 0.40, outputPer1M: 1.60, cacheReadPer1M: 0.10, cacheWritePer1M: 0.40 },
69
- 'gpt-4.1-nano': { inputPer1M: 0.10, outputPer1M: 0.40 },
70
- 'gpt-4o': { inputPer1M: 2.50, outputPer1M: 10, cacheReadPer1M: 1.25 },
71
- 'gpt-4o-mini': { inputPer1M: 0.15, outputPer1M: 0.60, cacheReadPer1M: 0.075 },
72
-
73
- // Google Gemini models — verified 2026-04-13 via Google AI pricing
74
- 'gemini-2.5-flash': { inputPer1M: 0.30, outputPer1M: 2.50 },
75
- 'gemini-2.5-flash-preview-05-20': { inputPer1M: 0.30, outputPer1M: 2.50 },
76
- 'gemini-2.5-flash-lite': { inputPer1M: 0.10, outputPer1M: 0.40 },
77
- 'gemini-2.5-pro': { inputPer1M: 1.25, outputPer1M: 10 },
78
- 'gemini-2.0-flash': { inputPer1M: 0.10, outputPer1M: 0.40 },
79
- 'gemini-2.0-flash-lite': { inputPer1M: 0.075, outputPer1M: 0.30 },
80
- // Gemini 3.x generation
81
- 'gemini-3.1-pro': { inputPer1M: 2, outputPer1M: 12 },
82
- 'gemini-3-flash': { inputPer1M: 0.50, outputPer1M: 3 },
83
-
84
- // OpenCode / open models — verified 2026-04-13
85
- 'deepseek-coder': { inputPer1M: 0.28, outputPer1M: 0.42 },
86
- 'deepseek-chat': { inputPer1M: 0.28, outputPer1M: 0.42 },
87
- };
88
-
89
- // ── Core ───────────────────────────────────────────────────────────
90
-
91
- /**
92
- * Calculate cost for a single model's token usage.
93
- * Returns null if the model is not in the price table.
94
- */
95
- export function calculateModelCost(
96
- model: string,
97
- usage: TokenUsage,
98
- customPrices?: Record<string, ModelPrice>,
99
- ): number | null {
100
- const prices = { ...DEFAULT_PRICES, ...customPrices };
101
- const price = findPrice(model, prices);
102
- if (!price) return null;
103
-
104
- const inputCost = (usage.inputTokens / 1_000_000) * price.inputPer1M;
105
- const outputCost = (usage.outputTokens / 1_000_000) * price.outputPer1M;
106
- const cacheReadCost = (usage.cacheReadTokens / 1_000_000) * (price.cacheReadPer1M ?? 0);
107
- const cacheWriteCost = (usage.cacheWriteTokens / 1_000_000) * (price.cacheWritePer1M ?? 0);
108
-
109
- return inputCost + outputCost + cacheReadCost + cacheWriteCost;
110
- }
111
-
112
- /**
113
- * Calculate total pipeline cost from accumulated token usage.
114
- */
115
- export function calculatePipelineCost(
116
- tokenUsage: Record<string, TokenUsage>,
117
- budgetUSD?: number,
118
- customPrices?: Record<string, ModelPrice>,
119
- ): CostSummary {
120
- const models: CostSummary['models'] = [];
121
- let totalUSD: number | null = 0;
122
- let hasAnyPrice = false;
123
-
124
- for (const [model, usage] of Object.entries(tokenUsage)) {
125
- const cost = calculateModelCost(model, usage, customPrices);
126
- if (cost !== null) {
127
- hasAnyPrice = true;
128
- totalUSD = (totalUSD ?? 0) + cost;
129
- }
130
-
131
- models.push({
132
- model,
133
- inputTokens: usage.inputTokens,
134
- outputTokens: usage.outputTokens,
135
- cacheReadTokens: usage.cacheReadTokens,
136
- cacheWriteTokens: usage.cacheWriteTokens,
137
- costUSD: cost !== null ? Math.round(cost * 10000) / 10000 : null, // Round to 4 decimals
138
- });
139
- }
140
-
141
- if (!hasAnyPrice) totalUSD = null;
142
- else totalUSD = Math.round(totalUSD! * 10000) / 10000;
143
-
144
- return {
145
- totalUSD,
146
- models,
147
- budgetExceeded: budgetUSD != null && totalUSD != null && totalUSD > budgetUSD,
148
- budgetUSD: budgetUSD ?? null,
149
- };
150
- }
151
-
152
- /**
153
- * Check if the current cost exceeds the budget.
154
- * Returns false if no budget is set or cost cannot be calculated.
155
- */
156
- export function isBudgetExceeded(
157
- tokenUsage: Record<string, TokenUsage>,
158
- budgetUSD?: number,
159
- customPrices?: Record<string, ModelPrice>,
160
- ): boolean {
161
- if (budgetUSD == null) return false;
162
- const summary = calculatePipelineCost(tokenUsage, budgetUSD, customPrices);
163
- return summary.budgetExceeded;
164
- }
165
-
166
- /**
167
- * Format cost summary for CLI display.
168
- */
169
- export function formatCostSummary(summary: CostSummary): string {
170
- const lines: string[] = [];
171
-
172
- for (const m of summary.models) {
173
- const tokens = `in=${fmtNum(m.inputTokens)} out=${fmtNum(m.outputTokens)}`;
174
- const cache = m.cacheReadTokens > 0 || m.cacheWriteTokens > 0
175
- ? ` cache_r=${fmtNum(m.cacheReadTokens)} cache_w=${fmtNum(m.cacheWriteTokens)}`
176
- : '';
177
- const cost = m.costUSD !== null ? ` ($${m.costUSD.toFixed(4)})` : '';
178
- lines.push(` ${m.model}: ${tokens}${cache}${cost}`);
179
- }
180
-
181
- if (summary.totalUSD !== null) {
182
- lines.push(` Total: $${summary.totalUSD.toFixed(4)}`);
183
- }
184
-
185
- if (summary.budgetUSD !== null) {
186
- const status = summary.budgetExceeded ? ' [WARN] EXCEEDED' : '';
187
- lines.push(` Budget: $${summary.budgetUSD.toFixed(2)}${status}`);
188
- }
189
-
190
- return lines.join('\n');
191
- }
192
-
193
- // ── Helpers ────────────────────────────────────────────────────────
194
-
195
- /**
196
- * Find price entry by model name. Supports fuzzy matching:
197
- * "claude-sonnet-4-20250514" matches "claude-sonnet-4-20250514" exactly,
198
- * but also tries prefix matching for versioned models.
199
- */
200
- function findPrice(model: string, prices: Record<string, ModelPrice>): ModelPrice | null {
201
- // Exact match
202
- if (prices[model]) return prices[model];
203
-
204
- // Prefix match: "claude-sonnet-4-20250514" → try "claude-sonnet-4"
205
- const normalized = model.toLowerCase();
206
- for (const [key, price] of Object.entries(prices)) {
207
- if (normalized.startsWith(key.toLowerCase()) || key.toLowerCase().startsWith(normalized)) {
208
- return price;
209
- }
210
- }
211
-
212
- return null;
213
- }
214
-
215
- function fmtNum(n: number): string {
216
- if (n >= 1_000_000) return `${(n / 1_000_000).toFixed(1)}M`;
217
- if (n >= 1_000) return `${(n / 1_000).toFixed(1)}K`;
218
- return String(n);
219
- }
1
+ /**
2
+ * Cost Tracker — Phase 7, Step 8: Token → USD conversion and budget enforcement.
3
+ *
4
+ * Converts accumulated token usage (per model) into USD using a configurable
5
+ * price table. Supports budget limits that abort the pipeline when exceeded.
6
+ *
7
+ * Design principle: if the price table doesn't have a model, report tokens only
8
+ * (skip USD calculation). Never crash because of a missing price entry.
9
+ */
10
+
11
+ import type { TokenUsage } from './adapters/types.js';
12
+
13
+ // ── Types ──────────────────────────────────────────────────────────
14
+
15
+ export interface ModelPrice {
16
+ /** Cost per 1M input tokens in USD */
17
+ inputPer1M: number;
18
+ /** Cost per 1M output tokens in USD */
19
+ outputPer1M: number;
20
+ /** Cost per 1M cache read tokens in USD (default: 0) */
21
+ cacheReadPer1M?: number;
22
+ /** Cost per 1M cache write tokens in USD (default: 0) */
23
+ cacheWritePer1M?: number;
24
+ }
25
+
26
+ export interface CostSummary {
27
+ /** Total cost in USD (null if no price data available) */
28
+ totalUSD: number | null;
29
+ /** Per-model breakdown */
30
+ models: {
31
+ model: string;
32
+ inputTokens: number;
33
+ outputTokens: number;
34
+ cacheReadTokens: number;
35
+ cacheWriteTokens: number;
36
+ costUSD: number | null;
37
+ }[];
38
+ /** Whether the budget has been exceeded */
39
+ budgetExceeded: boolean;
40
+ /** Configured budget (null if no budget) */
41
+ budgetUSD: number | null;
42
+ }
43
+
44
+ // ── Default Price Table ────────────────────────────────────────────
45
+
46
+ const DEFAULT_PRICES: Record<string, ModelPrice> = {
47
+ // Claude models — verified 2026-04-13 via Anthropic docs
48
+ 'claude-opus-4-6': { inputPer1M: 5, outputPer1M: 25, cacheReadPer1M: 0.50, cacheWritePer1M: 6.25 },
49
+ 'claude-sonnet-4-6': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
50
+ 'claude-haiku-4-5': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
51
+ 'claude-haiku-4-5-20251001': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
52
+ // Legacy Claude models (still in use by some adapters)
53
+ 'claude-sonnet-4-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
54
+ 'claude-sonnet-4-0-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
55
+ 'claude-opus-4-20250514': { inputPer1M: 15, outputPer1M: 75, cacheReadPer1M: 1.50, cacheWritePer1M: 18.75 },
56
+ 'claude-3-5-sonnet': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
57
+ 'claude-3-5-haiku': { inputPer1M: 0.80, outputPer1M: 4, cacheReadPer1M: 0.08, cacheWritePer1M: 1 },
58
+
59
+ // OpenAI / Codex models — verified 2026-04-13
60
+ 'gpt-5': { inputPer1M: 1.25, outputPer1M: 10 },
61
+ 'gpt-5-mini': { inputPer1M: 0.25, outputPer1M: 2 },
62
+ 'gpt-5-nano': { inputPer1M: 0.05, outputPer1M: 0.40 },
63
+ 'gpt-5.2': { inputPer1M: 1.75, outputPer1M: 14 },
64
+ 'o3': { inputPer1M: 2, outputPer1M: 8 },
65
+ 'o3-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
66
+ 'o4-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
67
+ 'gpt-4.1': { inputPer1M: 2, outputPer1M: 8, cacheReadPer1M: 0.50, cacheWritePer1M: 2 },
68
+ 'gpt-4.1-mini': { inputPer1M: 0.40, outputPer1M: 1.60, cacheReadPer1M: 0.10, cacheWritePer1M: 0.40 },
69
+ 'gpt-4.1-nano': { inputPer1M: 0.10, outputPer1M: 0.40 },
70
+ 'gpt-4o': { inputPer1M: 2.50, outputPer1M: 10, cacheReadPer1M: 1.25 },
71
+ 'gpt-4o-mini': { inputPer1M: 0.15, outputPer1M: 0.60, cacheReadPer1M: 0.075 },
72
+
73
+ // Google Gemini models — verified 2026-04-13 via Google AI pricing
74
+ 'gemini-2.5-flash': { inputPer1M: 0.30, outputPer1M: 2.50 },
75
+ 'gemini-2.5-flash-preview-05-20': { inputPer1M: 0.30, outputPer1M: 2.50 },
76
+ 'gemini-2.5-flash-lite': { inputPer1M: 0.10, outputPer1M: 0.40 },
77
+ 'gemini-2.5-pro': { inputPer1M: 1.25, outputPer1M: 10 },
78
+ 'gemini-2.0-flash': { inputPer1M: 0.10, outputPer1M: 0.40 },
79
+ 'gemini-2.0-flash-lite': { inputPer1M: 0.075, outputPer1M: 0.30 },
80
+ // Gemini 3.x generation
81
+ 'gemini-3.1-pro': { inputPer1M: 2, outputPer1M: 12 },
82
+ 'gemini-3-flash': { inputPer1M: 0.50, outputPer1M: 3 },
83
+
84
+ // OpenCode / open models — verified 2026-04-13
85
+ 'deepseek-coder': { inputPer1M: 0.28, outputPer1M: 0.42 },
86
+ 'deepseek-chat': { inputPer1M: 0.28, outputPer1M: 0.42 },
87
+ };
88
+
89
+ // ── Core ───────────────────────────────────────────────────────────
90
+
91
+ /**
92
+ * Calculate cost for a single model's token usage.
93
+ * Returns null if the model is not in the price table.
94
+ */
95
+ export function calculateModelCost(
96
+ model: string,
97
+ usage: TokenUsage,
98
+ customPrices?: Record<string, ModelPrice>,
99
+ ): number | null {
100
+ const prices = { ...DEFAULT_PRICES, ...customPrices };
101
+ const price = findPrice(model, prices);
102
+ if (!price) return null;
103
+
104
+ const inputCost = (usage.inputTokens / 1_000_000) * price.inputPer1M;
105
+ const outputCost = (usage.outputTokens / 1_000_000) * price.outputPer1M;
106
+ const cacheReadCost = (usage.cacheReadTokens / 1_000_000) * (price.cacheReadPer1M ?? 0);
107
+ const cacheWriteCost = (usage.cacheWriteTokens / 1_000_000) * (price.cacheWritePer1M ?? 0);
108
+
109
+ return inputCost + outputCost + cacheReadCost + cacheWriteCost;
110
+ }
111
+
112
+ /**
113
+ * Calculate total pipeline cost from accumulated token usage.
114
+ */
115
+ export function calculatePipelineCost(
116
+ tokenUsage: Record<string, TokenUsage>,
117
+ budgetUSD?: number,
118
+ customPrices?: Record<string, ModelPrice>,
119
+ ): CostSummary {
120
+ const models: CostSummary['models'] = [];
121
+ let totalUSD: number | null = 0;
122
+ let hasAnyPrice = false;
123
+
124
+ for (const [model, usage] of Object.entries(tokenUsage)) {
125
+ const cost = calculateModelCost(model, usage, customPrices);
126
+ if (cost !== null) {
127
+ hasAnyPrice = true;
128
+ totalUSD = (totalUSD ?? 0) + cost;
129
+ }
130
+
131
+ models.push({
132
+ model,
133
+ inputTokens: usage.inputTokens,
134
+ outputTokens: usage.outputTokens,
135
+ cacheReadTokens: usage.cacheReadTokens,
136
+ cacheWriteTokens: usage.cacheWriteTokens,
137
+ costUSD: cost !== null ? Math.round(cost * 10000) / 10000 : null, // Round to 4 decimals
138
+ });
139
+ }
140
+
141
+ if (!hasAnyPrice) totalUSD = null;
142
+ else totalUSD = Math.round(totalUSD! * 10000) / 10000;
143
+
144
+ return {
145
+ totalUSD,
146
+ models,
147
+ budgetExceeded: budgetUSD != null && totalUSD != null && totalUSD > budgetUSD,
148
+ budgetUSD: budgetUSD ?? null,
149
+ };
150
+ }
151
+
152
+ /**
153
+ * Check if the current cost exceeds the budget.
154
+ * Returns false if no budget is set or cost cannot be calculated.
155
+ */
156
+ export function isBudgetExceeded(
157
+ tokenUsage: Record<string, TokenUsage>,
158
+ budgetUSD?: number,
159
+ customPrices?: Record<string, ModelPrice>,
160
+ ): boolean {
161
+ if (budgetUSD == null) return false;
162
+ const summary = calculatePipelineCost(tokenUsage, budgetUSD, customPrices);
163
+ return summary.budgetExceeded;
164
+ }
165
+
166
+ /**
167
+ * Format cost summary for CLI display.
168
+ */
169
+ export function formatCostSummary(summary: CostSummary): string {
170
+ const lines: string[] = [];
171
+
172
+ for (const m of summary.models) {
173
+ const tokens = `in=${fmtNum(m.inputTokens)} out=${fmtNum(m.outputTokens)}`;
174
+ const cache = m.cacheReadTokens > 0 || m.cacheWriteTokens > 0
175
+ ? ` cache_r=${fmtNum(m.cacheReadTokens)} cache_w=${fmtNum(m.cacheWriteTokens)}`
176
+ : '';
177
+ const cost = m.costUSD !== null ? ` ($${m.costUSD.toFixed(4)})` : '';
178
+ lines.push(` ${m.model}: ${tokens}${cache}${cost}`);
179
+ }
180
+
181
+ if (summary.totalUSD !== null) {
182
+ lines.push(` Total: $${summary.totalUSD.toFixed(4)}`);
183
+ }
184
+
185
+ if (summary.budgetUSD !== null) {
186
+ const status = summary.budgetExceeded ? ' [WARN] EXCEEDED' : '';
187
+ lines.push(` Budget: $${summary.budgetUSD.toFixed(2)}${status}`);
188
+ }
189
+
190
+ return lines.join('\n');
191
+ }
192
+
193
+ // ── Helpers ────────────────────────────────────────────────────────
194
+
195
+ /**
196
+ * Find price entry by model name. Supports fuzzy matching:
197
+ * "claude-sonnet-4-20250514" matches "claude-sonnet-4-20250514" exactly,
198
+ * but also tries prefix matching for versioned models.
199
+ */
200
+ function findPrice(model: string, prices: Record<string, ModelPrice>): ModelPrice | null {
201
+ // Exact match
202
+ if (prices[model]) return prices[model];
203
+
204
+ // Prefix match: "claude-sonnet-4-20250514" → try "claude-sonnet-4"
205
+ const normalized = model.toLowerCase();
206
+ for (const [key, price] of Object.entries(prices)) {
207
+ if (normalized.startsWith(key.toLowerCase()) || key.toLowerCase().startsWith(normalized)) {
208
+ return price;
209
+ }
210
+ }
211
+
212
+ return null;
213
+ }
214
+
215
+ function fmtNum(n: number): string {
216
+ if (n >= 1_000_000) return `${(n / 1_000_000).toFixed(1)}M`;
217
+ if (n >= 1_000) return `${(n / 1_000).toFixed(1)}K`;
218
+ return String(n);
219
+ }