memorix 1.2.0 → 1.2.2

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (212) hide show
  1. package/CHANGELOG.md +30 -1
  2. package/README.md +18 -4
  3. package/README.zh-CN.md +18 -4
  4. package/TEAM.md +86 -86
  5. package/dist/cli/index.js +15919 -14055
  6. package/dist/cli/index.js.map +1 -1
  7. package/dist/index.js +1997 -1021
  8. package/dist/index.js.map +1 -1
  9. package/dist/maintenance-runner.d.ts +1 -1
  10. package/dist/maintenance-runner.js +8481 -8005
  11. package/dist/maintenance-runner.js.map +1 -1
  12. package/dist/memcode-runtime/CHANGELOG.md +30 -1
  13. package/dist/sdk.d.ts +7 -2
  14. package/dist/sdk.js +2022 -1024
  15. package/dist/sdk.js.map +1 -1
  16. package/dist/types.d.ts +49 -1
  17. package/dist/types.js.map +1 -1
  18. package/docs/1.2.2-MEMORY-CONTROL-PLANE.md +434 -0
  19. package/docs/AGENT_OPERATOR_PLAYBOOK.md +4 -0
  20. package/docs/API_REFERENCE.md +27 -5
  21. package/docs/DESIGN_DECISIONS.md +357 -357
  22. package/docs/DEVELOPMENT.md +4 -0
  23. package/docs/README.md +1 -1
  24. package/docs/SETUP.md +7 -1
  25. package/docs/dev-log/progress.txt +91 -11
  26. package/docs/knowledge/workflows/memorix-release.md +57 -0
  27. package/package.json +1 -1
  28. package/plugins/codex/memorix/.codex-plugin/plugin.json +1 -1
  29. package/src/audit/index.ts +156 -156
  30. package/src/cli/command-guide.ts +192 -0
  31. package/src/cli/commands/audit-list.ts +89 -89
  32. package/src/cli/commands/audit.ts +9 -4
  33. package/src/cli/commands/background.ts +659 -659
  34. package/src/cli/commands/cleanup.ts +5 -1
  35. package/src/cli/commands/codegraph.ts +17 -8
  36. package/src/cli/commands/context.ts +3 -2
  37. package/src/cli/commands/doctor.ts +4 -2
  38. package/src/cli/commands/explain.ts +9 -3
  39. package/src/cli/commands/formation.ts +48 -48
  40. package/src/cli/commands/git-hook-install.ts +111 -111
  41. package/src/cli/commands/handoff.ts +75 -61
  42. package/src/cli/commands/hooks-status.ts +63 -63
  43. package/src/cli/commands/identity.ts +116 -0
  44. package/src/cli/commands/ingest-commit.ts +153 -153
  45. package/src/cli/commands/ingest-image.ts +71 -69
  46. package/src/cli/commands/ingest-log.ts +180 -180
  47. package/src/cli/commands/ingest.ts +44 -44
  48. package/src/cli/commands/integrate-shared.ts +15 -15
  49. package/src/cli/commands/knowledge.ts +40 -0
  50. package/src/cli/commands/lock.ts +93 -92
  51. package/src/cli/commands/memory.ts +58 -21
  52. package/src/cli/commands/message.ts +123 -118
  53. package/src/cli/commands/operator-shared.ts +98 -3
  54. package/src/cli/commands/poll.ts +74 -64
  55. package/src/cli/commands/purge-all-memory.ts +85 -85
  56. package/src/cli/commands/purge-project-memory.ts +83 -83
  57. package/src/cli/commands/reasoning.ts +135 -121
  58. package/src/cli/commands/retention.ts +9 -4
  59. package/src/cli/commands/serve-http.ts +22 -43
  60. package/src/cli/commands/serve-shared.ts +118 -118
  61. package/src/cli/commands/session.ts +29 -3
  62. package/src/cli/commands/setup.ts +9 -3
  63. package/src/cli/commands/skills.ts +124 -119
  64. package/src/cli/commands/status.ts +4 -3
  65. package/src/cli/commands/task.ts +193 -184
  66. package/src/cli/commands/team.ts +14 -10
  67. package/src/cli/commands/transfer.ts +108 -55
  68. package/src/cli/commands/uninstall-project-artifacts.ts +85 -85
  69. package/src/cli/identity.ts +89 -0
  70. package/src/cli/index.ts +96 -19
  71. package/src/cli/invocation.ts +115 -0
  72. package/src/cli/tui/ChatView.tsx +234 -234
  73. package/src/cli/tui/CommandBar.tsx +312 -312
  74. package/src/cli/tui/ContextRail.tsx +118 -118
  75. package/src/cli/tui/HeaderBar.tsx +72 -72
  76. package/src/cli/tui/LogoBanner.tsx +51 -51
  77. package/src/cli/tui/Sidebar.tsx +179 -179
  78. package/src/cli/tui/chat-service.ts +41 -18
  79. package/src/cli/tui/data.ts +23 -44
  80. package/src/cli/tui/index.ts +41 -41
  81. package/src/cli/tui/markdown-render.tsx +371 -371
  82. package/src/cli/tui/operator-context.ts +60 -0
  83. package/src/cli/tui/use-mouse.ts +157 -157
  84. package/src/cli/tui/useNavigation.ts +56 -56
  85. package/src/cli/tui/views/MemoryView.tsx +10 -8
  86. package/src/cli/update-checker.ts +211 -211
  87. package/src/cli/version.ts +7 -7
  88. package/src/cli/workbench.ts +1 -1
  89. package/src/codegraph/auto-context.ts +34 -17
  90. package/src/codegraph/context-pack.ts +1 -0
  91. package/src/codegraph/current-facts.ts +19 -1
  92. package/src/codegraph/project-context.ts +2 -0
  93. package/src/codegraph/task-lens.ts +49 -5
  94. package/src/compact/engine.ts +26 -10
  95. package/src/compact/index-format.ts +25 -2
  96. package/src/compact/token-budget.ts +74 -74
  97. package/src/dashboard/project-classification.ts +64 -64
  98. package/src/dashboard/server.ts +58 -52
  99. package/src/embedding/fastembed-provider.ts +142 -142
  100. package/src/embedding/transformers-provider.ts +111 -111
  101. package/src/git/extractor.ts +209 -209
  102. package/src/git/hooks-path.ts +85 -85
  103. package/src/hooks/admission.ts +117 -0
  104. package/src/hooks/handler.ts +98 -91
  105. package/src/hooks/pattern-detector.ts +173 -173
  106. package/src/hooks/significance-filter.ts +250 -250
  107. package/src/knowledge/claims.ts +51 -1
  108. package/src/knowledge/context-assembly.ts +97 -0
  109. package/src/knowledge/types.ts +1 -0
  110. package/src/knowledge/workflows.ts +34 -3
  111. package/src/knowledge/workset.ts +179 -10
  112. package/src/llm/memory-manager.ts +328 -328
  113. package/src/llm/provider.ts +885 -885
  114. package/src/llm/quality.ts +248 -248
  115. package/src/memory/admission.ts +57 -0
  116. package/src/memory/attribution-guard.ts +249 -249
  117. package/src/memory/auto-relations.ts +21 -0
  118. package/src/memory/consolidation.ts +13 -2
  119. package/src/memory/disclosure-policy.ts +140 -135
  120. package/src/memory/entity-extractor.ts +197 -197
  121. package/src/memory/export-import.ts +11 -3
  122. package/src/memory/formation/evaluate.ts +217 -217
  123. package/src/memory/formation/extract.ts +361 -361
  124. package/src/memory/formation/index.ts +417 -417
  125. package/src/memory/formation/resolve.ts +344 -344
  126. package/src/memory/formation/types.ts +315 -315
  127. package/src/memory/freshness.ts +122 -122
  128. package/src/memory/graph-context.ts +8 -2
  129. package/src/memory/graph-scope.ts +46 -0
  130. package/src/memory/graph.ts +197 -197
  131. package/src/memory/observations.ts +162 -4
  132. package/src/memory/quality-audit.ts +2 -0
  133. package/src/memory/refs.ts +94 -94
  134. package/src/memory/retention.ts +22 -2
  135. package/src/memory/secret-filter.ts +79 -79
  136. package/src/memory/session.ts +5 -2
  137. package/src/memory/visibility.ts +80 -0
  138. package/src/multimodal/image-loader.ts +143 -143
  139. package/src/orchestrate/adapters/claude-stream.ts +192 -192
  140. package/src/orchestrate/adapters/claude.ts +111 -111
  141. package/src/orchestrate/adapters/codex-stream.ts +134 -134
  142. package/src/orchestrate/adapters/codex.ts +41 -41
  143. package/src/orchestrate/adapters/gemini-stream.ts +166 -166
  144. package/src/orchestrate/adapters/gemini.ts +42 -42
  145. package/src/orchestrate/adapters/index.ts +73 -73
  146. package/src/orchestrate/adapters/opencode-stream.ts +143 -143
  147. package/src/orchestrate/adapters/opencode.ts +47 -47
  148. package/src/orchestrate/adapters/spawn-helper.ts +286 -286
  149. package/src/orchestrate/adapters/types.ts +77 -77
  150. package/src/orchestrate/capability-router.ts +284 -284
  151. package/src/orchestrate/context-compact.ts +188 -188
  152. package/src/orchestrate/cost-tracker.ts +219 -219
  153. package/src/orchestrate/error-recovery.ts +191 -191
  154. package/src/orchestrate/evidence.ts +140 -140
  155. package/src/orchestrate/ledger.ts +110 -110
  156. package/src/orchestrate/memorix-bridge.ts +378 -340
  157. package/src/orchestrate/output-budget.ts +80 -80
  158. package/src/orchestrate/permission.ts +152 -152
  159. package/src/orchestrate/pipeline-trace.ts +131 -131
  160. package/src/orchestrate/prompt-builder.ts +155 -155
  161. package/src/orchestrate/ring-buffer.ts +37 -37
  162. package/src/orchestrate/task-graph.ts +389 -389
  163. package/src/orchestrate/verify-gate.ts +33 -10
  164. package/src/orchestrate/worktree.ts +232 -232
  165. package/src/project/aliases.ts +374 -374
  166. package/src/project/detector.ts +268 -268
  167. package/src/rules/adapters/claude-code.ts +99 -99
  168. package/src/rules/adapters/codex.ts +97 -97
  169. package/src/rules/adapters/copilot.ts +124 -124
  170. package/src/rules/adapters/cursor.ts +114 -114
  171. package/src/rules/adapters/kiro.ts +126 -126
  172. package/src/rules/adapters/trae.ts +56 -56
  173. package/src/rules/adapters/windsurf.ts +83 -83
  174. package/src/rules/syncer.ts +235 -235
  175. package/src/runtime/control-plane-maintenance.ts +1 -0
  176. package/src/runtime/isolated-maintenance.ts +1 -0
  177. package/src/runtime/lifecycle.ts +18 -0
  178. package/src/runtime/maintenance-jobs.ts +1 -0
  179. package/src/runtime/maintenance-runner.ts +2 -0
  180. package/src/runtime/project-maintenance.ts +89 -0
  181. package/src/sdk.ts +334 -304
  182. package/src/search/intent-detector.ts +289 -289
  183. package/src/search/query-expansion.ts +52 -52
  184. package/src/server/formation-timeout.ts +27 -27
  185. package/src/server.ts +334 -93
  186. package/src/skills/mini-skills.ts +386 -386
  187. package/src/store/chat-store.ts +119 -119
  188. package/src/store/graph-store.ts +249 -249
  189. package/src/store/mini-skill-store.ts +349 -349
  190. package/src/store/orama-store.ts +61 -6
  191. package/src/store/persistence-json.ts +212 -212
  192. package/src/store/persistence.ts +291 -291
  193. package/src/store/project-affinity.ts +195 -195
  194. package/src/store/sqlite-db.ts +23 -1
  195. package/src/store/sqlite-store.ts +12 -2
  196. package/src/team/event-bus.ts +76 -76
  197. package/src/team/file-locks.ts +173 -173
  198. package/src/team/handoff.ts +168 -161
  199. package/src/team/messages.ts +203 -203
  200. package/src/team/poll.ts +132 -132
  201. package/src/team/tasks.ts +211 -211
  202. package/src/types.ts +51 -0
  203. package/src/wiki/generator.ts +2 -0
  204. package/src/workspace/mcp-adapters/codex.ts +191 -191
  205. package/src/workspace/mcp-adapters/copilot.ts +105 -105
  206. package/src/workspace/mcp-adapters/cursor.ts +53 -53
  207. package/src/workspace/mcp-adapters/kiro.ts +64 -64
  208. package/src/workspace/mcp-adapters/opencode.ts +123 -123
  209. package/src/workspace/mcp-adapters/trae.ts +134 -134
  210. package/src/workspace/mcp-adapters/windsurf.ts +91 -91
  211. package/src/workspace/sanitizer.ts +60 -60
  212. package/src/workspace/workflow-sync.ts +131 -131
@@ -1,219 +1,219 @@
1
- /**
2
- * Cost Tracker — Phase 7, Step 8: Token → USD conversion and budget enforcement.
3
- *
4
- * Converts accumulated token usage (per model) into USD using a configurable
5
- * price table. Supports budget limits that abort the pipeline when exceeded.
6
- *
7
- * Design principle: if the price table doesn't have a model, report tokens only
8
- * (skip USD calculation). Never crash because of a missing price entry.
9
- */
10
-
11
- import type { TokenUsage } from './adapters/types.js';
12
-
13
- // ── Types ──────────────────────────────────────────────────────────
14
-
15
- export interface ModelPrice {
16
- /** Cost per 1M input tokens in USD */
17
- inputPer1M: number;
18
- /** Cost per 1M output tokens in USD */
19
- outputPer1M: number;
20
- /** Cost per 1M cache read tokens in USD (default: 0) */
21
- cacheReadPer1M?: number;
22
- /** Cost per 1M cache write tokens in USD (default: 0) */
23
- cacheWritePer1M?: number;
24
- }
25
-
26
- export interface CostSummary {
27
- /** Total cost in USD (null if no price data available) */
28
- totalUSD: number | null;
29
- /** Per-model breakdown */
30
- models: {
31
- model: string;
32
- inputTokens: number;
33
- outputTokens: number;
34
- cacheReadTokens: number;
35
- cacheWriteTokens: number;
36
- costUSD: number | null;
37
- }[];
38
- /** Whether the budget has been exceeded */
39
- budgetExceeded: boolean;
40
- /** Configured budget (null if no budget) */
41
- budgetUSD: number | null;
42
- }
43
-
44
- // ── Default Price Table ────────────────────────────────────────────
45
-
46
- const DEFAULT_PRICES: Record<string, ModelPrice> = {
47
- // Claude models — verified 2026-04-13 via Anthropic docs
48
- 'claude-opus-4-6': { inputPer1M: 5, outputPer1M: 25, cacheReadPer1M: 0.50, cacheWritePer1M: 6.25 },
49
- 'claude-sonnet-4-6': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
50
- 'claude-haiku-4-5': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
51
- 'claude-haiku-4-5-20251001': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
52
- // Legacy Claude models (still in use by some adapters)
53
- 'claude-sonnet-4-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
54
- 'claude-sonnet-4-0-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
55
- 'claude-opus-4-20250514': { inputPer1M: 15, outputPer1M: 75, cacheReadPer1M: 1.50, cacheWritePer1M: 18.75 },
56
- 'claude-3-5-sonnet': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
57
- 'claude-3-5-haiku': { inputPer1M: 0.80, outputPer1M: 4, cacheReadPer1M: 0.08, cacheWritePer1M: 1 },
58
-
59
- // OpenAI / Codex models — verified 2026-04-13
60
- 'gpt-5': { inputPer1M: 1.25, outputPer1M: 10 },
61
- 'gpt-5-mini': { inputPer1M: 0.25, outputPer1M: 2 },
62
- 'gpt-5-nano': { inputPer1M: 0.05, outputPer1M: 0.40 },
63
- 'gpt-5.2': { inputPer1M: 1.75, outputPer1M: 14 },
64
- 'o3': { inputPer1M: 2, outputPer1M: 8 },
65
- 'o3-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
66
- 'o4-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
67
- 'gpt-4.1': { inputPer1M: 2, outputPer1M: 8, cacheReadPer1M: 0.50, cacheWritePer1M: 2 },
68
- 'gpt-4.1-mini': { inputPer1M: 0.40, outputPer1M: 1.60, cacheReadPer1M: 0.10, cacheWritePer1M: 0.40 },
69
- 'gpt-4.1-nano': { inputPer1M: 0.10, outputPer1M: 0.40 },
70
- 'gpt-4o': { inputPer1M: 2.50, outputPer1M: 10, cacheReadPer1M: 1.25 },
71
- 'gpt-4o-mini': { inputPer1M: 0.15, outputPer1M: 0.60, cacheReadPer1M: 0.075 },
72
-
73
- // Google Gemini models — verified 2026-04-13 via Google AI pricing
74
- 'gemini-2.5-flash': { inputPer1M: 0.30, outputPer1M: 2.50 },
75
- 'gemini-2.5-flash-preview-05-20': { inputPer1M: 0.30, outputPer1M: 2.50 },
76
- 'gemini-2.5-flash-lite': { inputPer1M: 0.10, outputPer1M: 0.40 },
77
- 'gemini-2.5-pro': { inputPer1M: 1.25, outputPer1M: 10 },
78
- 'gemini-2.0-flash': { inputPer1M: 0.10, outputPer1M: 0.40 },
79
- 'gemini-2.0-flash-lite': { inputPer1M: 0.075, outputPer1M: 0.30 },
80
- // Gemini 3.x generation
81
- 'gemini-3.1-pro': { inputPer1M: 2, outputPer1M: 12 },
82
- 'gemini-3-flash': { inputPer1M: 0.50, outputPer1M: 3 },
83
-
84
- // OpenCode / open models — verified 2026-04-13
85
- 'deepseek-coder': { inputPer1M: 0.28, outputPer1M: 0.42 },
86
- 'deepseek-chat': { inputPer1M: 0.28, outputPer1M: 0.42 },
87
- };
88
-
89
- // ── Core ───────────────────────────────────────────────────────────
90
-
91
- /**
92
- * Calculate cost for a single model's token usage.
93
- * Returns null if the model is not in the price table.
94
- */
95
- export function calculateModelCost(
96
- model: string,
97
- usage: TokenUsage,
98
- customPrices?: Record<string, ModelPrice>,
99
- ): number | null {
100
- const prices = { ...DEFAULT_PRICES, ...customPrices };
101
- const price = findPrice(model, prices);
102
- if (!price) return null;
103
-
104
- const inputCost = (usage.inputTokens / 1_000_000) * price.inputPer1M;
105
- const outputCost = (usage.outputTokens / 1_000_000) * price.outputPer1M;
106
- const cacheReadCost = (usage.cacheReadTokens / 1_000_000) * (price.cacheReadPer1M ?? 0);
107
- const cacheWriteCost = (usage.cacheWriteTokens / 1_000_000) * (price.cacheWritePer1M ?? 0);
108
-
109
- return inputCost + outputCost + cacheReadCost + cacheWriteCost;
110
- }
111
-
112
- /**
113
- * Calculate total pipeline cost from accumulated token usage.
114
- */
115
- export function calculatePipelineCost(
116
- tokenUsage: Record<string, TokenUsage>,
117
- budgetUSD?: number,
118
- customPrices?: Record<string, ModelPrice>,
119
- ): CostSummary {
120
- const models: CostSummary['models'] = [];
121
- let totalUSD: number | null = 0;
122
- let hasAnyPrice = false;
123
-
124
- for (const [model, usage] of Object.entries(tokenUsage)) {
125
- const cost = calculateModelCost(model, usage, customPrices);
126
- if (cost !== null) {
127
- hasAnyPrice = true;
128
- totalUSD = (totalUSD ?? 0) + cost;
129
- }
130
-
131
- models.push({
132
- model,
133
- inputTokens: usage.inputTokens,
134
- outputTokens: usage.outputTokens,
135
- cacheReadTokens: usage.cacheReadTokens,
136
- cacheWriteTokens: usage.cacheWriteTokens,
137
- costUSD: cost !== null ? Math.round(cost * 10000) / 10000 : null, // Round to 4 decimals
138
- });
139
- }
140
-
141
- if (!hasAnyPrice) totalUSD = null;
142
- else totalUSD = Math.round(totalUSD! * 10000) / 10000;
143
-
144
- return {
145
- totalUSD,
146
- models,
147
- budgetExceeded: budgetUSD != null && totalUSD != null && totalUSD > budgetUSD,
148
- budgetUSD: budgetUSD ?? null,
149
- };
150
- }
151
-
152
- /**
153
- * Check if the current cost exceeds the budget.
154
- * Returns false if no budget is set or cost cannot be calculated.
155
- */
156
- export function isBudgetExceeded(
157
- tokenUsage: Record<string, TokenUsage>,
158
- budgetUSD?: number,
159
- customPrices?: Record<string, ModelPrice>,
160
- ): boolean {
161
- if (budgetUSD == null) return false;
162
- const summary = calculatePipelineCost(tokenUsage, budgetUSD, customPrices);
163
- return summary.budgetExceeded;
164
- }
165
-
166
- /**
167
- * Format cost summary for CLI display.
168
- */
169
- export function formatCostSummary(summary: CostSummary): string {
170
- const lines: string[] = [];
171
-
172
- for (const m of summary.models) {
173
- const tokens = `in=${fmtNum(m.inputTokens)} out=${fmtNum(m.outputTokens)}`;
174
- const cache = m.cacheReadTokens > 0 || m.cacheWriteTokens > 0
175
- ? ` cache_r=${fmtNum(m.cacheReadTokens)} cache_w=${fmtNum(m.cacheWriteTokens)}`
176
- : '';
177
- const cost = m.costUSD !== null ? ` ($${m.costUSD.toFixed(4)})` : '';
178
- lines.push(` ${m.model}: ${tokens}${cache}${cost}`);
179
- }
180
-
181
- if (summary.totalUSD !== null) {
182
- lines.push(` Total: $${summary.totalUSD.toFixed(4)}`);
183
- }
184
-
185
- if (summary.budgetUSD !== null) {
186
- const status = summary.budgetExceeded ? ' [WARN] EXCEEDED' : '';
187
- lines.push(` Budget: $${summary.budgetUSD.toFixed(2)}${status}`);
188
- }
189
-
190
- return lines.join('\n');
191
- }
192
-
193
- // ── Helpers ────────────────────────────────────────────────────────
194
-
195
- /**
196
- * Find price entry by model name. Supports fuzzy matching:
197
- * "claude-sonnet-4-20250514" matches "claude-sonnet-4-20250514" exactly,
198
- * but also tries prefix matching for versioned models.
199
- */
200
- function findPrice(model: string, prices: Record<string, ModelPrice>): ModelPrice | null {
201
- // Exact match
202
- if (prices[model]) return prices[model];
203
-
204
- // Prefix match: "claude-sonnet-4-20250514" → try "claude-sonnet-4"
205
- const normalized = model.toLowerCase();
206
- for (const [key, price] of Object.entries(prices)) {
207
- if (normalized.startsWith(key.toLowerCase()) || key.toLowerCase().startsWith(normalized)) {
208
- return price;
209
- }
210
- }
211
-
212
- return null;
213
- }
214
-
215
- function fmtNum(n: number): string {
216
- if (n >= 1_000_000) return `${(n / 1_000_000).toFixed(1)}M`;
217
- if (n >= 1_000) return `${(n / 1_000).toFixed(1)}K`;
218
- return String(n);
219
- }
1
+ /**
2
+ * Cost Tracker — Phase 7, Step 8: Token → USD conversion and budget enforcement.
3
+ *
4
+ * Converts accumulated token usage (per model) into USD using a configurable
5
+ * price table. Supports budget limits that abort the pipeline when exceeded.
6
+ *
7
+ * Design principle: if the price table doesn't have a model, report tokens only
8
+ * (skip USD calculation). Never crash because of a missing price entry.
9
+ */
10
+
11
+ import type { TokenUsage } from './adapters/types.js';
12
+
13
+ // ── Types ──────────────────────────────────────────────────────────
14
+
15
+ export interface ModelPrice {
16
+ /** Cost per 1M input tokens in USD */
17
+ inputPer1M: number;
18
+ /** Cost per 1M output tokens in USD */
19
+ outputPer1M: number;
20
+ /** Cost per 1M cache read tokens in USD (default: 0) */
21
+ cacheReadPer1M?: number;
22
+ /** Cost per 1M cache write tokens in USD (default: 0) */
23
+ cacheWritePer1M?: number;
24
+ }
25
+
26
+ export interface CostSummary {
27
+ /** Total cost in USD (null if no price data available) */
28
+ totalUSD: number | null;
29
+ /** Per-model breakdown */
30
+ models: {
31
+ model: string;
32
+ inputTokens: number;
33
+ outputTokens: number;
34
+ cacheReadTokens: number;
35
+ cacheWriteTokens: number;
36
+ costUSD: number | null;
37
+ }[];
38
+ /** Whether the budget has been exceeded */
39
+ budgetExceeded: boolean;
40
+ /** Configured budget (null if no budget) */
41
+ budgetUSD: number | null;
42
+ }
43
+
44
+ // ── Default Price Table ────────────────────────────────────────────
45
+
46
+ const DEFAULT_PRICES: Record<string, ModelPrice> = {
47
+ // Claude models — verified 2026-04-13 via Anthropic docs
48
+ 'claude-opus-4-6': { inputPer1M: 5, outputPer1M: 25, cacheReadPer1M: 0.50, cacheWritePer1M: 6.25 },
49
+ 'claude-sonnet-4-6': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
50
+ 'claude-haiku-4-5': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
51
+ 'claude-haiku-4-5-20251001': { inputPer1M: 1, outputPer1M: 5, cacheReadPer1M: 0.10, cacheWritePer1M: 1.25 },
52
+ // Legacy Claude models (still in use by some adapters)
53
+ 'claude-sonnet-4-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
54
+ 'claude-sonnet-4-0-20250514': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
55
+ 'claude-opus-4-20250514': { inputPer1M: 15, outputPer1M: 75, cacheReadPer1M: 1.50, cacheWritePer1M: 18.75 },
56
+ 'claude-3-5-sonnet': { inputPer1M: 3, outputPer1M: 15, cacheReadPer1M: 0.30, cacheWritePer1M: 3.75 },
57
+ 'claude-3-5-haiku': { inputPer1M: 0.80, outputPer1M: 4, cacheReadPer1M: 0.08, cacheWritePer1M: 1 },
58
+
59
+ // OpenAI / Codex models — verified 2026-04-13
60
+ 'gpt-5': { inputPer1M: 1.25, outputPer1M: 10 },
61
+ 'gpt-5-mini': { inputPer1M: 0.25, outputPer1M: 2 },
62
+ 'gpt-5-nano': { inputPer1M: 0.05, outputPer1M: 0.40 },
63
+ 'gpt-5.2': { inputPer1M: 1.75, outputPer1M: 14 },
64
+ 'o3': { inputPer1M: 2, outputPer1M: 8 },
65
+ 'o3-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
66
+ 'o4-mini': { inputPer1M: 1.10, outputPer1M: 4.40 },
67
+ 'gpt-4.1': { inputPer1M: 2, outputPer1M: 8, cacheReadPer1M: 0.50, cacheWritePer1M: 2 },
68
+ 'gpt-4.1-mini': { inputPer1M: 0.40, outputPer1M: 1.60, cacheReadPer1M: 0.10, cacheWritePer1M: 0.40 },
69
+ 'gpt-4.1-nano': { inputPer1M: 0.10, outputPer1M: 0.40 },
70
+ 'gpt-4o': { inputPer1M: 2.50, outputPer1M: 10, cacheReadPer1M: 1.25 },
71
+ 'gpt-4o-mini': { inputPer1M: 0.15, outputPer1M: 0.60, cacheReadPer1M: 0.075 },
72
+
73
+ // Google Gemini models — verified 2026-04-13 via Google AI pricing
74
+ 'gemini-2.5-flash': { inputPer1M: 0.30, outputPer1M: 2.50 },
75
+ 'gemini-2.5-flash-preview-05-20': { inputPer1M: 0.30, outputPer1M: 2.50 },
76
+ 'gemini-2.5-flash-lite': { inputPer1M: 0.10, outputPer1M: 0.40 },
77
+ 'gemini-2.5-pro': { inputPer1M: 1.25, outputPer1M: 10 },
78
+ 'gemini-2.0-flash': { inputPer1M: 0.10, outputPer1M: 0.40 },
79
+ 'gemini-2.0-flash-lite': { inputPer1M: 0.075, outputPer1M: 0.30 },
80
+ // Gemini 3.x generation
81
+ 'gemini-3.1-pro': { inputPer1M: 2, outputPer1M: 12 },
82
+ 'gemini-3-flash': { inputPer1M: 0.50, outputPer1M: 3 },
83
+
84
+ // OpenCode / open models — verified 2026-04-13
85
+ 'deepseek-coder': { inputPer1M: 0.28, outputPer1M: 0.42 },
86
+ 'deepseek-chat': { inputPer1M: 0.28, outputPer1M: 0.42 },
87
+ };
88
+
89
+ // ── Core ───────────────────────────────────────────────────────────
90
+
91
+ /**
92
+ * Calculate cost for a single model's token usage.
93
+ * Returns null if the model is not in the price table.
94
+ */
95
+ export function calculateModelCost(
96
+ model: string,
97
+ usage: TokenUsage,
98
+ customPrices?: Record<string, ModelPrice>,
99
+ ): number | null {
100
+ const prices = { ...DEFAULT_PRICES, ...customPrices };
101
+ const price = findPrice(model, prices);
102
+ if (!price) return null;
103
+
104
+ const inputCost = (usage.inputTokens / 1_000_000) * price.inputPer1M;
105
+ const outputCost = (usage.outputTokens / 1_000_000) * price.outputPer1M;
106
+ const cacheReadCost = (usage.cacheReadTokens / 1_000_000) * (price.cacheReadPer1M ?? 0);
107
+ const cacheWriteCost = (usage.cacheWriteTokens / 1_000_000) * (price.cacheWritePer1M ?? 0);
108
+
109
+ return inputCost + outputCost + cacheReadCost + cacheWriteCost;
110
+ }
111
+
112
+ /**
113
+ * Calculate total pipeline cost from accumulated token usage.
114
+ */
115
+ export function calculatePipelineCost(
116
+ tokenUsage: Record<string, TokenUsage>,
117
+ budgetUSD?: number,
118
+ customPrices?: Record<string, ModelPrice>,
119
+ ): CostSummary {
120
+ const models: CostSummary['models'] = [];
121
+ let totalUSD: number | null = 0;
122
+ let hasAnyPrice = false;
123
+
124
+ for (const [model, usage] of Object.entries(tokenUsage)) {
125
+ const cost = calculateModelCost(model, usage, customPrices);
126
+ if (cost !== null) {
127
+ hasAnyPrice = true;
128
+ totalUSD = (totalUSD ?? 0) + cost;
129
+ }
130
+
131
+ models.push({
132
+ model,
133
+ inputTokens: usage.inputTokens,
134
+ outputTokens: usage.outputTokens,
135
+ cacheReadTokens: usage.cacheReadTokens,
136
+ cacheWriteTokens: usage.cacheWriteTokens,
137
+ costUSD: cost !== null ? Math.round(cost * 10000) / 10000 : null, // Round to 4 decimals
138
+ });
139
+ }
140
+
141
+ if (!hasAnyPrice) totalUSD = null;
142
+ else totalUSD = Math.round(totalUSD! * 10000) / 10000;
143
+
144
+ return {
145
+ totalUSD,
146
+ models,
147
+ budgetExceeded: budgetUSD != null && totalUSD != null && totalUSD > budgetUSD,
148
+ budgetUSD: budgetUSD ?? null,
149
+ };
150
+ }
151
+
152
+ /**
153
+ * Check if the current cost exceeds the budget.
154
+ * Returns false if no budget is set or cost cannot be calculated.
155
+ */
156
+ export function isBudgetExceeded(
157
+ tokenUsage: Record<string, TokenUsage>,
158
+ budgetUSD?: number,
159
+ customPrices?: Record<string, ModelPrice>,
160
+ ): boolean {
161
+ if (budgetUSD == null) return false;
162
+ const summary = calculatePipelineCost(tokenUsage, budgetUSD, customPrices);
163
+ return summary.budgetExceeded;
164
+ }
165
+
166
+ /**
167
+ * Format cost summary for CLI display.
168
+ */
169
+ export function formatCostSummary(summary: CostSummary): string {
170
+ const lines: string[] = [];
171
+
172
+ for (const m of summary.models) {
173
+ const tokens = `in=${fmtNum(m.inputTokens)} out=${fmtNum(m.outputTokens)}`;
174
+ const cache = m.cacheReadTokens > 0 || m.cacheWriteTokens > 0
175
+ ? ` cache_r=${fmtNum(m.cacheReadTokens)} cache_w=${fmtNum(m.cacheWriteTokens)}`
176
+ : '';
177
+ const cost = m.costUSD !== null ? ` ($${m.costUSD.toFixed(4)})` : '';
178
+ lines.push(` ${m.model}: ${tokens}${cache}${cost}`);
179
+ }
180
+
181
+ if (summary.totalUSD !== null) {
182
+ lines.push(` Total: $${summary.totalUSD.toFixed(4)}`);
183
+ }
184
+
185
+ if (summary.budgetUSD !== null) {
186
+ const status = summary.budgetExceeded ? ' [WARN] EXCEEDED' : '';
187
+ lines.push(` Budget: $${summary.budgetUSD.toFixed(2)}${status}`);
188
+ }
189
+
190
+ return lines.join('\n');
191
+ }
192
+
193
+ // ── Helpers ────────────────────────────────────────────────────────
194
+
195
+ /**
196
+ * Find price entry by model name. Supports fuzzy matching:
197
+ * "claude-sonnet-4-20250514" matches "claude-sonnet-4-20250514" exactly,
198
+ * but also tries prefix matching for versioned models.
199
+ */
200
+ function findPrice(model: string, prices: Record<string, ModelPrice>): ModelPrice | null {
201
+ // Exact match
202
+ if (prices[model]) return prices[model];
203
+
204
+ // Prefix match: "claude-sonnet-4-20250514" → try "claude-sonnet-4"
205
+ const normalized = model.toLowerCase();
206
+ for (const [key, price] of Object.entries(prices)) {
207
+ if (normalized.startsWith(key.toLowerCase()) || key.toLowerCase().startsWith(normalized)) {
208
+ return price;
209
+ }
210
+ }
211
+
212
+ return null;
213
+ }
214
+
215
+ function fmtNum(n: number): string {
216
+ if (n >= 1_000_000) return `${(n / 1_000_000).toFixed(1)}M`;
217
+ if (n >= 1_000) return `${(n / 1_000).toFixed(1)}K`;
218
+ return String(n);
219
+ }