tokentracker-cli 0.84.3 → 0.84.5

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (90) hide show
  1. package/dashboard/dist/assets/{AchievementBadge-DT_Bt15h.js → AchievementBadge-DJvRfvu1.js} +1 -1
  2. package/dashboard/dist/assets/{AchievementsPage-BcfcEclS.js → AchievementsPage-Cr8WToxS.js} +2 -2
  3. package/dashboard/dist/assets/{AchievementsSection-BJpLEpRy.js → AchievementsSection-BCPYB46e.js} +1 -1
  4. package/dashboard/dist/assets/{BadgeDetailModal-CmlRAx0w.js → BadgeDetailModal-D-12Ek-H.js} +1 -1
  5. package/dashboard/dist/assets/{Card-D81iMbAy.js → Card-Ha3x9zIj.js} +1 -1
  6. package/dashboard/dist/assets/{ClawdAnimated-Ck6LE5S4.js → ClawdAnimated-DoGgAppT.js} +1 -1
  7. package/dashboard/dist/assets/{CommandPalette-BOy-RMrZ.js → CommandPalette-DIKocEAD.js} +1 -1
  8. package/dashboard/dist/assets/{CommunityStatsModal-FjeLcdqJ.js → CommunityStatsModal-BNuMbEWy.js} +1 -1
  9. package/dashboard/dist/assets/{ConfirmModal-CZ-NaD-D.js → ConfirmModal-F2gPIIS9.js} +1 -1
  10. package/dashboard/dist/assets/{DashboardPage-CRe207LA.js → DashboardPage-YvoCEGG9.js} +1 -1
  11. package/dashboard/dist/assets/{DevicePage-7Yt97g2k.js → DevicePage-CbC_J22s.js} +1 -1
  12. package/dashboard/dist/assets/{DialogClose-DJds92_A.js → DialogClose-VEynQCYg.js} +1 -1
  13. package/dashboard/dist/assets/DialogDescription-B9bb9Ar0.js +1 -0
  14. package/dashboard/dist/assets/{DialogTitle-CnIlDBWJ.js → DialogTitle-BbAVHI_v.js} +1 -1
  15. package/dashboard/dist/assets/{FadeIn-Bn_RaWs5.js → FadeIn-CZVqSnis.js} +1 -1
  16. package/dashboard/dist/assets/{HeaderGithubStar-oL5g8jVh.js → HeaderGithubStar-Cs_h332Y.js} +1 -1
  17. package/dashboard/dist/assets/{HoverTooltip-CzzS-IxA.js → HoverTooltip-DvvHTphK.js} +1 -1
  18. package/dashboard/dist/assets/{IpCheckPage-AyB_AJ8J.js → IpCheckPage-DPVmqSS_.js} +1 -1
  19. package/dashboard/dist/assets/{LandingPage-gKfdYwuu.js → LandingPage-BCvNh3zE.js} +2 -2
  20. package/dashboard/dist/assets/{LeaderboardAvatar-BvyNPxvh.js → LeaderboardAvatar-BAPZ2ROX.js} +1 -1
  21. package/dashboard/dist/assets/{LeaderboardPage-DOeS9D8n.js → LeaderboardPage-DZSxazHP.js} +2 -2
  22. package/dashboard/dist/assets/{LeaderboardProfileModal-C-Aqp1Gb.js → LeaderboardProfileModal-r03akfe2.js} +3 -3
  23. package/dashboard/dist/assets/{LeaderboardProfilePage-DEMKFbtF.js → LeaderboardProfilePage-CU63YW2M.js} +1 -1
  24. package/dashboard/dist/assets/{LimitsPage-C53tHDde.js → LimitsPage-JWZDa5_Y.js} +1 -1
  25. package/dashboard/dist/assets/{LocalOnlyNotice-D54HqxLG.js → LocalOnlyNotice-BHlcPCZk.js} +1 -1
  26. package/dashboard/dist/assets/{LoginCard-oPBw7LXy.js → LoginCard-C-dqUumR.js} +1 -1
  27. package/dashboard/dist/assets/{LoginModal-Dz0KZMxm.js → LoginModal-BFzRg4Iy.js} +1 -1
  28. package/dashboard/dist/assets/{LoginPage-DGRnQYze.js → LoginPage-QtlnGF1S.js} +1 -1
  29. package/dashboard/dist/assets/{PetPage-CgVLe9qk.js → PetPage-C6UMY2qD.js} +1 -1
  30. package/dashboard/dist/assets/{PopoverPopup-Ble3SNoD.js → PopoverPopup-CltmqcrT.js} +1 -1
  31. package/dashboard/dist/assets/{ProviderIcon-PRELoSZw.js → ProviderIcon-BqOI9EiD.js} +1 -1
  32. package/dashboard/dist/assets/{ResetPasswordPage-Es-ZQw-9.js → ResetPasswordPage-BsfCBUZD.js} +1 -1
  33. package/dashboard/dist/assets/{SegmentedControl-6SpvB0j6.js → SegmentedControl-D4biXw_Q.js} +1 -1
  34. package/dashboard/dist/assets/{Select-DheXInDh.js → Select-DTw-14iE.js} +1 -1
  35. package/dashboard/dist/assets/{SelectItemText-yYHWlZTF.js → SelectItemText-3tUrnu3f.js} +1 -1
  36. package/dashboard/dist/assets/{SessionsPage-a7OVoq4M.js → SessionsPage-DYYlNqHg.js} +1 -1
  37. package/dashboard/dist/assets/{SettingsPage-BDbj6b4P.js → SettingsPage-CYlEboSl.js} +1 -1
  38. package/dashboard/dist/assets/{SkillsPage-BPVMDPgz.js → SkillsPage-Bzq8V5Pd.js} +1 -1
  39. package/dashboard/dist/assets/{TokenGalaxy-C1Q39N9y.js → TokenGalaxy-DU77e8sM.js} +1 -1
  40. package/dashboard/dist/assets/{ToolbarRootContext-D9aM7lug.js → ToolbarRootContext-DV3KGf-a.js} +1 -1
  41. package/dashboard/dist/assets/{TrendMonitor-CY_h1Ri-.js → TrendMonitor-D6QYKmbl.js} +1 -1
  42. package/dashboard/dist/assets/{WidgetsPage-Bz-0mZAb.js → WidgetsPage-DSP4jCsH.js} +1 -1
  43. package/dashboard/dist/assets/{WrappedPage-BQfwpPEN.js → WrappedPage-B0x6O_DU.js} +1 -1
  44. package/dashboard/dist/assets/{__vite-browser-external-CETNCSn1.js → __vite-browser-external-D3llGq6p.js} +2 -2
  45. package/dashboard/dist/assets/{agent-logos-B78kWhvN.js → agent-logos-B2S_cdOL.js} +1 -1
  46. package/dashboard/dist/assets/{arrow-up-right-CV1xfEXw.js → arrow-up-right-DNUVuhx2.js} +1 -1
  47. package/dashboard/dist/assets/{check-DX9GQFI8.js → check-C3ARiiE_.js} +1 -1
  48. package/dashboard/dist/assets/{chevron-right-DliGgkYe.js → chevron-right-3bY7H8jO.js} +1 -1
  49. package/dashboard/dist/assets/{copy-BRVpNdYi.js → copy-C-UIMKoO.js} +1 -1
  50. package/dashboard/dist/assets/{download-DOY8JNGM.js → download-DtRzJfcr.js} +1 -1
  51. package/dashboard/dist/assets/{flame-D4d0bEEi.js → flame-Bas3ry-g.js} +1 -1
  52. package/dashboard/dist/assets/{format-tokens-DuG3QtIM.js → format-tokens-DFtRUIH8.js} +1 -1
  53. package/dashboard/dist/assets/{icons-BXUnN55n.js → icons-CHBt4Tmo.js} +1 -1
  54. package/dashboard/dist/assets/{index-D-gLs38W.js → index-4unZlm2w.js} +1 -1
  55. package/dashboard/dist/assets/{index-B5GFyGve.js → index-C1ADnvzm.js} +1 -1
  56. package/dashboard/dist/assets/{index-B5l9g9A4.js → index-CZRlrK7w.js} +1 -1
  57. package/dashboard/dist/assets/{info-BdpahGob.js → info-BkrO-fg3.js} +1 -1
  58. package/dashboard/dist/assets/{limits-providers-DNBMhmgr.js → limits-providers-BtwYmYHA.js} +1 -1
  59. package/dashboard/dist/assets/{link-2-B9ja2SoD.js → link-2-DTLFz8uG.js} +1 -1
  60. package/dashboard/dist/assets/{loader-circle-eLo_WTSz.js → loader-circle-BqQW_Ipq.js} +1 -1
  61. package/dashboard/dist/assets/{maximize-2-DpU-6LJp.js → maximize-2-YdpeE4dq.js} +1 -1
  62. package/dashboard/dist/assets/{provider-display-BFExlP4i.js → provider-display-CgBOpjq0.js} +1 -1
  63. package/dashboard/dist/assets/{react-D2-wZCT7.js → react-DD9jpv--.js} +1 -1
  64. package/dashboard/dist/assets/{search-CzwrTNr6.js → search--TvX-fig.js} +1 -1
  65. package/dashboard/dist/assets/{skills-api-Dvs1Ob20.js → skills-api-C5C8a1at.js} +1 -1
  66. package/dashboard/dist/assets/{store-6KvYORca.js → store-CRA8gPnP.js} +1 -1
  67. package/dashboard/dist/assets/{terminal-CLniV5ss.js → terminal-BXZyl2ek.js} +1 -1
  68. package/dashboard/dist/assets/{trash-2-BGDRYHet.js → trash-2-IFdBQ65v.js} +1 -1
  69. package/dashboard/dist/assets/{use-community-stats-C4mjay7y.js → use-community-stats-LyJ5aMpX.js} +1 -1
  70. package/dashboard/dist/assets/{use-limits-display-prefs-Bk5YiEg6.js → use-limits-display-prefs-CKvFS1Nn.js} +1 -1
  71. package/dashboard/dist/assets/{use-native-settings-DK4S67CD.js → use-native-settings-Cg-4j_0z.js} +1 -1
  72. package/dashboard/dist/assets/{use-ordered-list-BP6sp53s.js → use-ordered-list-2pX5PXZO.js} +1 -1
  73. package/dashboard/dist/assets/use-reduced-motion-BCR01kRg.js +1 -0
  74. package/dashboard/dist/assets/{use-session-efficiency-pref-CKeItZTw.js → use-session-efficiency-pref-DilJOODA.js} +1 -1
  75. package/dashboard/dist/assets/{use-transform-D2S8SkIw.js → use-transform-DiyYTLpI.js} +1 -1
  76. package/dashboard/dist/assets/{use-usage-limits-reVqkBqL.js → use-usage-limits-DkUmp_nW.js} +1 -1
  77. package/dashboard/dist/assets/{useCurrency-D8pQov2B.js → useCurrency-lm7KjGoD.js} +1 -1
  78. package/dashboard/dist/assets/{useScrollLock-BUlSrwAv.js → useScrollLock-DVSyJ-GE.js} +1 -1
  79. package/dashboard/dist/assets/{useTokenFormat-C4ieP4PU.js → useTokenFormat-Li8BANKO.js} +1 -1
  80. package/dashboard/dist/assets/{zap-FZb_tZVI.js → zap-CKZt1_i5.js} +1 -1
  81. package/dashboard/dist/index.html +1 -1
  82. package/dashboard/dist/ip-check.html +1 -1
  83. package/dashboard/dist/leaderboard.html +1 -1
  84. package/dashboard/dist/share.html +1 -1
  85. package/package.json +1 -1
  86. package/src/lib/pricing/curated-overrides.json +536 -101
  87. package/src/lib/pricing/seed-snapshot.json +1 -1
  88. package/src/lib/rollout.js +373 -72
  89. package/dashboard/dist/assets/DialogDescription-BwPmb3Uv.js +0 -1
  90. package/dashboard/dist/assets/use-reduced-motion-B_GOyvBE.js +0 -1
@@ -5,111 +5,546 @@
5
5
  "deepseek_v4_pro_pricing_note": "The 2026-05-31 promotional-expiry warning is obsolete: api-docs.deepseek.com/quick_start/pricing verified on 2026-06-10 still lists input $0.435/M, output $0.87/M, cache_read $0.003625/M as the official v4-pro rates (the discount became the list price). Do NOT 'restore' the old 4x rates; re-verify against the official page before any change."
6
6
  },
7
7
  "exact": {
8
- "claude-fable-5": { "input": 10, "output": 50, "cache_read": 1, "cache_write": 12.5, "note": "Fable 5 — Anthropic's top tier above Opus. Not yet in LiteLLM. $10/$50 per MTok; cache_read 0.1x, cache_write 1.25x. Remove once LiteLLM carries it." },
9
- "claude-opus-5": { "input": 5, "output": 25, "cache_read": 0.5, "cache_write": 6.25, "note": "Opus 5 (public 2026-07). Standard Opus tier, unchanged from 4.6/4.7/4.8: $5/$25 per MTok, cache_read 0.1x, cache_write 1.25x. Not yet in LiteLLM. Fast mode is billed at $10/$50 — see claude-opus-5-fast. Remove once LiteLLM carries it." },
10
- "claude-opus-5-fast":{ "input": 10, "output": 50, "cache_read": 1, "cache_write": 12.5, "note": "Opus 5 fast mode (speed=\"fast\", beta fast-mode-2026-02-01) is priced at $10/$50 per MTok — 2x standard Opus. Kept separate so a -fast model id is not silently priced at the standard tier. See claude-opus-5." },
11
- "claude-opus-4-8": { "input": 5, "output": 25, "cache_read": 0.5, "cache_write": 6.25, "note": "Opus 4.8 not yet in LiteLLM. Pin to the standard Opus tier (same as 4.6/4.7); remove once LiteLLM carries it." },
12
- "claude-sonnet-5": { "input": 3, "output": 15, "cache_read": 0.3, "cache_write": 3.75, "note": "Sonnet 5. LiteLLM added this model on ~2026-06-30 with matching rates, but pin explicitly so pricing is correct immediately rather than depending on the 24h local pricing-cache refresh cycle." },
13
- "claude-haiku-4-5": { "input": 1, "output": 5, "cache_read": 0.1, "cache_write": 1.25, "note": "Canonical undated alias emitted by Cursor. Pin to the same rates as claude-haiku-4-5-20251001 so local cost does not depend on which dated LiteLLM aliases are present in the current cache." },
14
- "gpt-5.6-sol": { "input": 5, "output": 30, "cache_read": 0.5, "cache_write": 6.25, "note": "GPT-5.6 family (public 2026-07-09). Flagship tier. Not yet in LiteLLM. developers.openai.com/api/docs/pricing: $5/$30 per MTok, cached input 0.1x, cache write 1.25x. Codex emits gpt-5.6-sol (+ reasoning-effort variants like gpt-5.6-solhigh, caught by fuzzy). Remove once LiteLLM carries it." },
15
- "gpt-5.6-terra": { "input": 2.5, "output": 15, "cache_read": 0.25, "cache_write": 3.125, "note": "GPT-5.6 balanced default tier. See gpt-5.6-sol." },
16
- "gpt-5.6-luna": { "input": 1, "output": 6, "cache_read": 0.1, "cache_write": 1.25, "note": "GPT-5.6 lightweight/cost-efficient tier. See gpt-5.6-sol." },
17
- "kiro-agent": { "input": 3, "output": 15, "cache_read": 0.3, "cache_write": 3.75 },
18
- "kiro-cli-agent": { "input": 3, "output": 15, "cache_read": 0.3, "cache_write": 3.75 },
19
- "hy3-preview-agent":{ "input": 0.167, "output": 0.556, "cache_read": 0.056, "cache_write": 0.167, "note": "Tencent Hunyuan Hy3 preview (CodeBuddy/WorkBuddy backend). Official TokenHub rate: 1.2 / 0.4 (cache hit) / 4.0 RMB per MTok in/read/out, converted at ~7.2 RMB/USD. DeepSeek-style cache: no write surcharge, so cache_write = input." },
20
- "hy3-preview": { "input": 0.167, "output": 0.556, "cache_read": 0.056, "cache_write": 0.167, "note": "See hy3-preview-agent." },
21
- "composer-1": { "input": 1.25, "output": 10, "cache_read": 0.125 },
22
- "composer-1.5": { "input": 3.5, "output": 17.5, "cache_read": 0.35 },
23
- "composer-2": { "input": 0.5, "output": 2.5, "cache_read": 0.2 },
24
- "composer-2-fast": { "input": 1.5, "output": 7.5, "cache_read": 0.15 },
25
- "MiniMax-M2.7": { "input": 0.3, "output": 1.2, "cache_read": 0.06, "cache_write": 0.375 },
26
- "MiniMax-M2.7-highspeed":{ "input": 0.6, "output": 2.4, "cache_read": 0.06, "cache_write": 0.375 },
27
- "minimax-m3": { "input": 0.3, "output": 1.2, "cache_read": 0.06, "cache_write": 0 },
28
- "deepseek-v4-flash":{ "input": 0.14, "output": 0.28, "cache_read": 0.0028, "cache_write": 0.14 },
29
- "deepseek-v4-pro": { "input": 0.435,"output": 0.87, "cache_read": 0.003625, "cache_write": 0.435 },
30
- "deepseek-chat": { "input": 0.14, "output": 0.28, "cache_read": 0.0028, "cache_write": 0.14 },
31
- "grok-build": { "input": 1.25, "output": 2.50, "cache_read": 0.20, "note": "Grok Build TUI estimate. Local telemetry currently exposes totalTokens without a stable prompt/output/cache split, so TokenTracker estimates input/output split until Grok exposes per-call usage details." },
32
- "cursor-grok-4.5": { "input": 2.00, "output": 6.00, "cache_read": 0.50, "cache_write": 0, "note": "Cursor Grok 4.5 base rate. Cursor Models & Pricing: $2/M input, $0.50/M cached input, $6/M output; no cache-write rate is published." },
33
- "cursor-grok-4.5-fast": { "input": 4.00, "output": 18.00, "cache_read": 1.00, "cache_write": 0, "note": "Cursor Grok 4.5 Fast rate. Cursor publishes $4/M input and $18/M output; the Cursor usage breakdown reports $1/M cached input. Kept separate from the cheaper base SKU." },
34
- "grok-4-0709": { "input": 3.00, "output": 15.00, "cache_read": 0.75 },
35
- "grok-4": { "input": 3.00, "output": 15.00, "cache_read": 0.75 },
36
- "grok-4-latest": { "input": 3.00, "output": 15.00, "cache_read": 0.75 },
37
- "grok-4-fast": { "input": 0.20, "output": 0.50, "cache_read": 0.05 },
38
- "grok-4-fast-reasoning": { "input": 0.20, "output": 0.50, "cache_read": 0.05 },
39
- "grok-4-fast-non-reasoning": { "input": 0.20, "output": 0.50, "cache_read": 0.05 },
40
- "grok-4-1-fast-non-reasoning": { "input": 0.20, "output": 0.50, "cache_read": 0.05 },
41
- "deepseek-reasoner":{ "input": 0.14, "output": 0.28, "cache_read": 0.0028, "cache_write": 0.14 },
42
- "kimi-for-coding": { "input": 0.6, "output": 2, "cache_read": 0.15 },
43
- "kimi-k2.5": { "input": 0.6, "output": 2, "cache_read": 0.15 },
44
- "kimi-k2.5-free": { "input": 0, "output": 0, "cache_read": 0 },
45
- "kimi-k2.6": { "input": 0.95, "output": 4, "cache_read": 0.16 },
46
- "kimi-k2.7-code": { "input": 0.95, "output": 4, "cache_read": 0.19 },
47
- "kimi-k3": { "input": 3, "output": 15, "cache_read": 0.3, "note": "Kimi K3 (released 2026-07-16). Reported API rates: $3/M input, $15/M output, $0.30/M cached input (0.1x). Not yet in LiteLLM; remove once it carries k3. Kimi Code records the alias as bare \"k3\" (kimi-code/k3), hence the separate exact key below." },
48
- "k3": { "input": 3, "output": 15, "cache_read": 0.3, "note": "Bare alias emitted by Kimi Code (modelAlias \"kimi-code/k3\" -> \"k3\"). Same rates as kimi-k3." },
49
- "glm-5.2": { "input": 1.4, "output": 4.4, "cache_read": 0.26 },
50
- "glm-5.1": { "input": 1.4, "output": 4.4, "cache_read": 0.26 },
51
- "glm-5": { "input": 1.0, "output": 3.2, "cache_read": 0.2 },
52
- "glm-5-turbo": { "input": 1.2, "output": 4.0, "cache_read": 0.24 },
53
- "glm-4.7": { "input": 0.6, "output": 2.2, "cache_read": 0.11 },
54
- "glm-4.7-flashx": { "input": 0.07, "output": 0.4, "cache_read": 0.01 },
55
- "glm-4.7-flash": { "input": 0, "output": 0, "cache_read": 0 },
56
- "glm-4.6": { "input": 0.6, "output": 2.2, "cache_read": 0.11 },
57
- "glm-4.5": { "input": 0.6, "output": 2.2, "cache_read": 0.11 },
58
- "glm-4.5-x": { "input": 2.2, "output": 8.9, "cache_read": 0.45 },
59
- "glm-4.5-air": { "input": 0.2, "output": 1.1, "cache_read": 0.03 },
60
- "glm-4.5-airx": { "input": 1.1, "output": 4.5, "cache_read": 0.22 },
61
- "glm-4.5-flash": { "input": 0, "output": 0, "cache_read": 0 },
62
- "glm-4.7-free": { "input": 0, "output": 0, "cache_read": 0 },
63
- "nemotron-3-super-free":{ "input": 0, "output": 0, "cache_read": 0 },
64
- "mimo-v2-pro-free": { "input": 0, "output": 0, "cache_read": 0 },
65
- "minimax-m2.1-free":{ "input": 0, "output": 0, "cache_read": 0 },
66
- "MiniMax-M2.1": { "input": 0.5, "output": 3, "cache_read": 0.05 },
67
- "antigravity-gpt-oss-120b": { "input": 2.5, "output": 10, "cache_read": 0, "note": "Antigravity bridge alias. Approximate with GPT-4o-class pricing until Google exposes route-specific billing metadata." },
68
- "sakana/fugu-ultra": { "input": 5, "output": 30, "cache_read": 0.5, "cache_write": 5, "note": "Sakana Fugu Ultra — multi-agent orchestration via an OpenAI-compatible API (sakana.ai PAYG + OpenRouter), used through Codex/Cursor/Cline/ZCode etc. OpenRouter rates: $5/$30 per MTok in/out, cache_read $0.5/M; no published cache-write surcharge, so cache_write = input. Subscriptions ($20/$100/$200) bill against this token-equivalent rate." },
69
- "longcat-2.0": { "input": 0.278, "output": 1.111, "cache_read": 0.00556, "cache_write": 0.278, "note": "Meituan LongCat-2.0, seen via ZCode custom-provider routing (issue #276). Official longcat.chat/platform launch-promo rate: RMB 2 / 0.04 (cache hit) / 8 per MTok in/read/out, converted at ~7.2 RMB/USD. Standard list price is RMB 5/0.10/20 (2.5x) once the promo ends — re-verify against longcat.chat/platform/docs/zh/Pricing/LongCat-2.0.html before assuming either is permanent. DeepSeek-style cache: no write surcharge, so cache_write = input." },
70
- "step-3.7-flash": { "input": 0.2, "output": 1.15, "cache_read": 0.04, "cache_write": 0.2, "note": "StepFun Step 3.7 Flash (issue #283). Official platform.stepfun.ai pricing: $0.20 / $0.04 (cache hit) / $1.15 per MTok in/read/out. No published cache-write surcharge, so cache_write = input." },
71
- "step-3.5-flash": { "input": 0.1, "output": 0.3, "cache_read": 0.02, "cache_write": 0.1, "note": "StepFun Step 3.5 Flash incl. dated snapshots like step-3.5-flash-2603 (issue #283). Official platform.stepfun.ai pricing: $0.10 / $0.02 (cache hit) / $0.30 per MTok in/read/out. No published cache-write surcharge, so cache_write = input." }
8
+ "claude-fable-5": {
9
+ "input": 10,
10
+ "output": 50,
11
+ "cache_read": 1,
12
+ "cache_write": 12.5,
13
+ "note": "Fable 5 \u2014 Anthropic's top tier above Opus. Not yet in LiteLLM. $10/$50 per MTok; cache_read 0.1x, cache_write 1.25x. Remove once LiteLLM carries it."
14
+ },
15
+ "claude-opus-5": {
16
+ "input": 5,
17
+ "output": 25,
18
+ "cache_read": 0.5,
19
+ "cache_write": 6.25,
20
+ "note": "Opus 5 (public 2026-07). Standard Opus tier, unchanged from 4.6/4.7/4.8: $5/$25 per MTok, cache_read 0.1x, cache_write 1.25x. Not yet in LiteLLM. Fast mode is billed at $10/$50 \u2014 see claude-opus-5-fast. Remove once LiteLLM carries it."
21
+ },
22
+ "claude-opus-5-fast": {
23
+ "input": 10,
24
+ "output": 50,
25
+ "cache_read": 1,
26
+ "cache_write": 12.5,
27
+ "note": "Opus 5 fast mode (speed=\"fast\", beta fast-mode-2026-02-01) is priced at $10/$50 per MTok \u2014 2x standard Opus. Kept separate so a -fast model id is not silently priced at the standard tier. See claude-opus-5."
28
+ },
29
+ "claude-opus-4-8": {
30
+ "input": 5,
31
+ "output": 25,
32
+ "cache_read": 0.5,
33
+ "cache_write": 6.25,
34
+ "note": "Opus 4.8 not yet in LiteLLM. Pin to the standard Opus tier (same as 4.6/4.7); remove once LiteLLM carries it."
35
+ },
36
+ "claude-sonnet-5": {
37
+ "input": 3,
38
+ "output": 15,
39
+ "cache_read": 0.3,
40
+ "cache_write": 3.75,
41
+ "note": "Sonnet 5. LiteLLM added this model on ~2026-06-30 with matching rates, but pin explicitly so pricing is correct immediately rather than depending on the 24h local pricing-cache refresh cycle."
42
+ },
43
+ "claude-haiku-4-5": {
44
+ "input": 1,
45
+ "output": 5,
46
+ "cache_read": 0.1,
47
+ "cache_write": 1.25,
48
+ "note": "Canonical undated alias emitted by Cursor. Pin to the same rates as claude-haiku-4-5-20251001 so local cost does not depend on which dated LiteLLM aliases are present in the current cache."
49
+ },
50
+ "gpt-5.6-sol": {
51
+ "input": 5,
52
+ "output": 30,
53
+ "cache_read": 0.5,
54
+ "cache_write": 6.25,
55
+ "note": "GPT-5.6 family (public 2026-07-09). Flagship tier. Not yet in LiteLLM. developers.openai.com/api/docs/pricing: $5/$30 per MTok, cached input 0.1x, cache write 1.25x. Codex emits gpt-5.6-sol (+ reasoning-effort variants like gpt-5.6-solhigh, caught by fuzzy). Remove once LiteLLM carries it."
56
+ },
57
+ "gpt-5.6-terra": {
58
+ "input": 2.5,
59
+ "output": 15,
60
+ "cache_read": 0.25,
61
+ "cache_write": 3.125,
62
+ "note": "GPT-5.6 balanced default tier. See gpt-5.6-sol."
63
+ },
64
+ "gpt-5.6-luna": {
65
+ "input": 1,
66
+ "output": 6,
67
+ "cache_read": 0.1,
68
+ "cache_write": 1.25,
69
+ "note": "GPT-5.6 lightweight/cost-efficient tier. See gpt-5.6-sol."
70
+ },
71
+ "kiro-agent": {
72
+ "input": 3,
73
+ "output": 15,
74
+ "cache_read": 0.3,
75
+ "cache_write": 3.75
76
+ },
77
+ "kiro-cli-agent": {
78
+ "input": 3,
79
+ "output": 15,
80
+ "cache_read": 0.3,
81
+ "cache_write": 3.75
82
+ },
83
+ "hy3-preview-agent": {
84
+ "input": 0.167,
85
+ "output": 0.556,
86
+ "cache_read": 0.056,
87
+ "cache_write": 0.167,
88
+ "note": "Tencent Hunyuan Hy3 preview (CodeBuddy/WorkBuddy backend). Official TokenHub rate: 1.2 / 0.4 (cache hit) / 4.0 RMB per MTok in/read/out, converted at ~7.2 RMB/USD. DeepSeek-style cache: no write surcharge, so cache_write = input."
89
+ },
90
+ "hy3-preview": {
91
+ "input": 0.167,
92
+ "output": 0.556,
93
+ "cache_read": 0.056,
94
+ "cache_write": 0.167,
95
+ "note": "See hy3-preview-agent."
96
+ },
97
+ "composer-1": {
98
+ "input": 1.25,
99
+ "output": 10,
100
+ "cache_read": 0.125
101
+ },
102
+ "composer-1.5": {
103
+ "input": 3.5,
104
+ "output": 17.5,
105
+ "cache_read": 0.35
106
+ },
107
+ "composer-2": {
108
+ "input": 0.5,
109
+ "output": 2.5,
110
+ "cache_read": 0.2
111
+ },
112
+ "composer-2-fast": {
113
+ "input": 1.5,
114
+ "output": 7.5,
115
+ "cache_read": 0.15
116
+ },
117
+ "MiniMax-M2.7": {
118
+ "input": 0.3,
119
+ "output": 1.2,
120
+ "cache_read": 0.06,
121
+ "cache_write": 0.375
122
+ },
123
+ "MiniMax-M2.7-highspeed": {
124
+ "input": 0.6,
125
+ "output": 2.4,
126
+ "cache_read": 0.06,
127
+ "cache_write": 0.375
128
+ },
129
+ "minimax-m3": {
130
+ "input": 0.3,
131
+ "output": 1.2,
132
+ "cache_read": 0.06,
133
+ "cache_write": 0
134
+ },
135
+ "deepseek-v4-flash": {
136
+ "input": 0.14,
137
+ "output": 0.28,
138
+ "cache_read": 0.0028,
139
+ "cache_write": 0.14
140
+ },
141
+ "deepseek-v4-pro": {
142
+ "input": 0.435,
143
+ "output": 0.87,
144
+ "cache_read": 0.003625,
145
+ "cache_write": 0.435
146
+ },
147
+ "deepseek-chat": {
148
+ "input": 0.14,
149
+ "output": 0.28,
150
+ "cache_read": 0.0028,
151
+ "cache_write": 0.14
152
+ },
153
+ "grok-build": {
154
+ "input": 1.25,
155
+ "output": 2.5,
156
+ "cache_read": 0.2,
157
+ "note": "Grok Build TUI fallback estimate for sessions without turn_completed.usage. Preferred path is turn_completed usage with per-model modelUsage splits."
158
+ },
159
+ "cursor-grok-4.5": {
160
+ "input": 2.0,
161
+ "output": 6.0,
162
+ "cache_read": 0.5,
163
+ "cache_write": 0,
164
+ "note": "Cursor Grok 4.5 base rate. Cursor Models & Pricing: $2/M input, $0.50/M cached input, $6/M output; no cache-write rate is published."
165
+ },
166
+ "cursor-grok-4.5-fast": {
167
+ "input": 4.0,
168
+ "output": 18.0,
169
+ "cache_read": 1.0,
170
+ "cache_write": 0,
171
+ "note": "Cursor Grok 4.5 Fast rate. Cursor publishes $4/M input and $18/M output; the Cursor usage breakdown reports $1/M cached input. Kept separate from the cheaper base SKU."
172
+ },
173
+ "grok-4-0709": {
174
+ "input": 3.0,
175
+ "output": 15.0,
176
+ "cache_read": 0.75
177
+ },
178
+ "grok-4": {
179
+ "input": 3.0,
180
+ "output": 15.0,
181
+ "cache_read": 0.75
182
+ },
183
+ "grok-4-latest": {
184
+ "input": 3.0,
185
+ "output": 15.0,
186
+ "cache_read": 0.75
187
+ },
188
+ "grok-4-fast": {
189
+ "input": 0.2,
190
+ "output": 0.5,
191
+ "cache_read": 0.05
192
+ },
193
+ "grok-4-fast-reasoning": {
194
+ "input": 0.2,
195
+ "output": 0.5,
196
+ "cache_read": 0.05
197
+ },
198
+ "grok-4-fast-non-reasoning": {
199
+ "input": 0.2,
200
+ "output": 0.5,
201
+ "cache_read": 0.05
202
+ },
203
+ "grok-4-1-fast-non-reasoning": {
204
+ "input": 0.2,
205
+ "output": 0.5,
206
+ "cache_read": 0.05
207
+ },
208
+ "deepseek-reasoner": {
209
+ "input": 0.14,
210
+ "output": 0.28,
211
+ "cache_read": 0.0028,
212
+ "cache_write": 0.14
213
+ },
214
+ "kimi-for-coding": {
215
+ "input": 0.6,
216
+ "output": 2,
217
+ "cache_read": 0.15
218
+ },
219
+ "kimi-k2.5": {
220
+ "input": 0.6,
221
+ "output": 2,
222
+ "cache_read": 0.15
223
+ },
224
+ "kimi-k2.5-free": {
225
+ "input": 0,
226
+ "output": 0,
227
+ "cache_read": 0
228
+ },
229
+ "kimi-k2.6": {
230
+ "input": 0.95,
231
+ "output": 4,
232
+ "cache_read": 0.16
233
+ },
234
+ "kimi-k2.7-code": {
235
+ "input": 0.95,
236
+ "output": 4,
237
+ "cache_read": 0.19
238
+ },
239
+ "kimi-k3": {
240
+ "input": 3,
241
+ "output": 15,
242
+ "cache_read": 0.3,
243
+ "note": "Kimi K3 (released 2026-07-16). Reported API rates: $3/M input, $15/M output, $0.30/M cached input (0.1x). Not yet in LiteLLM; remove once it carries k3. Kimi Code records the alias as bare \"k3\" (kimi-code/k3), hence the separate exact key below."
244
+ },
245
+ "k3": {
246
+ "input": 3,
247
+ "output": 15,
248
+ "cache_read": 0.3,
249
+ "note": "Bare alias emitted by Kimi Code (modelAlias \"kimi-code/k3\" -> \"k3\"). Same rates as kimi-k3."
250
+ },
251
+ "glm-5.2": {
252
+ "input": 1.4,
253
+ "output": 4.4,
254
+ "cache_read": 0.26
255
+ },
256
+ "glm-5.1": {
257
+ "input": 1.4,
258
+ "output": 4.4,
259
+ "cache_read": 0.26
260
+ },
261
+ "glm-5": {
262
+ "input": 1.0,
263
+ "output": 3.2,
264
+ "cache_read": 0.2
265
+ },
266
+ "glm-5-turbo": {
267
+ "input": 1.2,
268
+ "output": 4.0,
269
+ "cache_read": 0.24
270
+ },
271
+ "glm-4.7": {
272
+ "input": 0.6,
273
+ "output": 2.2,
274
+ "cache_read": 0.11
275
+ },
276
+ "glm-4.7-flashx": {
277
+ "input": 0.07,
278
+ "output": 0.4,
279
+ "cache_read": 0.01
280
+ },
281
+ "glm-4.7-flash": {
282
+ "input": 0,
283
+ "output": 0,
284
+ "cache_read": 0
285
+ },
286
+ "glm-4.6": {
287
+ "input": 0.6,
288
+ "output": 2.2,
289
+ "cache_read": 0.11
290
+ },
291
+ "glm-4.5": {
292
+ "input": 0.6,
293
+ "output": 2.2,
294
+ "cache_read": 0.11
295
+ },
296
+ "glm-4.5-x": {
297
+ "input": 2.2,
298
+ "output": 8.9,
299
+ "cache_read": 0.45
300
+ },
301
+ "glm-4.5-air": {
302
+ "input": 0.2,
303
+ "output": 1.1,
304
+ "cache_read": 0.03
305
+ },
306
+ "glm-4.5-airx": {
307
+ "input": 1.1,
308
+ "output": 4.5,
309
+ "cache_read": 0.22
310
+ },
311
+ "glm-4.5-flash": {
312
+ "input": 0,
313
+ "output": 0,
314
+ "cache_read": 0
315
+ },
316
+ "glm-4.7-free": {
317
+ "input": 0,
318
+ "output": 0,
319
+ "cache_read": 0
320
+ },
321
+ "nemotron-3-super-free": {
322
+ "input": 0,
323
+ "output": 0,
324
+ "cache_read": 0
325
+ },
326
+ "mimo-v2-pro-free": {
327
+ "input": 0,
328
+ "output": 0,
329
+ "cache_read": 0
330
+ },
331
+ "minimax-m2.1-free": {
332
+ "input": 0,
333
+ "output": 0,
334
+ "cache_read": 0
335
+ },
336
+ "MiniMax-M2.1": {
337
+ "input": 0.5,
338
+ "output": 3,
339
+ "cache_read": 0.05
340
+ },
341
+ "antigravity-gpt-oss-120b": {
342
+ "input": 2.5,
343
+ "output": 10,
344
+ "cache_read": 0,
345
+ "note": "Antigravity bridge alias. Approximate with GPT-4o-class pricing until Google exposes route-specific billing metadata."
346
+ },
347
+ "sakana/fugu-ultra": {
348
+ "input": 5,
349
+ "output": 30,
350
+ "cache_read": 0.5,
351
+ "cache_write": 5,
352
+ "note": "Sakana Fugu Ultra \u2014 multi-agent orchestration via an OpenAI-compatible API (sakana.ai PAYG + OpenRouter), used through Codex/Cursor/Cline/ZCode etc. OpenRouter rates: $5/$30 per MTok in/out, cache_read $0.5/M; no published cache-write surcharge, so cache_write = input. Subscriptions ($20/$100/$200) bill against this token-equivalent rate."
353
+ },
354
+ "longcat-2.0": {
355
+ "input": 0.278,
356
+ "output": 1.111,
357
+ "cache_read": 0.00556,
358
+ "cache_write": 0.278,
359
+ "note": "Meituan LongCat-2.0, seen via ZCode custom-provider routing (issue #276). Official longcat.chat/platform launch-promo rate: RMB 2 / 0.04 (cache hit) / 8 per MTok in/read/out, converted at ~7.2 RMB/USD. Standard list price is RMB 5/0.10/20 (2.5x) once the promo ends \u2014 re-verify against longcat.chat/platform/docs/zh/Pricing/LongCat-2.0.html before assuming either is permanent. DeepSeek-style cache: no write surcharge, so cache_write = input."
360
+ },
361
+ "step-3.7-flash": {
362
+ "input": 0.2,
363
+ "output": 1.15,
364
+ "cache_read": 0.04,
365
+ "cache_write": 0.2,
366
+ "note": "StepFun Step 3.7 Flash (issue #283). Official platform.stepfun.ai pricing: $0.20 / $0.04 (cache hit) / $1.15 per MTok in/read/out. No published cache-write surcharge, so cache_write = input."
367
+ },
368
+ "step-3.5-flash": {
369
+ "input": 0.1,
370
+ "output": 0.3,
371
+ "cache_read": 0.02,
372
+ "cache_write": 0.1,
373
+ "note": "StepFun Step 3.5 Flash incl. dated snapshots like step-3.5-flash-2603 (issue #283). Official platform.stepfun.ai pricing: $0.10 / $0.02 (cache hit) / $0.30 per MTok in/read/out. No published cache-write surcharge, so cache_write = input."
374
+ },
375
+ "grok-4.5-build": {
376
+ "input": 2.0,
377
+ "output": 6.0,
378
+ "cache_read": 0.5,
379
+ "cache_write": 0,
380
+ "note": "Grok Build paid SKU (grok-4.5-build from turn_completed.modelUsage). Rates aligned with xAI/Cursor Grok 4.5 until a dedicated Build price list is published."
381
+ },
382
+ "grok-build-free": {
383
+ "input": 0,
384
+ "output": 0,
385
+ "cache_read": 0,
386
+ "cache_write": 0,
387
+ "note": "Grok Build free tier (canonical model id). $0 marginal cost. Prefer this over grok-4.5-build-free so pricing lookup cannot fuzzy-match paid grok-4.5."
388
+ },
389
+ "grok-4.5-build-free": {
390
+ "input": 0,
391
+ "output": 0,
392
+ "cache_read": 0,
393
+ "cache_write": 0,
394
+ "note": "Alias for free Build SKU as labeled in Grok updates.jsonl modelUsage. Canonical storage id is grok-build-free."
395
+ }
72
396
  },
73
397
  "alias": {
74
398
  "auto": "composer-1"
75
399
  },
76
400
  "fuzzy": [
77
- { "match": "gpt-5.6-sol", "ref": "gpt-5.6-sol" },
78
- { "match": "gpt-5.6-terra", "ref": "gpt-5.6-terra" },
79
- { "match": "gpt-5.6-luna", "ref": "gpt-5.6-luna" },
80
- { "match": "gpt-5.6", "ref": "gpt-5.6-terra" },
81
- { "match": "kiro", "ref": "kiro-cli-agent" },
82
- { "match": "hy3", "ref": "hy3-preview-agent" },
83
- { "match": "composer", "ref": "composer-1" },
84
- { "match": "claude-opus-5-fast", "ref": "claude-opus-5-fast" },
85
- { "match": "claude-opus-5", "ref": "claude-opus-5" },
86
- { "match": "claude-opus-4-8", "ref": "claude-opus-4-8" },
87
- { "match": "minimax-m3", "ref": "minimax-m3" },
88
- { "match": "minimax-m2.7-highspeed", "ref": "MiniMax-M2.7-highspeed" },
89
- { "match": "minimax-m2.7", "ref": "MiniMax-M2.7" },
90
- { "match": "deepseek-v4-flash", "ref": "deepseek-v4-flash" },
91
- { "match": "deepseek-v4-pro", "ref": "deepseek-v4-pro" },
92
- { "match": "kimi-k2.7-code", "ref": "kimi-k2.7-code" },
93
- { "match": "kimi-k3", "ref": "kimi-k3" },
94
- { "match": "kimi-k2.6", "ref": "kimi-k2.6" },
95
- { "match": "kimi", "ref": "kimi-k2.5" },
96
- { "match": "glm-4.5-airx", "ref": "glm-4.5-airx" },
97
- { "match": "glm-4.5-air", "ref": "glm-4.5-air" },
98
- { "match": "glm-4.5-x", "ref": "glm-4.5-x" },
99
- { "match": "glm-4.5-flash", "ref": "glm-4.5-flash" },
100
- { "match": "glm-4.5", "ref": "glm-4.5" },
101
- { "match": "glm-4.7-flashx", "ref": "glm-4.7-flashx" },
102
- { "match": "glm-4.7-flash", "ref": "glm-4.7-flash" },
103
- { "match": "glm-4.7", "ref": "glm-4.7" },
104
- { "match": "glm-4.6", "ref": "glm-4.6" },
105
- { "match": "glm-5-turbo", "ref": "glm-5-turbo" },
106
- { "match": "glm-5.2", "ref": "glm-5.2" },
107
- { "match": "glm-5.1", "ref": "glm-5.1" },
108
- { "match": "glm-5", "ref": "glm-5" },
109
- { "match": "fugu", "ref": "sakana/fugu-ultra" },
110
- { "match": "longcat", "ref": "longcat-2.0" },
111
- { "match": "step-3.7-flash", "ref": "step-3.7-flash" },
112
- { "match": "step-3.5-flash", "ref": "step-3.5-flash" },
113
- { "match": "stepfun", "ref": "step-3.7-flash" }
401
+ {
402
+ "match": "gpt-5.6-sol",
403
+ "ref": "gpt-5.6-sol"
404
+ },
405
+ {
406
+ "match": "gpt-5.6-terra",
407
+ "ref": "gpt-5.6-terra"
408
+ },
409
+ {
410
+ "match": "gpt-5.6-luna",
411
+ "ref": "gpt-5.6-luna"
412
+ },
413
+ {
414
+ "match": "gpt-5.6",
415
+ "ref": "gpt-5.6-terra"
416
+ },
417
+ {
418
+ "match": "kiro",
419
+ "ref": "kiro-cli-agent"
420
+ },
421
+ {
422
+ "match": "hy3",
423
+ "ref": "hy3-preview-agent"
424
+ },
425
+ {
426
+ "match": "composer",
427
+ "ref": "composer-1"
428
+ },
429
+ {
430
+ "match": "claude-opus-5-fast",
431
+ "ref": "claude-opus-5-fast"
432
+ },
433
+ {
434
+ "match": "claude-opus-5",
435
+ "ref": "claude-opus-5"
436
+ },
437
+ {
438
+ "match": "claude-opus-4-8",
439
+ "ref": "claude-opus-4-8"
440
+ },
441
+ {
442
+ "match": "minimax-m3",
443
+ "ref": "minimax-m3"
444
+ },
445
+ {
446
+ "match": "minimax-m2.7-highspeed",
447
+ "ref": "MiniMax-M2.7-highspeed"
448
+ },
449
+ {
450
+ "match": "minimax-m2.7",
451
+ "ref": "MiniMax-M2.7"
452
+ },
453
+ {
454
+ "match": "deepseek-v4-flash",
455
+ "ref": "deepseek-v4-flash"
456
+ },
457
+ {
458
+ "match": "deepseek-v4-pro",
459
+ "ref": "deepseek-v4-pro"
460
+ },
461
+ {
462
+ "match": "kimi-k2.7-code",
463
+ "ref": "kimi-k2.7-code"
464
+ },
465
+ {
466
+ "match": "kimi-k3",
467
+ "ref": "kimi-k3"
468
+ },
469
+ {
470
+ "match": "kimi-k2.6",
471
+ "ref": "kimi-k2.6"
472
+ },
473
+ {
474
+ "match": "kimi",
475
+ "ref": "kimi-k2.5"
476
+ },
477
+ {
478
+ "match": "glm-4.5-airx",
479
+ "ref": "glm-4.5-airx"
480
+ },
481
+ {
482
+ "match": "glm-4.5-air",
483
+ "ref": "glm-4.5-air"
484
+ },
485
+ {
486
+ "match": "glm-4.5-x",
487
+ "ref": "glm-4.5-x"
488
+ },
489
+ {
490
+ "match": "glm-4.5-flash",
491
+ "ref": "glm-4.5-flash"
492
+ },
493
+ {
494
+ "match": "glm-4.5",
495
+ "ref": "glm-4.5"
496
+ },
497
+ {
498
+ "match": "glm-4.7-flashx",
499
+ "ref": "glm-4.7-flashx"
500
+ },
501
+ {
502
+ "match": "glm-4.7-flash",
503
+ "ref": "glm-4.7-flash"
504
+ },
505
+ {
506
+ "match": "glm-4.7",
507
+ "ref": "glm-4.7"
508
+ },
509
+ {
510
+ "match": "glm-4.6",
511
+ "ref": "glm-4.6"
512
+ },
513
+ {
514
+ "match": "glm-5-turbo",
515
+ "ref": "glm-5-turbo"
516
+ },
517
+ {
518
+ "match": "glm-5.2",
519
+ "ref": "glm-5.2"
520
+ },
521
+ {
522
+ "match": "glm-5.1",
523
+ "ref": "glm-5.1"
524
+ },
525
+ {
526
+ "match": "glm-5",
527
+ "ref": "glm-5"
528
+ },
529
+ {
530
+ "match": "fugu",
531
+ "ref": "sakana/fugu-ultra"
532
+ },
533
+ {
534
+ "match": "longcat",
535
+ "ref": "longcat-2.0"
536
+ },
537
+ {
538
+ "match": "step-3.7-flash",
539
+ "ref": "step-3.7-flash"
540
+ },
541
+ {
542
+ "match": "step-3.5-flash",
543
+ "ref": "step-3.5-flash"
544
+ },
545
+ {
546
+ "match": "stepfun",
547
+ "ref": "step-3.7-flash"
548
+ }
114
549
  ]
115
550
  }