@molecule/api-resource-ai-models 1.4.0 → 1.5.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/README.md CHANGED
@@ -3,7 +3,7 @@ AUTO-GENERATED — DO NOT EDIT THIS FILE.
3
3
  Generated by `mlcl sync-docs` from the package's src/index.ts JSDoc + mlcl/registry.json.
4
4
  Edits here are overwritten on the next commit (molecule's pre-commit hook regenerates).
5
5
  To change this document, edit the module-level JSDoc in src/index.ts.
6
- Generated: 2026-09-01T19:40:01.931Z
6
+ Generated: 2026-09-02T19:41:05.247Z
7
7
  -->
8
8
 
9
9
  # @molecule/api-resource-ai-models
@@ -846,6 +846,15 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
846
846
  at the post-promo list rate ($1.50/$7.50, cache read $0.15) under the same
847
847
  never-under-charge policy, each with a KNOWN_DIVERGENCES entry expiring
848
848
  2026-12-31. gemini-3.5-flash carries no promo: still $1.50/$9.00, $0.15)
849
+ (re-verified 2026-09-02: gemini-3.8-flash "New Stable" — supersedes
850
+ 3.7-flash, which the models page now calls "previous-generation". THIRD id
851
+ on the same flash-tier promo row, verbatim "$0.75 through December 31,
852
+ 2026. $1.50 starting January 1, 2027." / "$3.75 … $7.50" / "$0.075 …
853
+ $0.15", so it too is cataloged at list with a KNOWN_DIVERGENCES entry.
854
+ Specs from /docs/models/gemini-3.8-flash: 1M ctx / 65,536 out, thinking
855
+ low|medium|high, vision, tools, caching, search grounding, code execution,
856
+ url context — identical surface to 3.7-flash. Knowledge cutoff is published
857
+ by neither Google nor models.dev for this id)
849
858
  - xAI: https://docs.x.ai/developers/models + /developers/grok-4-5
850
859
  (grok-4.5 flagship 2026-07-08: $2/$6, 500K ctx, ≥200K prompts bill 2× —
851
860
  not modeled; reasoning_effort low|medium|high default high, image input;
package/dist/models.d.ts CHANGED
@@ -114,6 +114,15 @@ import type { ModelDefinition } from './types.js';
114
114
  * at the post-promo list rate ($1.50/$7.50, cache read $0.15) under the same
115
115
  * never-under-charge policy, each with a KNOWN_DIVERGENCES entry expiring
116
116
  * 2026-12-31. gemini-3.5-flash carries no promo: still $1.50/$9.00, $0.15)
117
+ * (re-verified 2026-09-02: gemini-3.8-flash "New Stable" — supersedes
118
+ * 3.7-flash, which the models page now calls "previous-generation". THIRD id
119
+ * on the same flash-tier promo row, verbatim "$0.75 through December 31,
120
+ * 2026. $1.50 starting January 1, 2027." / "$3.75 … $7.50" / "$0.075 …
121
+ * $0.15", so it too is cataloged at list with a KNOWN_DIVERGENCES entry.
122
+ * Specs from /docs/models/gemini-3.8-flash: 1M ctx / 65,536 out, thinking
123
+ * low|medium|high, vision, tools, caching, search grounding, code execution,
124
+ * url context — identical surface to 3.7-flash. Knowledge cutoff is published
125
+ * by neither Google nor models.dev for this id)
117
126
  * - xAI: https://docs.x.ai/developers/models + /developers/grok-4-5
118
127
  * (grok-4.5 flagship 2026-07-08: $2/$6, 500K ctx, ≥200K prompts bill 2× —
119
128
  * not modeled; reasoning_effort low|medium|high default high, image input;
@@ -1 +1 @@
1
- {"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAQjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAwLG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EAupDnC,CAAA"}
1
+ {"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAQjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAiMG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EAusDnC,CAAA"}
package/dist/models.js CHANGED
@@ -118,6 +118,15 @@ const WEEKDAYS_UTC = [1, 2, 3, 4, 5];
118
118
  * at the post-promo list rate ($1.50/$7.50, cache read $0.15) under the same
119
119
  * never-under-charge policy, each with a KNOWN_DIVERGENCES entry expiring
120
120
  * 2026-12-31. gemini-3.5-flash carries no promo: still $1.50/$9.00, $0.15)
121
+ * (re-verified 2026-09-02: gemini-3.8-flash "New Stable" — supersedes
122
+ * 3.7-flash, which the models page now calls "previous-generation". THIRD id
123
+ * on the same flash-tier promo row, verbatim "$0.75 through December 31,
124
+ * 2026. $1.50 starting January 1, 2027." / "$3.75 … $7.50" / "$0.075 …
125
+ * $0.15", so it too is cataloged at list with a KNOWN_DIVERGENCES entry.
126
+ * Specs from /docs/models/gemini-3.8-flash: 1M ctx / 65,536 out, thinking
127
+ * low|medium|high, vision, tools, caching, search grounding, code execution,
128
+ * url context — identical surface to 3.7-flash. Knowledge cutoff is published
129
+ * by neither Google nor models.dev for this id)
121
130
  * - xAI: https://docs.x.ai/developers/models + /developers/grok-4-5
122
131
  * (grok-4.5 flagship 2026-07-08: $2/$6, 500K ctx, ≥200K prompts bill 2× —
123
132
  * not modeled; reasoning_effort low|medium|high default high, image input;
@@ -776,11 +785,48 @@ export const MODELS = [
776
785
  // replace outright: the google bond has never been implemented/wired, so no
777
786
  // historical usage can reference the old ids.
778
787
  // ---------------------------------------------------------------------------
788
+ {
789
+ id: 'gemini-3.8-flash',
790
+ provider: 'google',
791
+ label: 'Gemini 3.8 Flash',
792
+ description: 'Google agentic flagship — long-horizon engineering & autonomous agents',
793
+ // Verified against /docs/models/gemini-3.8-flash (2026-09-02) — same limits
794
+ // and capability surface as 3.7-flash.
795
+ contextWindow: 1_048_576,
796
+ maxOutputTokens: 65_536,
797
+ supportsThinking: true,
798
+ thinkingBudgetTokens: 10_000,
799
+ thinkingConfigurable: true,
800
+ // thinking_level low|medium|high — minimal NOT supported on this model.
801
+ supportedEffortLevels: ['low', 'medium', 'high'],
802
+ defaultEffortLevel: 'medium',
803
+ supportsVision: true,
804
+ supportsPromptCaching: true,
805
+ supportsTools: true,
806
+ webSearchToolType: 'google_search',
807
+ codeExecutionToolType: 'code_execution',
808
+ webFetchToolType: 'url_context',
809
+ // "New Stable" 2026-09-02. LIST price $1.50/$7.50 — same card as 3.7- and
810
+ // 3.6-flash, and on the same flash-tier launch promo ($0.75/$3.75, cache
811
+ // read $0.075) through 2026-12-31; billed here at standard list so metering
812
+ // never under-charges (identical policy and expiry to the two entries
813
+ // below — see the matching KNOWN_DIVERGENCES entry in
814
+ // scripts/check-model-freshness.mjs).
815
+ inputPricePerMTok: 1.5,
816
+ outputPricePerMTok: 7.5,
817
+ // Gemini context cache: read $0.15/M (0.1× input), no write premium
818
+ // (storage billed separately per hour — not modeled).
819
+ cacheReadPricePerMTok: 0.15,
820
+ cacheWritePricePerMTok: 1.5,
821
+ // Published by neither Google's model page nor models.dev for this id —
822
+ // carried from 3.7-flash (same generation) as a best-effort estimate.
823
+ knowledgeCutoff: '2026-03-01',
824
+ },
779
825
  {
780
826
  id: 'gemini-3.7-flash',
781
827
  provider: 'google',
782
828
  label: 'Gemini 3.7 Flash',
783
- description: 'Google agentic flagship — complex coding & multi-step execution',
829
+ description: 'Previous Google agentic flagship — complex coding & multi-step execution',
784
830
  // Verified against /docs/models/gemini-3.7-flash (2026-08-13).
785
831
  contextWindow: 1_048_576,
786
832
  maxOutputTokens: 65_536,
@@ -810,6 +856,11 @@ export const MODELS = [
810
856
  cacheWritePricePerMTok: 1.5,
811
857
  // Not on Google's docs — models.dev reports 2026-03 (lead, not authority).
812
858
  knowledgeCutoff: '2026-03-01',
859
+ // Superseded by gemini-3.8-flash (2026-09-02) — same flash tier, same list
860
+ // price; Google's own models page now calls 3.7 "previous-generation".
861
+ // Still served upstream, so it stays priceable.
862
+ deprecatedAt: '2026-09-02',
863
+ supersededBy: 'gemini-3.8-flash',
813
864
  },
814
865
  {
815
866
  id: 'gemini-3.6-flash',
@@ -850,9 +901,11 @@ export const MODELS = [
850
901
  knowledgeCutoff: '2026-01-01',
851
902
  // Superseded by gemini-3.7-flash (2026-08-13) — same flash tier, same list
852
903
  // price; Google's own models page now calls 3.6 "previous-generation".
853
- // Still served upstream, so it stays priceable.
904
+ // Still served upstream, so it stays priceable. Re-pointed to
905
+ // gemini-3.8-flash on 2026-09-02 when 3.7 was itself deprecated —
906
+ // supersededBy must target a SELECTABLE model, never a chain.
854
907
  deprecatedAt: '2026-08-13',
855
- supersededBy: 'gemini-3.7-flash',
908
+ supersededBy: 'gemini-3.8-flash',
856
909
  },
857
910
  {
858
911
  id: 'gemini-3.5-flash',
@@ -883,11 +936,11 @@ export const MODELS = [
883
936
  cacheWritePricePerMTok: 1.5,
884
937
  knowledgeCutoff: '2025-01-01',
885
938
  // Superseded within the flash tier (first by 3.6-flash on 2026-07-21, now
886
- // pointed one hop to gemini-3.7-flash — supersededBy must target a
939
+ // pointed one hop to gemini-3.8-flash — supersededBy must target a
887
940
  // SELECTABLE model, never a chain). Still served upstream, so it stays
888
941
  // priceable.
889
942
  deprecatedAt: '2026-07-21',
890
- supersededBy: 'gemini-3.7-flash',
943
+ supersededBy: 'gemini-3.8-flash',
891
944
  },
892
945
  {
893
946
  id: 'gemini-3.1-pro-preview',
@@ -1743,9 +1796,13 @@ export const MODELS = [
1743
1796
  // GLM context cache: read ≈0.19× input, no write premium.
1744
1797
  cacheReadPricePerMTok: 0.26,
1745
1798
  cacheWritePricePerMTok: 1.4,
1746
- // No US re-host exists — DeepInfra serves GLM-5.2 and GLM-5.3-Flash but
1747
- // returns "model not found" for zai-org/GLM-5.3 (checked 2026-08-26), so
1748
- // this is pinned to the native host and bills the list card above.
1799
+ // Pinned to the native host, billing the list card above. DeepInfra has
1800
+ // since STARTED serving zai-org/GLM-5.3 (it 404'd when this was cataloged
1801
+ // on 2026-08-26; live 2026-09-02 at $1.20/$4.00, cache read 0.1× = $0.12,
1802
+ // i.e. below native). Adding 'us' here is not a catalog-only edit — it
1803
+ // needs a matching entry in molecule-dev's region-model-maps.ts and a
1804
+ // regionPricing block, moved together, or dispatch sends the canonical id
1805
+ // and 404s. Left as a human decision.
1749
1806
  regions: ['cn'],
1750
1807
  // Same base weights as glm-5.2, so the same best-effort estimate — Z.ai
1751
1808
  // publishes no cutoff.
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@molecule/api-resource-ai-models",
3
- "version": "1.4.0",
3
+ "version": "1.5.0",
4
4
  "description": "AI model catalog — server-side source of truth plus an authentication-gated discovery endpoint",
5
5
  "type": "module",
6
6
  "main": "dist/index.js",