@molecule/api-resource-ai-models 1.2.0 → 1.2.2

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/README.md CHANGED
@@ -3,7 +3,7 @@ AUTO-GENERATED — DO NOT EDIT THIS FILE.
3
3
  Generated by `mlcl sync-docs` from the package's src/index.ts JSDoc + mlcl/registry.json.
4
4
  Edits here are overwritten on the next commit (molecule's pre-commit hook regenerates).
5
5
  To change this document, edit the module-level JSDoc in src/index.ts.
6
- Generated: 2026-08-13T22:10:11.926Z
6
+ Generated: 2026-08-15T19:37:59.137Z
7
7
  -->
8
8
 
9
9
  # @molecule/api-resource-ai-models
@@ -817,6 +817,16 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
817
817
  125%) — the EN edition of that doc simply lags (zero qwen3.8 mentions),
818
818
  which an earlier pass misread as "excepted/console-only". qwen3.7-max
819
819
  still runs its 50%-off promo — billed here at list, $2.50/$7.50)
820
+ (re-verified 2026-08-15 against api.deepinfra.com/models/: qwen3.8-max's US
821
+ re-host rates are unchanged — $1.65/$4.951, cache read 0.1248× input =
822
+ $0.206. DeepInfra also began serving `Qwen/Qwen3.8-2.4T-A95B` on
823
+ 2026-08-12; by DeepInfra's own description it is the OPEN-WEIGHT variant of
824
+ Qwen3.8 Max (2.4T MoE, 95B active), 262,144 ctx / 131,072 out, $2/$6 with
825
+ cache read 0.1× input. NOT added: it is the same tier as qwen3.8-max, which
826
+ the catalog already offers, and it costs MORE on that very host ($2/$6 vs
827
+ $1.65/$4.951), so nothing would ever select it — and carrying both would
828
+ put two selectable Alibaba flagships in one family. Revisit only if Alibaba
829
+ publishes it as a distinct first-party DashScope model id.)
820
830
  - Zhipu: https://docs.z.ai/guides/overview/pricing (unchanged; glm-5.2 is
821
831
  the newest — "GLM-5.3/5.5" rumors have no released ids as of 2026-07-28)
822
832
 
package/dist/models.d.ts CHANGED
@@ -101,6 +101,16 @@ import type { ModelDefinition } from './types.js';
101
101
  * 125%) — the EN edition of that doc simply lags (zero qwen3.8 mentions),
102
102
  * which an earlier pass misread as "excepted/console-only". qwen3.7-max
103
103
  * still runs its 50%-off promo — billed here at list, $2.50/$7.50)
104
+ * (re-verified 2026-08-15 against api.deepinfra.com/models/: qwen3.8-max's US
105
+ * re-host rates are unchanged — $1.65/$4.951, cache read 0.1248× input =
106
+ * $0.206. DeepInfra also began serving `Qwen/Qwen3.8-2.4T-A95B` on
107
+ * 2026-08-12; by DeepInfra's own description it is the OPEN-WEIGHT variant of
108
+ * Qwen3.8 Max (2.4T MoE, 95B active), 262,144 ctx / 131,072 out, $2/$6 with
109
+ * cache read 0.1× input. NOT added: it is the same tier as qwen3.8-max, which
110
+ * the catalog already offers, and it costs MORE on that very host ($2/$6 vs
111
+ * $1.65/$4.951), so nothing would ever select it — and carrying both would
112
+ * put two selectable Alibaba flagships in one family. Revisit only if Alibaba
113
+ * publishes it as a distinct first-party DashScope model id.)
104
114
  * - Zhipu: https://docs.z.ai/guides/overview/pricing (unchanged; glm-5.2 is
105
115
  * the newest — "GLM-5.3/5.5" rumors have no released ids as of 2026-07-28)
106
116
  *
@@ -1 +1 @@
1
- {"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAEjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAmGG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EA+5CnC,CAAA"}
1
+ {"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAEjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GA6GG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EAk6CnC,CAAA"}
package/dist/models.js CHANGED
@@ -100,6 +100,16 @@
100
100
  * 125%) — the EN edition of that doc simply lags (zero qwen3.8 mentions),
101
101
  * which an earlier pass misread as "excepted/console-only". qwen3.7-max
102
102
  * still runs its 50%-off promo — billed here at list, $2.50/$7.50)
103
+ * (re-verified 2026-08-15 against api.deepinfra.com/models/: qwen3.8-max's US
104
+ * re-host rates are unchanged — $1.65/$4.951, cache read 0.1248× input =
105
+ * $0.206. DeepInfra also began serving `Qwen/Qwen3.8-2.4T-A95B` on
106
+ * 2026-08-12; by DeepInfra's own description it is the OPEN-WEIGHT variant of
107
+ * Qwen3.8 Max (2.4T MoE, 95B active), 262,144 ctx / 131,072 out, $2/$6 with
108
+ * cache read 0.1× input. NOT added: it is the same tier as qwen3.8-max, which
109
+ * the catalog already offers, and it costs MORE on that very host ($2/$6 vs
110
+ * $1.65/$4.951), so nothing would ever select it — and carrying both would
111
+ * put two selectable Alibaba flagships in one family. Revisit only if Alibaba
112
+ * publishes it as a distinct first-party DashScope model id.)
103
113
  * - Zhipu: https://docs.z.ai/guides/overview/pricing (unchanged; glm-5.2 is
104
114
  * the newest — "GLM-5.3/5.5" rumors have no released ids as of 2026-07-28)
105
115
  *
@@ -963,9 +973,12 @@ export const MODELS = [
963
973
  // `FREE_TIER_MODELS[mode] === modelId` before it looks at regions, so
964
974
  // leaving it here would widen nothing — it would just claim a free-tier
965
975
  // relationship that no longer exists.
966
- // US = DeepInfra, verified 2026-08-01 via api.deepinfra.com/models/
976
+ // US = DeepInfra, verified 2026-08-15 via api.deepinfra.com/models/
967
977
  // deepseek-ai/DeepSeek-V4-Pro. No cache-write premium (omitted → region
968
- // input rate).
978
+ // input rate). DeepInfra also serves the dated `-0813` snapshot (the
979
+ // official release superseding the preview weights the bare id carries) at
980
+ // IDENTICAL rates — so bumping the molecule-dev us modelMap to it moves no
981
+ // price here; it is a weights decision, not a pricing one.
969
982
  regionPricing: {
970
983
  us: { inputPricePerMTok: 1.3, outputPricePerMTok: 2.6, cacheReadPricePerMTok: 0.1 },
971
984
  },
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@molecule/api-resource-ai-models",
3
- "version": "1.2.0",
3
+ "version": "1.2.2",
4
4
  "description": "AI model catalog — server-side source of truth plus an authentication-gated discovery endpoint",
5
5
  "type": "module",
6
6
  "main": "dist/index.js",