@molecule/api-resource-ai-models 1.2.0 → 1.2.2
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +11 -1
- package/dist/models.d.ts +10 -0
- package/dist/models.d.ts.map +1 -1
- package/dist/models.js +15 -2
- package/package.json +1 -1
package/README.md
CHANGED
|
@@ -3,7 +3,7 @@ AUTO-GENERATED — DO NOT EDIT THIS FILE.
|
|
|
3
3
|
Generated by `mlcl sync-docs` from the package's src/index.ts JSDoc + mlcl/registry.json.
|
|
4
4
|
Edits here are overwritten on the next commit (molecule's pre-commit hook regenerates).
|
|
5
5
|
To change this document, edit the module-level JSDoc in src/index.ts.
|
|
6
|
-
Generated: 2026-08-
|
|
6
|
+
Generated: 2026-08-15T19:37:59.137Z
|
|
7
7
|
-->
|
|
8
8
|
|
|
9
9
|
# @molecule/api-resource-ai-models
|
|
@@ -817,6 +817,16 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
|
|
|
817
817
|
125%) — the EN edition of that doc simply lags (zero qwen3.8 mentions),
|
|
818
818
|
which an earlier pass misread as "excepted/console-only". qwen3.7-max
|
|
819
819
|
still runs its 50%-off promo — billed here at list, $2.50/$7.50)
|
|
820
|
+
(re-verified 2026-08-15 against api.deepinfra.com/models/: qwen3.8-max's US
|
|
821
|
+
re-host rates are unchanged — $1.65/$4.951, cache read 0.1248× input =
|
|
822
|
+
$0.206. DeepInfra also began serving `Qwen/Qwen3.8-2.4T-A95B` on
|
|
823
|
+
2026-08-12; by DeepInfra's own description it is the OPEN-WEIGHT variant of
|
|
824
|
+
Qwen3.8 Max (2.4T MoE, 95B active), 262,144 ctx / 131,072 out, $2/$6 with
|
|
825
|
+
cache read 0.1× input. NOT added: it is the same tier as qwen3.8-max, which
|
|
826
|
+
the catalog already offers, and it costs MORE on that very host ($2/$6 vs
|
|
827
|
+
$1.65/$4.951), so nothing would ever select it — and carrying both would
|
|
828
|
+
put two selectable Alibaba flagships in one family. Revisit only if Alibaba
|
|
829
|
+
publishes it as a distinct first-party DashScope model id.)
|
|
820
830
|
- Zhipu: https://docs.z.ai/guides/overview/pricing (unchanged; glm-5.2 is
|
|
821
831
|
the newest — "GLM-5.3/5.5" rumors have no released ids as of 2026-07-28)
|
|
822
832
|
|
package/dist/models.d.ts
CHANGED
|
@@ -101,6 +101,16 @@ import type { ModelDefinition } from './types.js';
|
|
|
101
101
|
* 125%) — the EN edition of that doc simply lags (zero qwen3.8 mentions),
|
|
102
102
|
* which an earlier pass misread as "excepted/console-only". qwen3.7-max
|
|
103
103
|
* still runs its 50%-off promo — billed here at list, $2.50/$7.50)
|
|
104
|
+
* (re-verified 2026-08-15 against api.deepinfra.com/models/: qwen3.8-max's US
|
|
105
|
+
* re-host rates are unchanged — $1.65/$4.951, cache read 0.1248× input =
|
|
106
|
+
* $0.206. DeepInfra also began serving `Qwen/Qwen3.8-2.4T-A95B` on
|
|
107
|
+
* 2026-08-12; by DeepInfra's own description it is the OPEN-WEIGHT variant of
|
|
108
|
+
* Qwen3.8 Max (2.4T MoE, 95B active), 262,144 ctx / 131,072 out, $2/$6 with
|
|
109
|
+
* cache read 0.1× input. NOT added: it is the same tier as qwen3.8-max, which
|
|
110
|
+
* the catalog already offers, and it costs MORE on that very host ($2/$6 vs
|
|
111
|
+
* $1.65/$4.951), so nothing would ever select it — and carrying both would
|
|
112
|
+
* put two selectable Alibaba flagships in one family. Revisit only if Alibaba
|
|
113
|
+
* publishes it as a distinct first-party DashScope model id.)
|
|
104
114
|
* - Zhipu: https://docs.z.ai/guides/overview/pricing (unchanged; glm-5.2 is
|
|
105
115
|
* the newest — "GLM-5.3/5.5" rumors have no released ids as of 2026-07-28)
|
|
106
116
|
*
|
package/dist/models.d.ts.map
CHANGED
|
@@ -1 +1 @@
|
|
|
1
|
-
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAEjD
|
|
1
|
+
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAEjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GA6GG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EAk6CnC,CAAA"}
|
package/dist/models.js
CHANGED
|
@@ -100,6 +100,16 @@
|
|
|
100
100
|
* 125%) — the EN edition of that doc simply lags (zero qwen3.8 mentions),
|
|
101
101
|
* which an earlier pass misread as "excepted/console-only". qwen3.7-max
|
|
102
102
|
* still runs its 50%-off promo — billed here at list, $2.50/$7.50)
|
|
103
|
+
* (re-verified 2026-08-15 against api.deepinfra.com/models/: qwen3.8-max's US
|
|
104
|
+
* re-host rates are unchanged — $1.65/$4.951, cache read 0.1248× input =
|
|
105
|
+
* $0.206. DeepInfra also began serving `Qwen/Qwen3.8-2.4T-A95B` on
|
|
106
|
+
* 2026-08-12; by DeepInfra's own description it is the OPEN-WEIGHT variant of
|
|
107
|
+
* Qwen3.8 Max (2.4T MoE, 95B active), 262,144 ctx / 131,072 out, $2/$6 with
|
|
108
|
+
* cache read 0.1× input. NOT added: it is the same tier as qwen3.8-max, which
|
|
109
|
+
* the catalog already offers, and it costs MORE on that very host ($2/$6 vs
|
|
110
|
+
* $1.65/$4.951), so nothing would ever select it — and carrying both would
|
|
111
|
+
* put two selectable Alibaba flagships in one family. Revisit only if Alibaba
|
|
112
|
+
* publishes it as a distinct first-party DashScope model id.)
|
|
103
113
|
* - Zhipu: https://docs.z.ai/guides/overview/pricing (unchanged; glm-5.2 is
|
|
104
114
|
* the newest — "GLM-5.3/5.5" rumors have no released ids as of 2026-07-28)
|
|
105
115
|
*
|
|
@@ -963,9 +973,12 @@ export const MODELS = [
|
|
|
963
973
|
// `FREE_TIER_MODELS[mode] === modelId` before it looks at regions, so
|
|
964
974
|
// leaving it here would widen nothing — it would just claim a free-tier
|
|
965
975
|
// relationship that no longer exists.
|
|
966
|
-
// US = DeepInfra, verified 2026-08-
|
|
976
|
+
// US = DeepInfra, verified 2026-08-15 via api.deepinfra.com/models/
|
|
967
977
|
// deepseek-ai/DeepSeek-V4-Pro. No cache-write premium (omitted → region
|
|
968
|
-
// input rate).
|
|
978
|
+
// input rate). DeepInfra also serves the dated `-0813` snapshot (the
|
|
979
|
+
// official release superseding the preview weights the bare id carries) at
|
|
980
|
+
// IDENTICAL rates — so bumping the molecule-dev us modelMap to it moves no
|
|
981
|
+
// price here; it is a weights decision, not a pricing one.
|
|
969
982
|
regionPricing: {
|
|
970
983
|
us: { inputPricePerMTok: 1.3, outputPricePerMTok: 2.6, cacheReadPricePerMTok: 0.1 },
|
|
971
984
|
},
|
package/package.json
CHANGED