@molecule/api-resource-ai-models 1.5.0 → 1.5.1
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +6 -3
- package/dist/models.d.ts +5 -2
- package/dist/models.d.ts.map +1 -1
- package/dist/models.js +18 -9
- package/package.json +1 -1
package/README.md
CHANGED
|
@@ -3,7 +3,7 @@ AUTO-GENERATED — DO NOT EDIT THIS FILE.
|
|
|
3
3
|
Generated by `mlcl sync-docs` from the package's src/index.ts JSDoc + mlcl/registry.json.
|
|
4
4
|
Edits here are overwritten on the next commit (molecule's pre-commit hook regenerates).
|
|
5
5
|
To change this document, edit the module-level JSDoc in src/index.ts.
|
|
6
|
-
Generated: 2026-09-
|
|
6
|
+
Generated: 2026-09-06T11:38:29.688Z
|
|
7
7
|
-->
|
|
8
8
|
|
|
9
9
|
# @molecule/api-resource-ai-models
|
|
@@ -861,7 +861,7 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
|
|
|
861
861
|
grok-4.3 still served at $1.25/$2.50 with the bigger 1M window;
|
|
862
862
|
grok-code-fast-1 no longer listed — retires 2026-08-15)
|
|
863
863
|
- DeepSeek: https://api-docs.deepseek.com/quick_start/pricing (verified
|
|
864
|
-
2026-
|
|
864
|
+
2026-09-06; legacy deepseek-chat/-reasoner ids fully retired 2026-07-24 —
|
|
865
865
|
never in this catalog. The V4-Pro-GA price RISE effective 2026-08-16T16:00Z
|
|
866
866
|
has LANDED and is folded into the base fields, along with the peak-hour 2×
|
|
867
867
|
the same card introduced; every rate re-read on the card 2026-08-31 and
|
|
@@ -870,7 +870,10 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
|
|
|
870
870
|
where on 2026-08-18 it carried no day qualifier at all, so the windows are
|
|
871
871
|
`daysOfWeekUtc`-restricted rather than daily. `deepseek-v4-flash-vision-exp`
|
|
872
872
|
also appears on the card at flash's rates: EXPERIMENTAL and vision-only-new,
|
|
873
|
-
deliberately not catalogued.
|
|
873
|
+
deliberately not catalogued. The 2026-09-06 re-read found every native rate
|
|
874
|
+
unchanged; what moved was the US re-host — DeepInfra cut
|
|
875
|
+
`DeepSeek-V4-Flash-0731` from $0.08 to $0.06 input, cache read $0.016 to
|
|
876
|
+
$0.015, output held at $0.18. See that entry's `regionPricing`.)
|
|
874
877
|
- Moonshot: https://platform.kimi.ai/docs/models + DeepInfra's model API for
|
|
875
878
|
the US re-host (kimi-k3 flagship 2026-07-16
|
|
876
879
|
— 2.8T MoE, 1M ctx, $3/$15 — NOT added: thinking is forced-on with
|
package/dist/models.d.ts
CHANGED
|
@@ -129,7 +129,7 @@ import type { ModelDefinition } from './types.js';
|
|
|
129
129
|
* grok-4.3 still served at $1.25/$2.50 with the bigger 1M window;
|
|
130
130
|
* grok-code-fast-1 no longer listed — retires 2026-08-15)
|
|
131
131
|
* - DeepSeek: https://api-docs.deepseek.com/quick_start/pricing (verified
|
|
132
|
-
* 2026-
|
|
132
|
+
* 2026-09-06; legacy deepseek-chat/-reasoner ids fully retired 2026-07-24 —
|
|
133
133
|
* never in this catalog. The V4-Pro-GA price RISE effective 2026-08-16T16:00Z
|
|
134
134
|
* has LANDED and is folded into the base fields, along with the peak-hour 2×
|
|
135
135
|
* the same card introduced; every rate re-read on the card 2026-08-31 and
|
|
@@ -138,7 +138,10 @@ import type { ModelDefinition } from './types.js';
|
|
|
138
138
|
* where on 2026-08-18 it carried no day qualifier at all, so the windows are
|
|
139
139
|
* `daysOfWeekUtc`-restricted rather than daily. `deepseek-v4-flash-vision-exp`
|
|
140
140
|
* also appears on the card at flash's rates: EXPERIMENTAL and vision-only-new,
|
|
141
|
-
* deliberately not catalogued.
|
|
141
|
+
* deliberately not catalogued. The 2026-09-06 re-read found every native rate
|
|
142
|
+
* unchanged; what moved was the US re-host — DeepInfra cut
|
|
143
|
+
* `DeepSeek-V4-Flash-0731` from $0.08 to $0.06 input, cache read $0.016 to
|
|
144
|
+
* $0.015, output held at $0.18. See that entry's `regionPricing`.)
|
|
142
145
|
* - Moonshot: https://platform.kimi.ai/docs/models + DeepInfra's model API for
|
|
143
146
|
* the US re-host (kimi-k3 flagship 2026-07-16
|
|
144
147
|
* — 2.8T MoE, 1M ctx, $3/$15 — NOT added: thinking is forced-on with
|
package/dist/models.d.ts.map
CHANGED
|
@@ -1 +1 @@
|
|
|
1
|
-
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAQjD
|
|
1
|
+
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAQjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAoMG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EA6sDnC,CAAA"}
|
package/dist/models.js
CHANGED
|
@@ -133,7 +133,7 @@ const WEEKDAYS_UTC = [1, 2, 3, 4, 5];
|
|
|
133
133
|
* grok-4.3 still served at $1.25/$2.50 with the bigger 1M window;
|
|
134
134
|
* grok-code-fast-1 no longer listed — retires 2026-08-15)
|
|
135
135
|
* - DeepSeek: https://api-docs.deepseek.com/quick_start/pricing (verified
|
|
136
|
-
* 2026-
|
|
136
|
+
* 2026-09-06; legacy deepseek-chat/-reasoner ids fully retired 2026-07-24 —
|
|
137
137
|
* never in this catalog. The V4-Pro-GA price RISE effective 2026-08-16T16:00Z
|
|
138
138
|
* has LANDED and is folded into the base fields, along with the peak-hour 2×
|
|
139
139
|
* the same card introduced; every rate re-read on the card 2026-08-31 and
|
|
@@ -142,7 +142,10 @@ const WEEKDAYS_UTC = [1, 2, 3, 4, 5];
|
|
|
142
142
|
* where on 2026-08-18 it carried no day qualifier at all, so the windows are
|
|
143
143
|
* `daysOfWeekUtc`-restricted rather than daily. `deepseek-v4-flash-vision-exp`
|
|
144
144
|
* also appears on the card at flash's rates: EXPERIMENTAL and vision-only-new,
|
|
145
|
-
* deliberately not catalogued.
|
|
145
|
+
* deliberately not catalogued. The 2026-09-06 re-read found every native rate
|
|
146
|
+
* unchanged; what moved was the US re-host — DeepInfra cut
|
|
147
|
+
* `DeepSeek-V4-Flash-0731` from $0.08 to $0.06 input, cache read $0.016 to
|
|
148
|
+
* $0.015, output held at $0.18. See that entry's `regionPricing`.)
|
|
146
149
|
* - Moonshot: https://platform.kimi.ai/docs/models + DeepInfra's model API for
|
|
147
150
|
* the US re-host (kimi-k3 flagship 2026-07-16
|
|
148
151
|
* — 2.8T MoE, 1M ctx, $3/$15 — NOT added: thinking is forced-on with
|
|
@@ -1274,23 +1277,29 @@ export const MODELS = [
|
|
|
1274
1277
|
// US (DeepInfra) DEFAULT as of 2026-08-16 — flipped from CN when DeepSeek's
|
|
1275
1278
|
// rise landed (owner decision 2026-08-14). CN was cheaper on real traffic
|
|
1276
1279
|
// only because of its cache reads; the rise takes those from $0.0028 to
|
|
1277
|
-
// $0.007 (peak $0.014) against DeepInfra's flat $0.
|
|
1278
|
-
// enough to carry the
|
|
1280
|
+
// $0.007 (peak $0.014) against DeepInfra's flat $0.015, which is no longer
|
|
1281
|
+
// enough to carry the 3.7x CN now loses on BOTH fresh input (0.22 vs 0.06)
|
|
1282
|
+
// and output (0.66 vs 0.18). On
|
|
1279
1283
|
// the agentic mix this model actually serves (~94% cache hits) US is
|
|
1280
|
-
// cheaper at EVERY hour: 0.
|
|
1284
|
+
// cheaper at EVERY hour: 0.103c/turn flat vs 0.152c off-peak and 0.303c at
|
|
1281
1285
|
// peak. It is also flat-rate, so free-tier cost stops varying by Beijing
|
|
1282
1286
|
// business hours. Re-derive if the cache-hit ratio drops much below ~90%,
|
|
1283
1287
|
// where CN's cheaper reads start winning again. This deliberately splits
|
|
1284
1288
|
// the plan/execute pair across regions — Pro stays CN because its US
|
|
1285
1289
|
// re-host is ~2.3x its own native rate even after the rise.
|
|
1286
1290
|
regions: ['us', 'cn'],
|
|
1287
|
-
// US = DeepInfra, verified 2026-
|
|
1291
|
+
// US = DeepInfra, verified 2026-09-06 against the id the bond actually
|
|
1288
1292
|
// sends: `deepseek-ai/DeepSeek-V4-Flash-0731`, the official release that
|
|
1289
1293
|
// supersedes the preview weights still served under the un-dated id
|
|
1290
|
-
// (cents_per_input_token 0.
|
|
1291
|
-
// rate_per_input_token_cached 0.
|
|
1294
|
+
// (cents_per_input_token 0.000006, cents_per_output_token 0.000018,
|
|
1295
|
+
// rate_per_input_token_cached 0.25 → cache read = 0.25 × input).
|
|
1296
|
+
// DeepInfra CUT the 0731 input rate 0.08 → 0.06 (cache read 0.016 →
|
|
1297
|
+
// 0.015); output held at 0.18. The un-dated `DeepSeek-V4-Flash` id moved
|
|
1298
|
+
// the other way (0.09 input, cache 0.2× = 0.018), so the two ids no longer
|
|
1299
|
+
// price alike — this entry tracks the dated one the modelMap sends, and the
|
|
1300
|
+
// gap widens the case for the US default this model already carries.
|
|
1292
1301
|
regionPricing: {
|
|
1293
|
-
us: { inputPricePerMTok: 0.
|
|
1302
|
+
us: { inputPricePerMTok: 0.06, outputPricePerMTok: 0.18, cacheReadPricePerMTok: 0.015 },
|
|
1294
1303
|
},
|
|
1295
1304
|
// Peak-hour surcharge live since 2026-08-16T16:00Z, Mon-Fri (see
|
|
1296
1305
|
// deepseek-v4-pro). It applies to the NATIVE CN card only — this model
|
package/package.json
CHANGED