@molecule/api-resource-ai-models 1.5.0 → 1.5.1

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/README.md CHANGED
@@ -3,7 +3,7 @@ AUTO-GENERATED — DO NOT EDIT THIS FILE.
3
3
  Generated by `mlcl sync-docs` from the package's src/index.ts JSDoc + mlcl/registry.json.
4
4
  Edits here are overwritten on the next commit (molecule's pre-commit hook regenerates).
5
5
  To change this document, edit the module-level JSDoc in src/index.ts.
6
- Generated: 2026-09-02T19:41:05.247Z
6
+ Generated: 2026-09-06T11:38:29.688Z
7
7
  -->
8
8
 
9
9
  # @molecule/api-resource-ai-models
@@ -861,7 +861,7 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
861
861
  grok-4.3 still served at $1.25/$2.50 with the bigger 1M window;
862
862
  grok-code-fast-1 no longer listed — retires 2026-08-15)
863
863
  - DeepSeek: https://api-docs.deepseek.com/quick_start/pricing (verified
864
- 2026-08-31; legacy deepseek-chat/-reasoner ids fully retired 2026-07-24 —
864
+ 2026-09-06; legacy deepseek-chat/-reasoner ids fully retired 2026-07-24 —
865
865
  never in this catalog. The V4-Pro-GA price RISE effective 2026-08-16T16:00Z
866
866
  has LANDED and is folded into the base fields, along with the peak-hour 2×
867
867
  the same card introduced; every rate re-read on the card 2026-08-31 and
@@ -870,7 +870,10 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
870
870
  where on 2026-08-18 it carried no day qualifier at all, so the windows are
871
871
  `daysOfWeekUtc`-restricted rather than daily. `deepseek-v4-flash-vision-exp`
872
872
  also appears on the card at flash's rates: EXPERIMENTAL and vision-only-new,
873
- deliberately not catalogued.)
873
+ deliberately not catalogued. The 2026-09-06 re-read found every native rate
874
+ unchanged; what moved was the US re-host — DeepInfra cut
875
+ `DeepSeek-V4-Flash-0731` from $0.08 to $0.06 input, cache read $0.016 to
876
+ $0.015, output held at $0.18. See that entry's `regionPricing`.)
874
877
  - Moonshot: https://platform.kimi.ai/docs/models + DeepInfra's model API for
875
878
  the US re-host (kimi-k3 flagship 2026-07-16
876
879
  — 2.8T MoE, 1M ctx, $3/$15 — NOT added: thinking is forced-on with
package/dist/models.d.ts CHANGED
@@ -129,7 +129,7 @@ import type { ModelDefinition } from './types.js';
129
129
  * grok-4.3 still served at $1.25/$2.50 with the bigger 1M window;
130
130
  * grok-code-fast-1 no longer listed — retires 2026-08-15)
131
131
  * - DeepSeek: https://api-docs.deepseek.com/quick_start/pricing (verified
132
- * 2026-08-31; legacy deepseek-chat/-reasoner ids fully retired 2026-07-24 —
132
+ * 2026-09-06; legacy deepseek-chat/-reasoner ids fully retired 2026-07-24 —
133
133
  * never in this catalog. The V4-Pro-GA price RISE effective 2026-08-16T16:00Z
134
134
  * has LANDED and is folded into the base fields, along with the peak-hour 2×
135
135
  * the same card introduced; every rate re-read on the card 2026-08-31 and
@@ -138,7 +138,10 @@ import type { ModelDefinition } from './types.js';
138
138
  * where on 2026-08-18 it carried no day qualifier at all, so the windows are
139
139
  * `daysOfWeekUtc`-restricted rather than daily. `deepseek-v4-flash-vision-exp`
140
140
  * also appears on the card at flash's rates: EXPERIMENTAL and vision-only-new,
141
- * deliberately not catalogued.)
141
+ * deliberately not catalogued. The 2026-09-06 re-read found every native rate
142
+ * unchanged; what moved was the US re-host — DeepInfra cut
143
+ * `DeepSeek-V4-Flash-0731` from $0.08 to $0.06 input, cache read $0.016 to
144
+ * $0.015, output held at $0.18. See that entry's `regionPricing`.)
142
145
  * - Moonshot: https://platform.kimi.ai/docs/models + DeepInfra's model API for
143
146
  * the US re-host (kimi-k3 flagship 2026-07-16
144
147
  * — 2.8T MoE, 1M ctx, $3/$15 — NOT added: thinking is forced-on with
@@ -1 +1 @@
1
- {"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAQjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAiMG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EAusDnC,CAAA"}
1
+ {"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAQjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GAoMG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EA6sDnC,CAAA"}
package/dist/models.js CHANGED
@@ -133,7 +133,7 @@ const WEEKDAYS_UTC = [1, 2, 3, 4, 5];
133
133
  * grok-4.3 still served at $1.25/$2.50 with the bigger 1M window;
134
134
  * grok-code-fast-1 no longer listed — retires 2026-08-15)
135
135
  * - DeepSeek: https://api-docs.deepseek.com/quick_start/pricing (verified
136
- * 2026-08-31; legacy deepseek-chat/-reasoner ids fully retired 2026-07-24 —
136
+ * 2026-09-06; legacy deepseek-chat/-reasoner ids fully retired 2026-07-24 —
137
137
  * never in this catalog. The V4-Pro-GA price RISE effective 2026-08-16T16:00Z
138
138
  * has LANDED and is folded into the base fields, along with the peak-hour 2×
139
139
  * the same card introduced; every rate re-read on the card 2026-08-31 and
@@ -142,7 +142,10 @@ const WEEKDAYS_UTC = [1, 2, 3, 4, 5];
142
142
  * where on 2026-08-18 it carried no day qualifier at all, so the windows are
143
143
  * `daysOfWeekUtc`-restricted rather than daily. `deepseek-v4-flash-vision-exp`
144
144
  * also appears on the card at flash's rates: EXPERIMENTAL and vision-only-new,
145
- * deliberately not catalogued.)
145
+ * deliberately not catalogued. The 2026-09-06 re-read found every native rate
146
+ * unchanged; what moved was the US re-host — DeepInfra cut
147
+ * `DeepSeek-V4-Flash-0731` from $0.08 to $0.06 input, cache read $0.016 to
148
+ * $0.015, output held at $0.18. See that entry's `regionPricing`.)
146
149
  * - Moonshot: https://platform.kimi.ai/docs/models + DeepInfra's model API for
147
150
  * the US re-host (kimi-k3 flagship 2026-07-16
148
151
  * — 2.8T MoE, 1M ctx, $3/$15 — NOT added: thinking is forced-on with
@@ -1274,23 +1277,29 @@ export const MODELS = [
1274
1277
  // US (DeepInfra) DEFAULT as of 2026-08-16 — flipped from CN when DeepSeek's
1275
1278
  // rise landed (owner decision 2026-08-14). CN was cheaper on real traffic
1276
1279
  // only because of its cache reads; the rise takes those from $0.0028 to
1277
- // $0.007 (peak $0.014) against DeepInfra's flat $0.016, which is no longer
1278
- // enough to carry the 1.6-3.1x it now loses on fresh input and output. On
1280
+ // $0.007 (peak $0.014) against DeepInfra's flat $0.015, which is no longer
1281
+ // enough to carry the 3.7x CN now loses on BOTH fresh input (0.22 vs 0.06)
1282
+ // and output (0.66 vs 0.18). On
1279
1283
  // the agentic mix this model actually serves (~94% cache hits) US is
1280
- // cheaper at EVERY hour: 0.114c/turn flat vs 0.152c off-peak and 0.303c at
1284
+ // cheaper at EVERY hour: 0.103c/turn flat vs 0.152c off-peak and 0.303c at
1281
1285
  // peak. It is also flat-rate, so free-tier cost stops varying by Beijing
1282
1286
  // business hours. Re-derive if the cache-hit ratio drops much below ~90%,
1283
1287
  // where CN's cheaper reads start winning again. This deliberately splits
1284
1288
  // the plan/execute pair across regions — Pro stays CN because its US
1285
1289
  // re-host is ~2.3x its own native rate even after the rise.
1286
1290
  regions: ['us', 'cn'],
1287
- // US = DeepInfra, verified 2026-08-13 against the id the bond actually
1291
+ // US = DeepInfra, verified 2026-09-06 against the id the bond actually
1288
1292
  // sends: `deepseek-ai/DeepSeek-V4-Flash-0731`, the official release that
1289
1293
  // supersedes the preview weights still served under the un-dated id
1290
- // (cents_per_input_token 0.000008, cents_per_output_token 0.000018,
1291
- // rate_per_input_token_cached 0.2 → cache read = 0.2 × input).
1294
+ // (cents_per_input_token 0.000006, cents_per_output_token 0.000018,
1295
+ // rate_per_input_token_cached 0.25 → cache read = 0.25 × input).
1296
+ // DeepInfra CUT the 0731 input rate 0.08 → 0.06 (cache read 0.016 →
1297
+ // 0.015); output held at 0.18. The un-dated `DeepSeek-V4-Flash` id moved
1298
+ // the other way (0.09 input, cache 0.2× = 0.018), so the two ids no longer
1299
+ // price alike — this entry tracks the dated one the modelMap sends, and the
1300
+ // gap widens the case for the US default this model already carries.
1292
1301
  regionPricing: {
1293
- us: { inputPricePerMTok: 0.08, outputPricePerMTok: 0.18, cacheReadPricePerMTok: 0.016 },
1302
+ us: { inputPricePerMTok: 0.06, outputPricePerMTok: 0.18, cacheReadPricePerMTok: 0.015 },
1294
1303
  },
1295
1304
  // Peak-hour surcharge live since 2026-08-16T16:00Z, Mon-Fri (see
1296
1305
  // deepseek-v4-pro). It applies to the NATIVE CN card only — this model
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@molecule/api-resource-ai-models",
3
- "version": "1.5.0",
3
+ "version": "1.5.1",
4
4
  "description": "AI model catalog — server-side source of truth plus an authentication-gated discovery endpoint",
5
5
  "type": "module",
6
6
  "main": "dist/index.js",