@molecule/api-resource-ai-models 1.6.2 → 1.6.3
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +13 -1
- package/dist/models.d.ts +12 -0
- package/dist/models.d.ts.map +1 -1
- package/dist/models.js +39 -0
- package/package.json +1 -1
package/README.md
CHANGED
|
@@ -3,7 +3,7 @@ AUTO-GENERATED — DO NOT EDIT THIS FILE.
|
|
|
3
3
|
Generated by `mlcl sync-docs` from the package's src/index.ts JSDoc + mlcl/registry.json.
|
|
4
4
|
Edits here are overwritten on the next commit (molecule's pre-commit hook regenerates).
|
|
5
5
|
To change this document, edit the module-level JSDoc in src/index.ts.
|
|
6
|
-
Generated: 2026-09-
|
|
6
|
+
Generated: 2026-09-21T16:34:40.015Z
|
|
7
7
|
-->
|
|
8
8
|
|
|
9
9
|
# @molecule/api-resource-ai-models
|
|
@@ -828,6 +828,18 @@ Sources (verified 2026-07-28; OpenAI re-verified 2026-07-31 after the
|
|
|
828
828
|
never under-charges when it lapses, per KNOWN_DIVERGENCES; -terra and
|
|
829
829
|
-luna unchanged and matching the page; cache-read 0.1× and cache-write
|
|
830
830
|
1.25× are now published first-party for all three tiers)
|
|
831
|
+
(re-verified 2026-09-21 on the pricing page: every cataloged OpenAI entry
|
|
832
|
+
is unchanged, and the -sol promo footnote still reads "available at least
|
|
833
|
+
through November 21, 2026". gpt-6-astra (released 2026-09-04) is listed at
|
|
834
|
+
$10/$50, cached input $1, cache writes $12.50 — and, like the Gemini/Grok
|
|
835
|
+
tiers, a long-context band above 272K prompt tokens that reprices the WHOLE
|
|
836
|
+
request ($20/$75, cached $2, cache writes $25): 2× input/cache, 1.5×
|
|
837
|
+
output. The catalog's price fields are flat per-MTok rates with no
|
|
838
|
+
context-band dimension, so that band is NOT modeled, exactly as the
|
|
839
|
+
> 200K tiers above are not. It is moot for now — astra is deliberately NOT
|
|
840
|
+
> in the catalog because it cannot serve a tool-carrying request on the
|
|
841
|
+
> bond's /v1/chat/completions endpoint; see the DO NOT ADD block above the
|
|
842
|
+
> OpenAI entries for the live 400s and the condition that lifts it.)
|
|
831
843
|
- Google: https://ai.google.dev/gemini-api/docs/pricing (gemini-3.6-flash GA
|
|
832
844
|
2026-07-21 $1.50/$7.50 supersedes 3.5-flash as the agentic flagship;
|
|
833
845
|
gemini-3.1-pro-preview still the pro tier — "3.5 Pro" has NOT shipped as
|
package/dist/models.d.ts
CHANGED
|
@@ -96,6 +96,18 @@ import type { ModelDefinition } from './types.js';
|
|
|
96
96
|
* never under-charges when it lapses, per KNOWN_DIVERGENCES; -terra and
|
|
97
97
|
* -luna unchanged and matching the page; cache-read 0.1× and cache-write
|
|
98
98
|
* 1.25× are now published first-party for all three tiers)
|
|
99
|
+
* (re-verified 2026-09-21 on the pricing page: every cataloged OpenAI entry
|
|
100
|
+
* is unchanged, and the -sol promo footnote still reads "available at least
|
|
101
|
+
* through November 21, 2026". gpt-6-astra (released 2026-09-04) is listed at
|
|
102
|
+
* $10/$50, cached input $1, cache writes $12.50 — and, like the Gemini/Grok
|
|
103
|
+
* tiers, a long-context band above 272K prompt tokens that reprices the WHOLE
|
|
104
|
+
* request ($20/$75, cached $2, cache writes $25): 2× input/cache, 1.5×
|
|
105
|
+
* output. The catalog's price fields are flat per-MTok rates with no
|
|
106
|
+
* context-band dimension, so that band is NOT modeled, exactly as the
|
|
107
|
+
* >200K tiers above are not. It is moot for now — astra is deliberately NOT
|
|
108
|
+
* in the catalog because it cannot serve a tool-carrying request on the
|
|
109
|
+
* bond's /v1/chat/completions endpoint; see the DO NOT ADD block above the
|
|
110
|
+
* OpenAI entries for the live 400s and the condition that lifts it.)
|
|
99
111
|
* - Google: https://ai.google.dev/gemini-api/docs/pricing (gemini-3.6-flash GA
|
|
100
112
|
* 2026-07-21 $1.50/$7.50 supersedes 3.5-flash as the agentic flagship;
|
|
101
113
|
* gemini-3.1-pro-preview still the pro tier — "3.5 Pro" has NOT shipped as
|
package/dist/models.d.ts.map
CHANGED
|
@@ -1 +1 @@
|
|
|
1
|
-
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAQjD
|
|
1
|
+
{"version":3,"file":"models.d.ts","sourceRoot":"","sources":["../src/models.ts"],"names":[],"mappings":"AAAA;;;;;;;;GAQG;AAEH,OAAO,KAAK,EAAE,eAAe,EAAE,MAAM,YAAY,CAAA;AAQjD;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;;GA8NG;AACH,eAAO,MAAM,MAAM,EAAE,SAAS,eAAe,EAg3DnC,CAAA"}
|
package/dist/models.js
CHANGED
|
@@ -100,6 +100,18 @@ const WEEKDAYS_UTC = [1, 2, 3, 4, 5];
|
|
|
100
100
|
* never under-charges when it lapses, per KNOWN_DIVERGENCES; -terra and
|
|
101
101
|
* -luna unchanged and matching the page; cache-read 0.1× and cache-write
|
|
102
102
|
* 1.25× are now published first-party for all three tiers)
|
|
103
|
+
* (re-verified 2026-09-21 on the pricing page: every cataloged OpenAI entry
|
|
104
|
+
* is unchanged, and the -sol promo footnote still reads "available at least
|
|
105
|
+
* through November 21, 2026". gpt-6-astra (released 2026-09-04) is listed at
|
|
106
|
+
* $10/$50, cached input $1, cache writes $12.50 — and, like the Gemini/Grok
|
|
107
|
+
* tiers, a long-context band above 272K prompt tokens that reprices the WHOLE
|
|
108
|
+
* request ($20/$75, cached $2, cache writes $25): 2× input/cache, 1.5×
|
|
109
|
+
* output. The catalog's price fields are flat per-MTok rates with no
|
|
110
|
+
* context-band dimension, so that band is NOT modeled, exactly as the
|
|
111
|
+
* >200K tiers above are not. It is moot for now — astra is deliberately NOT
|
|
112
|
+
* in the catalog because it cannot serve a tool-carrying request on the
|
|
113
|
+
* bond's /v1/chat/completions endpoint; see the DO NOT ADD block above the
|
|
114
|
+
* OpenAI entries for the live 400s and the condition that lifts it.)
|
|
103
115
|
* - Google: https://ai.google.dev/gemini-api/docs/pricing (gemini-3.6-flash GA
|
|
104
116
|
* 2026-07-21 $1.50/$7.50 supersedes 3.5-flash as the agentic flagship;
|
|
105
117
|
* gemini-3.1-pro-preview still the pro tier — "3.5 Pro" has NOT shipped as
|
|
@@ -573,6 +585,33 @@ export const MODELS = [
|
|
|
573
585
|
// high|xhigh, default medium); re-verify. Long-context 2× price variants
|
|
574
586
|
// exist upstream — not modeled (same as the Gemini/Grok >200K tiers), and
|
|
575
587
|
// neither are the Batch (0.5×) or Sol "Fast mode" (2×) cards.
|
|
588
|
+
//
|
|
589
|
+
// DO NOT ADD gpt-6-astra (OpenAI's flagship since 2026-09-04) UNTIL THE BOND
|
|
590
|
+
// MOVES TO /v1/responses. It is deliberately absent, not overlooked. The
|
|
591
|
+
// openai bond posts to /v1/chat/completions, and on that endpoint this model
|
|
592
|
+
// rejects EVERY request Synthase can send, because Synthase always carries
|
|
593
|
+
// function tools (probed live 2026-09-21 on our own key):
|
|
594
|
+
// - tools + reasoning_effort low|medium|high|xhigh → 400 "Function tools
|
|
595
|
+
// with reasoning_effort are not supported for gpt-6-astra in
|
|
596
|
+
// /v1/chat/completions. To use function tools, use /v1/responses or set
|
|
597
|
+
// reasoning_effort to 'none'."
|
|
598
|
+
// - tools, field omitted → the same 400 (the model applies its own default)
|
|
599
|
+
// - tools + reasoning_effort 'none' → 400 "Unsupported value: … does not
|
|
600
|
+
// support 'none' with this model. Supported values are: 'low', 'medium',
|
|
601
|
+
// 'high', and 'xhigh'."
|
|
602
|
+
// So the gpt-5.6 workaround (`toolsRequireReasoningOff`, which pins an
|
|
603
|
+
// explicit 'none') does NOT carry over: OpenAI removed 'none' from this
|
|
604
|
+
// model's ladder, which closes the one door that made the 5.6 family usable.
|
|
605
|
+
// Nothing is wrong with the model or the account — no-tools chat and
|
|
606
|
+
// /v1/responses WITH tools both return 200. Both candidate entries were built
|
|
607
|
+
// and run through molecule-dev's verify:model-dispatch; both FAILED, so the
|
|
608
|
+
// entry was withheld rather than shipped broken (the glm-5.3-flash lesson:
|
|
609
|
+
// a catalog entry that cannot serve a Synthase-shaped turn breaks every turn
|
|
610
|
+
// on it). The freshness gate WILL keep listing it as a new-model candidate —
|
|
611
|
+
// that is correct; it becomes addable the day the bond speaks /v1/responses.
|
|
612
|
+
// Note also that the docs page advertises effort 'max', which
|
|
613
|
+
// /v1/chat/completions rejects for this model — verify the ladder against the
|
|
614
|
+
// endpoint, not the docs, when this is revisited.
|
|
576
615
|
// ---------------------------------------------------------------------------
|
|
577
616
|
{
|
|
578
617
|
id: 'gpt-5.6-sol',
|
package/package.json
CHANGED