@prestyj/core 5.16.1 → 5.16.2
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/{chunk-OUE2GRO6.js → chunk-3MQNCB44.js} +80 -28
- package/dist/{chunk-OUE2GRO6.js.map → chunk-3MQNCB44.js.map} +1 -1
- package/dist/index.cjs +138 -54
- package/dist/index.cjs.map +1 -1
- package/dist/index.d.cts +3 -3
- package/dist/index.d.ts +3 -3
- package/dist/index.js +60 -28
- package/dist/index.js.map +1 -1
- package/dist/model-registry.cjs +79 -27
- package/dist/model-registry.cjs.map +1 -1
- package/dist/model-registry.d.cts +2 -3
- package/dist/model-registry.d.ts +2 -3
- package/dist/model-registry.js +1 -1
- package/package.json +2 -2
|
@@ -1972,6 +1972,29 @@ var MODELS = [
|
|
|
1972
1972
|
maxThinkingLevel: "high"
|
|
1973
1973
|
},
|
|
1974
1974
|
// ── OpenAI (Codex) ─────────────────────────────────────
|
|
1975
|
+
{
|
|
1976
|
+
// GPT-6 Astra — "Our most capable model for complex, demanding work."
|
|
1977
|
+
// (Codex catalog priority 1, listed for every ChatGPT plan, requires a
|
|
1978
|
+
// Codex client >= 0.153.0 — see CODEX_CLIENT_VERSION). Same split as 5.6:
|
|
1979
|
+
// 1.05M on the public Responses API, 272K on the ChatGPT OAuth route
|
|
1980
|
+
// (openai/codex models.json, `gpt-6-astra`). Reasoning ladder low → medium
|
|
1981
|
+
// → high → xhigh → max → ultra; `ultra` is the Codex orchestration preset
|
|
1982
|
+
// (multi_agent v2) and is Codex-only — the public API tops out at `max`.
|
|
1983
|
+
// Note: through a plain API key OpenAI requires the Responses API for tool
|
|
1984
|
+
// calling on Astra, so the Chat Completions path is text-only; the OAuth
|
|
1985
|
+
// Codex route is the supported way to use it as an agent.
|
|
1986
|
+
id: "gpt-6-astra",
|
|
1987
|
+
name: "GPT-6 Astra",
|
|
1988
|
+
provider: "openai",
|
|
1989
|
+
contextWindow: 105e4,
|
|
1990
|
+
codexContextWindow: 272e3,
|
|
1991
|
+
maxOutputTokens: 128e3,
|
|
1992
|
+
supportsThinking: true,
|
|
1993
|
+
supportsImages: true,
|
|
1994
|
+
supportsVideo: false,
|
|
1995
|
+
costTier: "high",
|
|
1996
|
+
maxThinkingLevel: "ultra"
|
|
1997
|
+
},
|
|
1975
1998
|
// GPT-5.6 family — three agentic coding tiers launched July 2026. The public
|
|
1976
1999
|
// Responses API advertises a 1.05M context window; OpenAI's Codex product
|
|
1977
2000
|
// catalog advertises 272K on the ChatGPT OAuth route (corrected from the
|
|
@@ -2025,24 +2048,11 @@ var MODELS = [
|
|
|
2025
2048
|
costTier: "low",
|
|
2026
2049
|
maxThinkingLevel: "max"
|
|
2027
2050
|
},
|
|
2028
|
-
{
|
|
2029
|
-
id: "gpt-5.5",
|
|
2030
|
-
name: "GPT-5.5",
|
|
2031
|
-
provider: "openai",
|
|
2032
|
-
contextWindow: 105e4,
|
|
2033
|
-
codexContextWindow: 272e3,
|
|
2034
|
-
maxOutputTokens: 128e3,
|
|
2035
|
-
supportsThinking: true,
|
|
2036
|
-
supportsImages: true,
|
|
2037
|
-
supportsVideo: false,
|
|
2038
|
-
costTier: "high",
|
|
2039
|
-
maxThinkingLevel: "xhigh"
|
|
2040
|
-
},
|
|
2041
2051
|
// ── Sakana (Fugu) ──────────────────────────────────────
|
|
2042
2052
|
// Sakana Fugu is a multi-agent system surfaced as a standard LLM via the
|
|
2043
2053
|
// OpenAI-compatible Sakana API (https://api.sakana.ai/v1). Both models take
|
|
2044
|
-
// text + image input
|
|
2045
|
-
//
|
|
2054
|
+
// text + image input. Plain Fugu stops at xhigh; Ultra v1.1 also supports max.
|
|
2055
|
+
// `fugu` routes across all providers; `fugu-ultra` is
|
|
2046
2056
|
// the heavier tier (may need larger client timeouts on complex tasks).
|
|
2047
2057
|
{
|
|
2048
2058
|
id: "fugu",
|
|
@@ -2066,7 +2076,8 @@ var MODELS = [
|
|
|
2066
2076
|
supportsImages: true,
|
|
2067
2077
|
supportsVideo: false,
|
|
2068
2078
|
costTier: "high",
|
|
2069
|
-
|
|
2079
|
+
// The rolling alias now serves v1.1, which adds a distinct max effort.
|
|
2080
|
+
maxThinkingLevel: "max"
|
|
2070
2081
|
},
|
|
2071
2082
|
// ── xAI (Grok) ─────────────────────────────────────────
|
|
2072
2083
|
// Grok 4.6 (released 2026-08-12) is xAI's flagship for coding, agentic tasks,
|
|
@@ -2120,6 +2131,34 @@ var MODELS = [
|
|
|
2120
2131
|
costTier: "low",
|
|
2121
2132
|
maxThinkingLevel: "high"
|
|
2122
2133
|
},
|
|
2134
|
+
// Keep 3.1 Flash Lite first for the working OAuth default and fast-model routing.
|
|
2135
|
+
// New GA models are opt-in; Code Assist access varies by account.
|
|
2136
|
+
{
|
|
2137
|
+
id: "gemini-3.8-flash",
|
|
2138
|
+
name: "Gemini 3.8 Flash",
|
|
2139
|
+
provider: "gemini",
|
|
2140
|
+
contextWindow: 1048576,
|
|
2141
|
+
maxOutputTokens: 65536,
|
|
2142
|
+
supportsThinking: true,
|
|
2143
|
+
supportsImages: true,
|
|
2144
|
+
supportsVideo: true,
|
|
2145
|
+
maxVideoBytes: 20 * 1024 * 1024,
|
|
2146
|
+
costTier: "low",
|
|
2147
|
+
maxThinkingLevel: "high"
|
|
2148
|
+
},
|
|
2149
|
+
{
|
|
2150
|
+
id: "gemini-3.5-flash-lite",
|
|
2151
|
+
name: "Gemini 3.5 Flash Lite",
|
|
2152
|
+
provider: "gemini",
|
|
2153
|
+
contextWindow: 1048576,
|
|
2154
|
+
maxOutputTokens: 65536,
|
|
2155
|
+
supportsThinking: true,
|
|
2156
|
+
supportsImages: true,
|
|
2157
|
+
supportsVideo: true,
|
|
2158
|
+
maxVideoBytes: 20 * 1024 * 1024,
|
|
2159
|
+
costTier: "low",
|
|
2160
|
+
maxThinkingLevel: "high"
|
|
2161
|
+
},
|
|
2123
2162
|
{
|
|
2124
2163
|
// Gemini 3.7 Flash (released 2026-08-13) — Google's most capable Flash for
|
|
2125
2164
|
// coding, agents, and multi-step execution; GA-stable on the Gemini API as
|
|
@@ -2127,7 +2166,7 @@ var MODELS = [
|
|
|
2127
2166
|
// Sent over our Code Assist (OAuth) transport ahead of gemini-cli — upstream
|
|
2128
2167
|
// hasn't listed 3.7 yet (google-gemini/gemini-cli#28802, still open) — so
|
|
2129
2168
|
// free/personal accounts 404 (entitlement-gated) while Code Assist
|
|
2130
|
-
// Standard/Enterprise accounts get it.
|
|
2169
|
+
// Standard/Enterprise accounts get it. Kept after the working flash-lite:
|
|
2131
2170
|
// getFastModel picks the first low-tier entry, and flash-lite is the one
|
|
2132
2171
|
// that works on every account.
|
|
2133
2172
|
id: "gemini-3.7-flash",
|
|
@@ -2319,21 +2358,19 @@ var MODELS = [
|
|
|
2319
2358
|
{
|
|
2320
2359
|
// `deepseek-v4-pro` now serves DeepSeek-V4-Pro-0813 (released 2026-08-13,
|
|
2321
2360
|
// first STABLE V4 Pro — supersedes the April preview; calling name
|
|
2322
|
-
// unchanged, same 1.6T/49B MoE). 1M context,
|
|
2323
|
-
//
|
|
2324
|
-
//
|
|
2325
|
-
// price band rather than the preview's top band.
|
|
2361
|
+
// unchanged, same 1.6T/49B MoE). 1M context, text-only, low/high/max effort.
|
|
2362
|
+
// Docs abbreviate output as 384K; use the same conservative 384,000-token
|
|
2363
|
+
// application cap across V4 models rather than mixing decimal/binary units.
|
|
2326
2364
|
id: "deepseek-v4-pro",
|
|
2327
2365
|
name: "DeepSeek V4 Pro",
|
|
2328
2366
|
provider: "deepseek",
|
|
2329
2367
|
contextWindow: 1048576,
|
|
2330
|
-
maxOutputTokens:
|
|
2368
|
+
maxOutputTokens: 384e3,
|
|
2331
2369
|
supportsThinking: true,
|
|
2332
2370
|
supportsImages: false,
|
|
2333
2371
|
supportsVideo: false,
|
|
2334
2372
|
costTier: "medium",
|
|
2335
|
-
|
|
2336
|
-
maxThinkingLevel: "xhigh"
|
|
2373
|
+
maxThinkingLevel: "max"
|
|
2337
2374
|
},
|
|
2338
2375
|
{
|
|
2339
2376
|
id: "deepseek-v4-flash",
|
|
@@ -2345,7 +2382,20 @@ var MODELS = [
|
|
|
2345
2382
|
supportsImages: false,
|
|
2346
2383
|
supportsVideo: false,
|
|
2347
2384
|
costTier: "low",
|
|
2348
|
-
maxThinkingLevel: "
|
|
2385
|
+
maxThinkingLevel: "max"
|
|
2386
|
+
},
|
|
2387
|
+
// Opt-in experimental vision sibling; never replaces the stable summary model.
|
|
2388
|
+
{
|
|
2389
|
+
id: "deepseek-v4-flash-vision-exp",
|
|
2390
|
+
name: "DeepSeek V4 Flash Vision (Experimental)",
|
|
2391
|
+
provider: "deepseek",
|
|
2392
|
+
contextWindow: 1048576,
|
|
2393
|
+
maxOutputTokens: 384e3,
|
|
2394
|
+
supportsThinking: true,
|
|
2395
|
+
supportsImages: true,
|
|
2396
|
+
supportsVideo: false,
|
|
2397
|
+
costTier: "low",
|
|
2398
|
+
maxThinkingLevel: "max"
|
|
2349
2399
|
},
|
|
2350
2400
|
// ── OpenRouter ─────────────────────────────────────────
|
|
2351
2401
|
{
|
|
@@ -2355,8 +2405,10 @@ var MODELS = [
|
|
|
2355
2405
|
contextWindow: 1e6,
|
|
2356
2406
|
maxOutputTokens: 65536,
|
|
2357
2407
|
supportsThinking: true,
|
|
2358
|
-
supportsImages:
|
|
2359
|
-
supportsVideo:
|
|
2408
|
+
supportsImages: true,
|
|
2409
|
+
supportsVideo: true,
|
|
2410
|
+
// Practical inline-payload cap, not an asserted provider maximum.
|
|
2411
|
+
maxVideoBytes: 20 * 1024 * 1024,
|
|
2360
2412
|
costTier: "medium",
|
|
2361
2413
|
maxThinkingLevel: "high"
|
|
2362
2414
|
},
|
|
@@ -2559,4 +2611,4 @@ export {
|
|
|
2559
2611
|
getSummaryModel,
|
|
2560
2612
|
getFastModel
|
|
2561
2613
|
};
|
|
2562
|
-
//# sourceMappingURL=chunk-
|
|
2614
|
+
//# sourceMappingURL=chunk-3MQNCB44.js.map
|