@prestyj/core 5.16.1 → 5.16.2

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
@@ -1972,6 +1972,29 @@ var MODELS = [
1972
1972
  maxThinkingLevel: "high"
1973
1973
  },
1974
1974
  // ── OpenAI (Codex) ─────────────────────────────────────
1975
+ {
1976
+ // GPT-6 Astra — "Our most capable model for complex, demanding work."
1977
+ // (Codex catalog priority 1, listed for every ChatGPT plan, requires a
1978
+ // Codex client >= 0.153.0 — see CODEX_CLIENT_VERSION). Same split as 5.6:
1979
+ // 1.05M on the public Responses API, 272K on the ChatGPT OAuth route
1980
+ // (openai/codex models.json, `gpt-6-astra`). Reasoning ladder low → medium
1981
+ // → high → xhigh → max → ultra; `ultra` is the Codex orchestration preset
1982
+ // (multi_agent v2) and is Codex-only — the public API tops out at `max`.
1983
+ // Note: through a plain API key OpenAI requires the Responses API for tool
1984
+ // calling on Astra, so the Chat Completions path is text-only; the OAuth
1985
+ // Codex route is the supported way to use it as an agent.
1986
+ id: "gpt-6-astra",
1987
+ name: "GPT-6 Astra",
1988
+ provider: "openai",
1989
+ contextWindow: 105e4,
1990
+ codexContextWindow: 272e3,
1991
+ maxOutputTokens: 128e3,
1992
+ supportsThinking: true,
1993
+ supportsImages: true,
1994
+ supportsVideo: false,
1995
+ costTier: "high",
1996
+ maxThinkingLevel: "ultra"
1997
+ },
1975
1998
  // GPT-5.6 family — three agentic coding tiers launched July 2026. The public
1976
1999
  // Responses API advertises a 1.05M context window; OpenAI's Codex product
1977
2000
  // catalog advertises 272K on the ChatGPT OAuth route (corrected from the
@@ -2025,24 +2048,11 @@ var MODELS = [
2025
2048
  costTier: "low",
2026
2049
  maxThinkingLevel: "max"
2027
2050
  },
2028
- {
2029
- id: "gpt-5.5",
2030
- name: "GPT-5.5",
2031
- provider: "openai",
2032
- contextWindow: 105e4,
2033
- codexContextWindow: 272e3,
2034
- maxOutputTokens: 128e3,
2035
- supportsThinking: true,
2036
- supportsImages: true,
2037
- supportsVideo: false,
2038
- costTier: "high",
2039
- maxThinkingLevel: "xhigh"
2040
- },
2041
2051
  // ── Sakana (Fugu) ──────────────────────────────────────
2042
2052
  // Sakana Fugu is a multi-agent system surfaced as a standard LLM via the
2043
2053
  // OpenAI-compatible Sakana API (https://api.sakana.ai/v1). Both models take
2044
- // text + image input and only accept "high"/"xhigh" reasoning effort, so the
2045
- // top tier is `xhigh`. `fugu` routes across all providers; `fugu-ultra` is
2054
+ // text + image input. Plain Fugu stops at xhigh; Ultra v1.1 also supports max.
2055
+ // `fugu` routes across all providers; `fugu-ultra` is
2046
2056
  // the heavier tier (may need larger client timeouts on complex tasks).
2047
2057
  {
2048
2058
  id: "fugu",
@@ -2066,7 +2076,8 @@ var MODELS = [
2066
2076
  supportsImages: true,
2067
2077
  supportsVideo: false,
2068
2078
  costTier: "high",
2069
- maxThinkingLevel: "xhigh"
2079
+ // The rolling alias now serves v1.1, which adds a distinct max effort.
2080
+ maxThinkingLevel: "max"
2070
2081
  },
2071
2082
  // ── xAI (Grok) ─────────────────────────────────────────
2072
2083
  // Grok 4.6 (released 2026-08-12) is xAI's flagship for coding, agentic tasks,
@@ -2120,6 +2131,34 @@ var MODELS = [
2120
2131
  costTier: "low",
2121
2132
  maxThinkingLevel: "high"
2122
2133
  },
2134
+ // Keep 3.1 Flash Lite first for the working OAuth default and fast-model routing.
2135
+ // New GA models are opt-in; Code Assist access varies by account.
2136
+ {
2137
+ id: "gemini-3.8-flash",
2138
+ name: "Gemini 3.8 Flash",
2139
+ provider: "gemini",
2140
+ contextWindow: 1048576,
2141
+ maxOutputTokens: 65536,
2142
+ supportsThinking: true,
2143
+ supportsImages: true,
2144
+ supportsVideo: true,
2145
+ maxVideoBytes: 20 * 1024 * 1024,
2146
+ costTier: "low",
2147
+ maxThinkingLevel: "high"
2148
+ },
2149
+ {
2150
+ id: "gemini-3.5-flash-lite",
2151
+ name: "Gemini 3.5 Flash Lite",
2152
+ provider: "gemini",
2153
+ contextWindow: 1048576,
2154
+ maxOutputTokens: 65536,
2155
+ supportsThinking: true,
2156
+ supportsImages: true,
2157
+ supportsVideo: true,
2158
+ maxVideoBytes: 20 * 1024 * 1024,
2159
+ costTier: "low",
2160
+ maxThinkingLevel: "high"
2161
+ },
2123
2162
  {
2124
2163
  // Gemini 3.7 Flash (released 2026-08-13) — Google's most capable Flash for
2125
2164
  // coding, agents, and multi-step execution; GA-stable on the Gemini API as
@@ -2127,7 +2166,7 @@ var MODELS = [
2127
2166
  // Sent over our Code Assist (OAuth) transport ahead of gemini-cli — upstream
2128
2167
  // hasn't listed 3.7 yet (google-gemini/gemini-cli#28802, still open) — so
2129
2168
  // free/personal accounts 404 (entitlement-gated) while Code Assist
2130
- // Standard/Enterprise accounts get it. Listed SECOND, after flash-lite:
2169
+ // Standard/Enterprise accounts get it. Kept after the working flash-lite:
2131
2170
  // getFastModel picks the first low-tier entry, and flash-lite is the one
2132
2171
  // that works on every account.
2133
2172
  id: "gemini-3.7-flash",
@@ -2319,21 +2358,19 @@ var MODELS = [
2319
2358
  {
2320
2359
  // `deepseek-v4-pro` now serves DeepSeek-V4-Pro-0813 (released 2026-08-13,
2321
2360
  // first STABLE V4 Pro — supersedes the April preview; calling name
2322
- // unchanged, same 1.6T/49B MoE). 1M context, 384K (393,216) max output,
2323
- // text-only, reasoning ladder low/high plus Think Max — mapped from our
2324
- // `xhigh`. ~$0.43/$0.87 per MTok on DeepSeek's own API, so a mid-tier
2325
- // price band rather than the preview's top band.
2361
+ // unchanged, same 1.6T/49B MoE). 1M context, text-only, low/high/max effort.
2362
+ // Docs abbreviate output as 384K; use the same conservative 384,000-token
2363
+ // application cap across V4 models rather than mixing decimal/binary units.
2326
2364
  id: "deepseek-v4-pro",
2327
2365
  name: "DeepSeek V4 Pro",
2328
2366
  provider: "deepseek",
2329
2367
  contextWindow: 1048576,
2330
- maxOutputTokens: 393216,
2368
+ maxOutputTokens: 384e3,
2331
2369
  supportsThinking: true,
2332
2370
  supportsImages: false,
2333
2371
  supportsVideo: false,
2334
2372
  costTier: "medium",
2335
- // DeepSeek V4 maps `xhigh` → its internal `max` tier.
2336
- maxThinkingLevel: "xhigh"
2373
+ maxThinkingLevel: "max"
2337
2374
  },
2338
2375
  {
2339
2376
  id: "deepseek-v4-flash",
@@ -2345,7 +2382,20 @@ var MODELS = [
2345
2382
  supportsImages: false,
2346
2383
  supportsVideo: false,
2347
2384
  costTier: "low",
2348
- maxThinkingLevel: "xhigh"
2385
+ maxThinkingLevel: "max"
2386
+ },
2387
+ // Opt-in experimental vision sibling; never replaces the stable summary model.
2388
+ {
2389
+ id: "deepseek-v4-flash-vision-exp",
2390
+ name: "DeepSeek V4 Flash Vision (Experimental)",
2391
+ provider: "deepseek",
2392
+ contextWindow: 1048576,
2393
+ maxOutputTokens: 384e3,
2394
+ supportsThinking: true,
2395
+ supportsImages: true,
2396
+ supportsVideo: false,
2397
+ costTier: "low",
2398
+ maxThinkingLevel: "max"
2349
2399
  },
2350
2400
  // ── OpenRouter ─────────────────────────────────────────
2351
2401
  {
@@ -2355,8 +2405,10 @@ var MODELS = [
2355
2405
  contextWindow: 1e6,
2356
2406
  maxOutputTokens: 65536,
2357
2407
  supportsThinking: true,
2358
- supportsImages: false,
2359
- supportsVideo: false,
2408
+ supportsImages: true,
2409
+ supportsVideo: true,
2410
+ // Practical inline-payload cap, not an asserted provider maximum.
2411
+ maxVideoBytes: 20 * 1024 * 1024,
2360
2412
  costTier: "medium",
2361
2413
  maxThinkingLevel: "high"
2362
2414
  },
@@ -2559,4 +2611,4 @@ export {
2559
2611
  getSummaryModel,
2560
2612
  getFastModel
2561
2613
  };
2562
- //# sourceMappingURL=chunk-OUE2GRO6.js.map
2614
+ //# sourceMappingURL=chunk-3MQNCB44.js.map