@prestyj/core 5.23.0 → 5.25.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
@@ -1935,11 +1935,37 @@ var MODELS = [
1935
1935
  // costTier: "high",
1936
1936
  // maxThinkingLevel: "max",
1937
1937
  // },
1938
+ {
1939
+ // Released 2026-09-22 — "For long-running agentic coding and knowledge
1940
+ // work". Fable-class capability at $4/$20 MTok (cheaper than the Opus 5 it
1941
+ // replaces, $5/$25). Adaptive thinking with the full effort ladder
1942
+ // (low→max, xhigh included), but thinking can no longer be disabled: a
1943
+ // `thinking: {type: "disabled"}` or budget_tokens request 400s. @prestyj/ai
1944
+ // omits the field entirely when thinking is off, so that path is safe.
1945
+ // Forced tool use (`tool_choice` any/tool) also 400s — see
1946
+ // `toAnthropicToolChoice`, which downgrades it to `auto` for this model.
1947
+ // Anthropic declares the server-side default effort as `medium` (Opus 5
1948
+ // ran `high`), and 5.5 thinks more per turn at a given level, so a fresh
1949
+ // session starts at `medium` rather than the ladder ceiling.
1950
+ id: "claude-opus-5-5",
1951
+ name: "Claude Opus 5.5",
1952
+ provider: "anthropic",
1953
+ contextWindow: 1e6,
1954
+ maxOutputTokens: 128e3,
1955
+ supportsThinking: true,
1956
+ defaultThinkingLevel: "medium",
1957
+ supportsImages: true,
1958
+ supportsVideo: false,
1959
+ costTier: "high",
1960
+ maxThinkingLevel: "max"
1961
+ },
1938
1962
  {
1939
1963
  // Released 2026-07-24 — "For complex agentic coding and enterprise work".
1940
1964
  // Near-Fable capability at half the price ($5/$25 vs $10/$50). Adaptive
1941
1965
  // thinking with the full effort ladder (low→max, xhigh included); dateless
1942
- // ID is the canonical pinned snapshot (post-4.6 naming scheme).
1966
+ // ID is the canonical pinned snapshot (post-4.6 naming scheme). Kept as a
1967
+ // legacy option now that Opus 5.5 leads the line: it's the last Opus that
1968
+ // accepts disabled thinking and forced tool use.
1943
1969
  id: "claude-opus-5",
1944
1970
  name: "Claude Opus 5",
1945
1971
  provider: "anthropic",
@@ -2088,16 +2114,21 @@ var MODELS = [
2088
2114
  maxThinkingLevel: "max"
2089
2115
  },
2090
2116
  // ── xAI (Grok) ─────────────────────────────────────────
2091
- // Grok 4.6 (released 2026-08-12) is xAI's flagship for coding, agentic tasks,
2092
- // and knowledge work, with a focus on long-running agents — 500K context,
2093
- // text+image input, and a `reasoning_effort` ladder that adds a new `xhigh`
2094
- // top rung (low/medium/high default/xhigh; reasoning still can't be fully
2117
+ // Grok 4.7 (released 2026-09-21) — xAI's flagship for coding, agentic tasks,
2118
+ // and knowledge work: a new, larger base model with a longer RL run weighted
2119
+ // toward hours-long tasks, plus stronger self-verification and long-context
2120
+ // management. 500K context, text+image input, and a `reasoning_effort`
2121
+ // ladder of low/medium/high default/xhigh (reasoning still can't be fully
2095
2122
  // disabled). $2/$6 per MTok under 200K prompt tokens ($4/$12 at or above),
2096
- // and it's the default model of the Grok Build coding agent. xAI advertises "no text output limit"; we keep the same
2097
- // 131K practical cap as 4.5 for budget predictability and input headroom.
2123
+ // and it's the default model of the Grok Build coding agent. xAI advertises
2124
+ // "no fixed text output limit"; we keep the 131K practical cap for budget
2125
+ // predictability and input headroom. (A faster "Grok 4.7 Fast" variant
2126
+ // exists but is Cursor/Grok Build-only — not on the public API — so it isn't
2127
+ // registered.) Only the newest Grok ships — 4.6/4.5 are superseded and
2128
+ // retired; saved sessions on them fall back to this default.
2098
2129
  {
2099
- id: "grok-4.6",
2100
- name: "Grok 4.6",
2130
+ id: "grok-4.7",
2131
+ name: "Grok 4.7",
2101
2132
  provider: "xai",
2102
2133
  contextWindow: 5e5,
2103
2134
  maxOutputTokens: 131072,
@@ -2107,24 +2138,6 @@ var MODELS = [
2107
2138
  costTier: "medium",
2108
2139
  maxThinkingLevel: "xhigh"
2109
2140
  },
2110
- // Grok 4.5 (released 2026-07-08) — superseded by 4.6 but retained as an explicit option. 500K context, text+image input,
2111
- // configurable `reasoning_effort` (low/medium/high, server default high;
2112
- // reasoning can't be fully disabled). Served over the OpenAI-compatible API
2113
- // at https://api.x.ai/v1 (API key from console.x.ai). xAI hasn't published an
2114
- // official max-output cap for 4.5; 131K matches the Grok Responses ceiling
2115
- // third-party integrations use.
2116
- {
2117
- id: "grok-4.5",
2118
- name: "Grok 4.5",
2119
- provider: "xai",
2120
- contextWindow: 5e5,
2121
- maxOutputTokens: 131072,
2122
- supportsThinking: true,
2123
- supportsImages: true,
2124
- supportsVideo: false,
2125
- costTier: "medium",
2126
- maxThinkingLevel: "high"
2127
- },
2128
2141
  // ── Gemini ─────────────────────────────────────────
2129
2142
  {
2130
2143
  id: "gemini-3.1-flash-lite",
@@ -2313,11 +2326,15 @@ var MODELS = [
2313
2326
  maxThinkingLevel: "high"
2314
2327
  },
2315
2328
  // ── Xiaomi (MiMo) ──────────────────────────────────────
2316
- // V2.6 series (2026-09) supersedes V2.5 one-for-one: pro → pro, the omni
2317
- // `mimo-v2.5` → flash, ultraspeed → ultraspeed. The whole series is now
2318
- // full-modality, so unlike V2.5-Pro the flagship no longer needs a separate
2319
- // omni sibling for attachments. V2.5 entries are retired here — a session
2320
- // that still has one saved falls back to the provider default on next start.
2329
+ // V2.6 series (released 2026-09-22, open-weight: Pro 1.02T/42B-A, Flash
2330
+ // 309B/15B-A, plus a 9B Qwen distill not served over the API) supersedes V2.5
2331
+ // one-for-one: pro → pro, the omni `mimo-v2.5` → flash, ultraspeed →
2332
+ // ultraspeed. Every V2.6 text model is natively full-modality, so unlike
2333
+ // V2.5-Pro the flagship no longer needs a separate omni sibling for
2334
+ // attachments — image/video ride the same OpenAI-compatible base64 transport
2335
+ // the old omni model used. API prices are unchanged from V2.5. The V2.5 ids
2336
+ // deprecate on the platform 2026-10-21 and are retired here — a session that
2337
+ // still has one saved falls back to the provider default on next start.
2321
2338
  //
2322
2339
  // Capabilities below are measured against the Token Plan host, not taken
2323
2340
  // from marketing copy: image and video both come back with `image_tokens` /
@@ -2327,6 +2344,8 @@ var MODELS = [
2327
2344
  // binary 1M below (2^20) and not the decimal 1e6 V2.5 was listed with — the
2328
2345
  // few-token gap is the chat envelope the server adds on top of the content.
2329
2346
  {
2347
+ // Coding/agentic flagship — highest open-weight score on Artificial
2348
+ // Analysis at launch (46, tied with Grok 4.7).
2330
2349
  id: "mimo-v2.6-pro",
2331
2350
  name: "MiMo-V2.6-Pro",
2332
2351
  provider: "xiaomi",
@@ -2340,9 +2359,10 @@ var MODELS = [
2340
2359
  maxThinkingLevel: "high",
2341
2360
  authStorageKeys: ["xiaomi", XIAOMI_CREDITS_KEY]
2342
2361
  },
2343
- // Flash: the cheap, high-frequency sibling at the same modality surface and
2344
- // window as Pro. It is the provider's `low` tier, so scout sub-agents and
2345
- // compaction summaries route here instead of paying Pro rates.
2362
+ // Flash: the cheap, high-frequency sibling (~10% of Pro's price class) at the
2363
+ // same modality surface and window as Pro. It is the provider's `low` tier,
2364
+ // so scout sub-agents and compaction summaries route here instead of paying
2365
+ // Pro rates.
2346
2366
  {
2347
2367
  id: "mimo-v2.6-flash",
2348
2368
  name: "MiMo-V2.6-Flash",
@@ -2363,9 +2383,9 @@ var MODELS = [
2363
2383
  // authStorageKeys doc). The Token Plan host rejects it with "Not supported
2364
2384
  // model" — the known-model/wrong-host reply — where an invented id gets
2365
2385
  // "Unsupported model", which is how this id was confirmed without a
2366
- // Credits key. Attachment support is inferred from the series announcement
2367
- // ("full modality across the series") rather than measured: images are
2368
- // enabled, video stays off until it can be verified on the platform host.
2386
+ // Credits key. Attachments can't be probed directly for the same reason, so
2387
+ // this entry tracks the rest of the V2.6 series ("full modality across the
2388
+ // series") and mirrors Pro's verified image+video surface.
2369
2389
  {
2370
2390
  id: "mimo-v2.6-pro-ultraspeed",
2371
2391
  name: "MiMo-V2.6-Pro-UltraSpeed",
@@ -2374,7 +2394,8 @@ var MODELS = [
2374
2394
  maxOutputTokens: 131072,
2375
2395
  supportsThinking: true,
2376
2396
  supportsImages: true,
2377
- supportsVideo: false,
2397
+ supportsVideo: true,
2398
+ maxVideoBytes: 36 * 1024 * 1024,
2378
2399
  costTier: "high",
2379
2400
  maxThinkingLevel: "high",
2380
2401
  authStorageKeys: [XIAOMI_CREDITS_KEY]
@@ -2524,7 +2545,7 @@ function getDefaultModel(provider) {
2524
2545
  return MODELS.find((m) => m.id === "Qwen/Qwen3-Coder-480B-A35B-Instruct");
2525
2546
  if (provider === "openrouter") return MODELS.find((m) => m.id === "qwen/qwen3.6-plus");
2526
2547
  if (provider === "sakana") return MODELS.find((m) => m.id === "fugu");
2527
- if (provider === "xai") return MODELS.find((m) => m.id === "grok-4.6");
2548
+ if (provider === "xai") return MODELS.find((m) => m.id === "grok-4.7");
2528
2549
  if (provider === "local") {
2529
2550
  return getModelsForProvider("local")[0] ?? PLACEHOLDER_LOCAL_MODEL;
2530
2551
  }
@@ -2636,4 +2657,4 @@ export {
2636
2657
  getSummaryModel,
2637
2658
  getFastModel
2638
2659
  };
2639
- //# sourceMappingURL=chunk-SQLBQOQH.js.map
2660
+ //# sourceMappingURL=chunk-GDM5Q7SA.js.map