crosscheck-mcp 0.2.18 → 0.2.20
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/browser-ext.cjs +47 -14
- package/dist/browser-ext.cjs.map +1 -1
- package/dist/browser-ext.js +47 -14
- package/dist/browser-ext.js.map +1 -1
- package/dist/node-stdio.cjs +30 -12
- package/dist/node-stdio.cjs.map +1 -1
- package/dist/node-stdio.js +30 -12
- package/dist/node-stdio.js.map +1 -1
- package/dist/pricing.json +17 -2
- package/package.json +1 -1
package/dist/browser-ext.js
CHANGED
|
@@ -11,8 +11,8 @@ var pricing_default = {
|
|
|
11
11
|
_meta: {
|
|
12
12
|
currency: "USD",
|
|
13
13
|
unit: "per_1k_tokens",
|
|
14
|
-
updated_at: "2026-
|
|
15
|
-
notes: "Prices reflect public list pricing per provider. cached_per_1k applies to prompt tokens served from prompt-cache (Anthropic) or input-cache (OpenAI) when supported. Update when providers change rates. Missing models fall back to cost=0 with estimated=true. 2026-06-07: removed decommissioned gemini-2.0-flash / gemini-2.0-flash-lite (Google returns 404) and replaced the low-tier lite model with gemini-2.5-flash-lite. 2026-06-08: bumped the default Anthropic model to claude-opus-4-8. Kept gemini-2.5-pro as the Gemini default (gemini-3.1-pro-preview measured ~2.2x the cost at $2/$12 per M \u2014 output-token dominated by ~1.5k thinking tokens \u2014 so not worth it for a 1-of-4 panel voice). 2026-06-09: corrected rates verified against live provider pricing pages \u2014 gemini-2.5-pro output $5->$10/M; gemini-2.5-flash $0.15/$0.60 -> $0.30/$2.50; gpt-5 $5/$15 -> $1.25/$10 (cached $0.125); grok-4-latest input $5->$3/M (legacy grok-4 tier). Anthropic opus/sonnet/haiku rates confirmed unchanged. Older/unused models (gpt-4o, o1, o3, grok-3) not all re-verified. 2026-07-02: added claude-fable-5 as the optional reasoning-upgrade model (used only for the lead reasoning seat when the operator bumps a heavy task). 2026-07-04: set claude-fable-5 to its real pay-as-you-go API price effective 2026-07-07 \u2014 $10/M input ($0.01/1k), $50/M output ($0.05/1k), cached at 10% of input ($0.001/1k). Dropped the estimated flag. Note: Fable 5 is cheaper than opus-4-8 on both input and output. 2026-07-09: Google retired gemini-2.5-pro AND gemini-3-pro-preview (both now 404 'no longer available'). Moved the Gemini default to gemini-3.1-pro-preview ($2/M in, $12/M out, cache $0.20/M \u2014 the newest working pro; note it's a preview and may need another bump) and added gemini-3.5-flash ($1.50/M in, $9/M out, cache $0.15/M \u2014 GA) to the low/cheap tier. Removed the dead gemini-2.5-pro rate. IMPORTANT: gemini-3.x are heavy thinking models \u2014 thoughtsTokenCount (billed as output) is now folded into completion_tokens in the adapter, so cost reflects reasoning tokens. 2026-07-13: bumped the default OpenAI model to gpt-5.5 ($5/M in, $30/M out, cached $0.50/M \u2014 verified live). Note gpt-5.5 is markedly pricier than gpt-5 ($1.25/$10): ~4x input, 3x output. Kept gpt-5 in the table for cost lookups on pinned installs. 2026-07-16: added gpt-5.6 at the Sol (flagship) tier \u2014 $5/M in, $30/M out, cached $0.50/M (cache-write $6.25/M not modeled). gpt-5.6 is the second premium voice for the max-reasoning co-reasoner (runs alongside claude-fable-5 when an OpenAI key is set). Terra/Luna tiers not added (unused). 2026-07-17: added Moonshot AI (Kimi) provider with kimi-k3, base URL https://api.moonshot.ai/v1. Confirmed real rates \u2014 $3.00/M input (cache miss, $0.003/1k), $15.00/M output ($0.015/1k), $0.30/M cache-hit input ($0.0003/1k); 1M-token context. Kimi K3 is reasoning-class (temperature must be 1; uses max_completion_tokens). 2026-07-24: added claude-opus-5 ($5/M in, $25/M out, cached $0.50/M -- 10% of input, matching this table's existing cache-discount ratio) and made it the default anthropic model (was claude-opus-4-8). Fable 5 remains the reasoning-upgrade/super-mode target, unchanged. 2026-08-04: added Alibaba Cloud Qwen provider (OpenAI-compatible, DashScope international endpoint https://dashscope-intl.aliyuncs.com/compatible-mode/v1) with default model qwen3.8-max -- confirmed business pricing $2.00/M input ($0.002/1k), $6.00/M output ($0.006/1k), cache estimated at 10% of input ($0.0002/1k, not independently confirmed). Also added qwen3.7-max as a fallback tier at $1.25/M in ($0.00125/1k) / $3.75/M out ($0.00375/1k). Note: the model string initially configured (qwen-v3.8 / qwen-v3.7) does not exist on DashScope and 404s -- live-probed the real key against qwen3.8-max, qwen3.7-max, qwen3-max, qwen-max, and qwen-plus, all of which work; corrected the default and pricing keys to qwen3.8-max/qwen3.7-max and fixed QWEN_MODEL in both .env files to match. Qwen's Max tier is a standard instruct model, not reasoning-class -- normal temperature/max_tokens handling."
|
|
14
|
+
updated_at: "2026-09-06",
|
|
15
|
+
notes: "Prices reflect public list pricing per provider. cached_per_1k applies to prompt tokens served from prompt-cache (Anthropic) or input-cache (OpenAI) when supported. Update when providers change rates. Missing models fall back to cost=0 with estimated=true. 2026-06-07: removed decommissioned gemini-2.0-flash / gemini-2.0-flash-lite (Google returns 404) and replaced the low-tier lite model with gemini-2.5-flash-lite. 2026-06-08: bumped the default Anthropic model to claude-opus-4-8. Kept gemini-2.5-pro as the Gemini default (gemini-3.1-pro-preview measured ~2.2x the cost at $2/$12 per M \u2014 output-token dominated by ~1.5k thinking tokens \u2014 so not worth it for a 1-of-4 panel voice). 2026-06-09: corrected rates verified against live provider pricing pages \u2014 gemini-2.5-pro output $5->$10/M; gemini-2.5-flash $0.15/$0.60 -> $0.30/$2.50; gpt-5 $5/$15 -> $1.25/$10 (cached $0.125); grok-4-latest input $5->$3/M (legacy grok-4 tier). Anthropic opus/sonnet/haiku rates confirmed unchanged. Older/unused models (gpt-4o, o1, o3, grok-3) not all re-verified. 2026-07-02: added claude-fable-5 as the optional reasoning-upgrade model (used only for the lead reasoning seat when the operator bumps a heavy task). 2026-07-04: set claude-fable-5 to its real pay-as-you-go API price effective 2026-07-07 \u2014 $10/M input ($0.01/1k), $50/M output ($0.05/1k), cached at 10% of input ($0.001/1k). Dropped the estimated flag. Note: Fable 5 is cheaper than opus-4-8 on both input and output. 2026-07-09: Google retired gemini-2.5-pro AND gemini-3-pro-preview (both now 404 'no longer available'). Moved the Gemini default to gemini-3.1-pro-preview ($2/M in, $12/M out, cache $0.20/M \u2014 the newest working pro; note it's a preview and may need another bump) and added gemini-3.5-flash ($1.50/M in, $9/M out, cache $0.15/M \u2014 GA) to the low/cheap tier. Removed the dead gemini-2.5-pro rate. IMPORTANT: gemini-3.x are heavy thinking models \u2014 thoughtsTokenCount (billed as output) is now folded into completion_tokens in the adapter, so cost reflects reasoning tokens. 2026-07-13: bumped the default OpenAI model to gpt-5.5 ($5/M in, $30/M out, cached $0.50/M \u2014 verified live). Note gpt-5.5 is markedly pricier than gpt-5 ($1.25/$10): ~4x input, 3x output. Kept gpt-5 in the table for cost lookups on pinned installs. 2026-07-16: added gpt-5.6 at the Sol (flagship) tier \u2014 $5/M in, $30/M out, cached $0.50/M (cache-write $6.25/M not modeled). gpt-5.6 is the second premium voice for the max-reasoning co-reasoner (runs alongside claude-fable-5 when an OpenAI key is set). Terra/Luna tiers not added (unused). 2026-07-17: added Moonshot AI (Kimi) provider with kimi-k3, base URL https://api.moonshot.ai/v1. Confirmed real rates \u2014 $3.00/M input (cache miss, $0.003/1k), $15.00/M output ($0.015/1k), $0.30/M cache-hit input ($0.0003/1k); 1M-token context. Kimi K3 is reasoning-class (temperature must be 1; uses max_completion_tokens). 2026-07-24: added claude-opus-5 ($5/M in, $25/M out, cached $0.50/M -- 10% of input, matching this table's existing cache-discount ratio) and made it the default anthropic model (was claude-opus-4-8). Fable 5 remains the reasoning-upgrade/super-mode target, unchanged. 2026-08-04: added Alibaba Cloud Qwen provider (OpenAI-compatible, DashScope international endpoint https://dashscope-intl.aliyuncs.com/compatible-mode/v1) with default model qwen3.8-max -- confirmed business pricing $2.00/M input ($0.002/1k), $6.00/M output ($0.006/1k), cache estimated at 10% of input ($0.0002/1k, not independently confirmed). Also added qwen3.7-max as a fallback tier at $1.25/M in ($0.00125/1k) / $3.75/M out ($0.00375/1k). Note: the model string initially configured (qwen-v3.8 / qwen-v3.7) does not exist on DashScope and 404s -- live-probed the real key against qwen3.8-max, qwen3.7-max, qwen3-max, qwen-max, and qwen-plus, all of which work; corrected the default and pricing keys to qwen3.8-max/qwen3.7-max and fixed QWEN_MODEL in both .env files to match. Qwen's Max tier is a standard instruct model, not reasoning-class -- normal temperature/max_tokens handling. 2026-09-06: added gpt-6-astra at published list rates ($10/$50 per Mtok, $1/Mtok cache read; the $12.50/Mtok cache-WRITE rate has no field in this schema and is not modelled). gemini-3.8-flash now carries Google's published paid-tier rates ($0.75/$3.75/$0.075 per Mtok). THESE DOUBLE ON 2027-01-01 ($1.50/$7.50/$0.15) and this schema has no effective-date support, so they must be updated by hand on that date or every Gemini cost will read half of actual. claude-fable-5-1 now carries published rates ($10/$50/$0.25 per Mtok). Note its cache-read rate is 4x cheaper than claude-fable-5's, so the two must not be assumed to track each other. No inferred rows remain. All three were previously absent, which silently recorded their cost as $0."
|
|
16
16
|
},
|
|
17
17
|
anthropic: {
|
|
18
18
|
"claude-opus-4-5": {
|
|
@@ -64,6 +64,11 @@ var pricing_default = {
|
|
|
64
64
|
prompt_per_1k: 8e-4,
|
|
65
65
|
completion_per_1k: 4e-3,
|
|
66
66
|
cached_per_1k: 8e-5
|
|
67
|
+
},
|
|
68
|
+
"claude-fable-5-1": {
|
|
69
|
+
prompt_per_1k: 0.01,
|
|
70
|
+
completion_per_1k: 0.05,
|
|
71
|
+
cached_per_1k: 25e-5
|
|
67
72
|
}
|
|
68
73
|
},
|
|
69
74
|
openai: {
|
|
@@ -116,6 +121,11 @@ var pricing_default = {
|
|
|
116
121
|
prompt_per_1k: 11e-4,
|
|
117
122
|
completion_per_1k: 44e-4,
|
|
118
123
|
cached_per_1k: 55e-5
|
|
124
|
+
},
|
|
125
|
+
"gpt-6-astra": {
|
|
126
|
+
prompt_per_1k: 0.01,
|
|
127
|
+
completion_per_1k: 0.05,
|
|
128
|
+
cached_per_1k: 1e-3
|
|
119
129
|
}
|
|
120
130
|
},
|
|
121
131
|
xai: {
|
|
@@ -155,6 +165,11 @@ var pricing_default = {
|
|
|
155
165
|
prompt_per_1k: 1e-4,
|
|
156
166
|
completion_per_1k: 4e-4,
|
|
157
167
|
cached_per_1k: 25e-6
|
|
168
|
+
},
|
|
169
|
+
"gemini-3.8-flash": {
|
|
170
|
+
prompt_per_1k: 75e-5,
|
|
171
|
+
completion_per_1k: 375e-5,
|
|
172
|
+
cached_per_1k: 75e-6
|
|
158
173
|
}
|
|
159
174
|
},
|
|
160
175
|
mistral: {
|
|
@@ -390,7 +405,13 @@ var PROVIDER_CAPS = {
|
|
|
390
405
|
family: "openai_chat",
|
|
391
406
|
system_role: "inline",
|
|
392
407
|
supports_temperature: "model",
|
|
393
|
-
|
|
408
|
+
// gpt-6 added 2026-09-06: GPT-6 Astra rejects `max_tokens` outright
|
|
409
|
+
// ("Unsupported parameter ... use max_completion_tokens"), which is a 400,
|
|
410
|
+
// NOT a model-access failure — so the fallback chain correctly declines to
|
|
411
|
+
// fire and the provider silently drops out of the panel instead. Reasoning
|
|
412
|
+
// -class routing is what selects max_completion_tokens and omits
|
|
413
|
+
// temperature, so the prefix is the fix.
|
|
414
|
+
reasoning_prefixes: ["gpt-5", "gpt-6", "o1", "o3", "o4"]
|
|
394
415
|
},
|
|
395
416
|
xai: { family: "openai_chat", system_role: "inline", supports_temperature: true },
|
|
396
417
|
mistral: { family: "openai_chat", system_role: "inline", supports_temperature: true },
|
|
@@ -1200,18 +1221,24 @@ async function sendOpenAICompatible(args) {
|
|
|
1200
1221
|
// src/providers/registry.ts
|
|
1201
1222
|
var DEFAULT_MODELS = {
|
|
1202
1223
|
anthropic: "claude-opus-5",
|
|
1203
|
-
|
|
1204
|
-
// capped at gpt-5) — gpt-5 is the safer, more broadly available default
|
|
1205
|
-
// until a user/org explicitly opts into 5.5 via `crosscheck models set`.
|
|
1206
|
-
openai: "gpt-5",
|
|
1224
|
+
openai: "gpt-5.6-sol",
|
|
1207
1225
|
xai: "grok-4-latest",
|
|
1208
1226
|
mistral: "mistral-large-latest",
|
|
1209
1227
|
groq: "llama-3.3-70b-versatile",
|
|
1210
1228
|
deepseek: "deepseek-chat",
|
|
1211
|
-
gemini: "gemini-3.
|
|
1229
|
+
gemini: "gemini-3.8-flash",
|
|
1212
1230
|
kimi: "kimi-k3",
|
|
1213
1231
|
qwen: "qwen3.8-max"
|
|
1214
1232
|
};
|
|
1233
|
+
var DEFAULT_FALLBACKS = {
|
|
1234
|
+
// gpt-5.6 requires a billing-enabled account and is excluded from the free
|
|
1235
|
+
// tier, so gpt-5 remains the broadly-available floor. One hop, not a
|
|
1236
|
+
// ladder — an intermediate 5.5 step would just be another model to reason
|
|
1237
|
+
// about.
|
|
1238
|
+
openai: ["gpt-5"],
|
|
1239
|
+
// 3.8 Flash is recent enough that some projects won't have it yet.
|
|
1240
|
+
gemini: ["gemini-3.1-pro-preview"]
|
|
1241
|
+
};
|
|
1215
1242
|
var MODEL_ENV_VARS = {
|
|
1216
1243
|
anthropic: "ANTHROPIC_MODEL",
|
|
1217
1244
|
openai: "OPENAI_MODEL",
|
|
@@ -1229,7 +1256,11 @@ function parseFallbackModels(raw) {
|
|
|
1229
1256
|
}
|
|
1230
1257
|
function fallbackModelsFor(env, provider) {
|
|
1231
1258
|
const envVar = MODEL_ENV_VARS[provider];
|
|
1232
|
-
|
|
1259
|
+
const fromEnv = envVar ? parseFallbackModels(env[`${envVar}_FALLBACKS`]) : [];
|
|
1260
|
+
if (fromEnv.length > 0) return fromEnv;
|
|
1261
|
+
const pinned = envVar ? env[envVar] : void 0;
|
|
1262
|
+
if (pinned && pinned !== DEFAULT_MODELS[provider]) return [];
|
|
1263
|
+
return DEFAULT_FALLBACKS[provider] ?? [];
|
|
1233
1264
|
}
|
|
1234
1265
|
function buildProviders(opts) {
|
|
1235
1266
|
const out = {};
|
|
@@ -1320,7 +1351,7 @@ import { z } from "zod";
|
|
|
1320
1351
|
|
|
1321
1352
|
// src/server-meta.ts
|
|
1322
1353
|
var SERVER_NAME = "crosscheck-agent";
|
|
1323
|
-
var SERVER_VERSION = true ? "0.2.
|
|
1354
|
+
var SERVER_VERSION = true ? "0.2.20" : "0.0.0-dev";
|
|
1324
1355
|
|
|
1325
1356
|
// src/tools/audit.ts
|
|
1326
1357
|
import { readdirSync, readFileSync as readFileSync3, statSync } from "fs";
|
|
@@ -2890,7 +2921,7 @@ function numberOrNull(v) {
|
|
|
2890
2921
|
import { performance as performance2 } from "perf_hooks";
|
|
2891
2922
|
|
|
2892
2923
|
// src/core/co-reason.ts
|
|
2893
|
-
var CO_REASON_MODEL = "gpt-5.6";
|
|
2924
|
+
var CO_REASON_MODEL = "gpt-5.6-sol";
|
|
2894
2925
|
var CO_REASON_PROVIDER = "openai";
|
|
2895
2926
|
var CO_REASON_LABEL = "GPT-5.6";
|
|
2896
2927
|
function pickCoReasoner(providers, upgradeRequested) {
|
|
@@ -4248,8 +4279,10 @@ var SUPER_MODELS = {
|
|
|
4248
4279
|
// on it succeeded as recently as 2026-08-30 — so it is an unlisted alias
|
|
4249
4280
|
// rather than a dead id.
|
|
4250
4281
|
xai: [DEFAULT_MODELS["xai"]],
|
|
4251
|
-
// 3.8 Flash
|
|
4252
|
-
|
|
4282
|
+
// 3.8 Flash is now BOTH the everyday default and the super pick for Gemini,
|
|
4283
|
+
// so this can no longer borrow DEFAULT_MODELS as its fallback — that would
|
|
4284
|
+
// repeat the same model and retry one that just failed. Named explicitly.
|
|
4285
|
+
gemini: ["gemini-3.8-flash", "gemini-3.1-pro-preview"],
|
|
4253
4286
|
kimi: [DEFAULT_MODELS["kimi"]],
|
|
4254
4287
|
qwen: [DEFAULT_MODELS["qwen"]]
|
|
4255
4288
|
};
|
|
@@ -11737,7 +11770,7 @@ var CHECK_INTERVAL_SECONDS = 3 * 24 * 60 * 60;
|
|
|
11737
11770
|
var DEFAULT_PACKAGE = "crosscheck-cli";
|
|
11738
11771
|
var FETCH_TIMEOUT_MS = 3e3;
|
|
11739
11772
|
function engineVersion() {
|
|
11740
|
-
return true ? "0.2.
|
|
11773
|
+
return true ? "0.2.20" : "0.0.0-dev";
|
|
11741
11774
|
}
|
|
11742
11775
|
function defaultUpdateCachePath() {
|
|
11743
11776
|
const base = process.env["CROSSCHECK_DATA_DIR"] || path9.join(os.homedir() || os.tmpdir(), ".crosscheck");
|