crosscheck-mcp 0.2.10 → 0.2.11
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/browser-ext.cjs +38 -11
- package/dist/browser-ext.cjs.map +1 -1
- package/dist/browser-ext.js +38 -11
- package/dist/browser-ext.js.map +1 -1
- package/dist/node-stdio.cjs +28 -9
- package/dist/node-stdio.cjs.map +1 -1
- package/dist/node-stdio.js +28 -9
- package/dist/node-stdio.js.map +1 -1
- package/dist/pricing.json +14 -2
- package/package.json +2 -2
package/dist/browser-ext.js
CHANGED
|
@@ -11,8 +11,8 @@ var pricing_default = {
|
|
|
11
11
|
_meta: {
|
|
12
12
|
currency: "USD",
|
|
13
13
|
unit: "per_1k_tokens",
|
|
14
|
-
updated_at: "2026-
|
|
15
|
-
notes: "Prices reflect public list pricing per provider. cached_per_1k applies to prompt tokens served from prompt-cache (Anthropic) or input-cache (OpenAI) when supported. Update when providers change rates. Missing models fall back to cost=0 with estimated=true. 2026-06-07: removed decommissioned gemini-2.0-flash / gemini-2.0-flash-lite (Google returns 404) and replaced the low-tier lite model with gemini-2.5-flash-lite. 2026-06-08: bumped the default Anthropic model to claude-opus-4-8. Kept gemini-2.5-pro as the Gemini default (gemini-3.1-pro-preview measured ~2.2x the cost at $2/$12 per M \u2014 output-token dominated by ~1.5k thinking tokens \u2014 so not worth it for a 1-of-4 panel voice). 2026-06-09: corrected rates verified against live provider pricing pages \u2014 gemini-2.5-pro output $5->$10/M; gemini-2.5-flash $0.15/$0.60 -> $0.30/$2.50; gpt-5 $5/$15 -> $1.25/$10 (cached $0.125); grok-4-latest input $5->$3/M (legacy grok-4 tier). Anthropic opus/sonnet/haiku rates confirmed unchanged. Older/unused models (gpt-4o, o1, o3, grok-3) not all re-verified. 2026-07-02: added claude-fable-5 as the optional reasoning-upgrade model (used only for the lead reasoning seat when the operator bumps a heavy task). 2026-07-04: set claude-fable-5 to its real pay-as-you-go API price effective 2026-07-07 \u2014 $10/M input ($0.01/1k), $50/M output ($0.05/1k), cached at 10% of input ($0.001/1k). Dropped the estimated flag. Note: Fable 5 is cheaper than opus-4-8 on both input and output. 2026-07-09: Google retired gemini-2.5-pro AND gemini-3-pro-preview (both now 404 'no longer available'). Moved the Gemini default to gemini-3.1-pro-preview ($2/M in, $12/M out, cache $0.20/M \u2014 the newest working pro; note it's a preview and may need another bump) and added gemini-3.5-flash ($1.50/M in, $9/M out, cache $0.15/M \u2014 GA) to the low/cheap tier. Removed the dead gemini-2.5-pro rate. IMPORTANT: gemini-3.x are heavy thinking models \u2014 thoughtsTokenCount (billed as output) is now folded into completion_tokens in the adapter, so cost reflects reasoning tokens. 2026-07-13: bumped the default OpenAI model to gpt-5.5 ($5/M in, $30/M out, cached $0.50/M \u2014 verified live). Note gpt-5.5 is markedly pricier than gpt-5 ($1.25/$10): ~4x input, 3x output. Kept gpt-5 in the table for cost lookups on pinned installs. 2026-07-16: added gpt-5.6 at the Sol (flagship) tier \u2014 $5/M in, $30/M out, cached $0.50/M (cache-write $6.25/M not modeled). gpt-5.6 is the second premium voice for the max-reasoning co-reasoner (runs alongside claude-fable-5 when an OpenAI key is set). Terra/Luna tiers not added (unused). 2026-07-17: added Moonshot AI (Kimi) provider with kimi-k3, base URL https://api.moonshot.ai/v1. Confirmed real rates \u2014 $3.00/M input (cache miss, $0.003/1k), $15.00/M output ($0.015/1k), $0.30/M cache-hit input ($0.0003/1k); 1M-token context. Kimi K3 is reasoning-class (temperature must be 1; uses max_completion_tokens). 2026-07-24: added claude-opus-5 ($5/M in, $25/M out, cached $0.50/M -- 10% of input, matching this table's existing cache-discount ratio) and made it the default anthropic model (was claude-opus-4-8). Fable 5 remains the reasoning-upgrade/super-mode target, unchanged."
|
|
14
|
+
updated_at: "2026-08-04",
|
|
15
|
+
notes: "Prices reflect public list pricing per provider. cached_per_1k applies to prompt tokens served from prompt-cache (Anthropic) or input-cache (OpenAI) when supported. Update when providers change rates. Missing models fall back to cost=0 with estimated=true. 2026-06-07: removed decommissioned gemini-2.0-flash / gemini-2.0-flash-lite (Google returns 404) and replaced the low-tier lite model with gemini-2.5-flash-lite. 2026-06-08: bumped the default Anthropic model to claude-opus-4-8. Kept gemini-2.5-pro as the Gemini default (gemini-3.1-pro-preview measured ~2.2x the cost at $2/$12 per M \u2014 output-token dominated by ~1.5k thinking tokens \u2014 so not worth it for a 1-of-4 panel voice). 2026-06-09: corrected rates verified against live provider pricing pages \u2014 gemini-2.5-pro output $5->$10/M; gemini-2.5-flash $0.15/$0.60 -> $0.30/$2.50; gpt-5 $5/$15 -> $1.25/$10 (cached $0.125); grok-4-latest input $5->$3/M (legacy grok-4 tier). Anthropic opus/sonnet/haiku rates confirmed unchanged. Older/unused models (gpt-4o, o1, o3, grok-3) not all re-verified. 2026-07-02: added claude-fable-5 as the optional reasoning-upgrade model (used only for the lead reasoning seat when the operator bumps a heavy task). 2026-07-04: set claude-fable-5 to its real pay-as-you-go API price effective 2026-07-07 \u2014 $10/M input ($0.01/1k), $50/M output ($0.05/1k), cached at 10% of input ($0.001/1k). Dropped the estimated flag. Note: Fable 5 is cheaper than opus-4-8 on both input and output. 2026-07-09: Google retired gemini-2.5-pro AND gemini-3-pro-preview (both now 404 'no longer available'). Moved the Gemini default to gemini-3.1-pro-preview ($2/M in, $12/M out, cache $0.20/M \u2014 the newest working pro; note it's a preview and may need another bump) and added gemini-3.5-flash ($1.50/M in, $9/M out, cache $0.15/M \u2014 GA) to the low/cheap tier. Removed the dead gemini-2.5-pro rate. IMPORTANT: gemini-3.x are heavy thinking models \u2014 thoughtsTokenCount (billed as output) is now folded into completion_tokens in the adapter, so cost reflects reasoning tokens. 2026-07-13: bumped the default OpenAI model to gpt-5.5 ($5/M in, $30/M out, cached $0.50/M \u2014 verified live). Note gpt-5.5 is markedly pricier than gpt-5 ($1.25/$10): ~4x input, 3x output. Kept gpt-5 in the table for cost lookups on pinned installs. 2026-07-16: added gpt-5.6 at the Sol (flagship) tier \u2014 $5/M in, $30/M out, cached $0.50/M (cache-write $6.25/M not modeled). gpt-5.6 is the second premium voice for the max-reasoning co-reasoner (runs alongside claude-fable-5 when an OpenAI key is set). Terra/Luna tiers not added (unused). 2026-07-17: added Moonshot AI (Kimi) provider with kimi-k3, base URL https://api.moonshot.ai/v1. Confirmed real rates \u2014 $3.00/M input (cache miss, $0.003/1k), $15.00/M output ($0.015/1k), $0.30/M cache-hit input ($0.0003/1k); 1M-token context. Kimi K3 is reasoning-class (temperature must be 1; uses max_completion_tokens). 2026-07-24: added claude-opus-5 ($5/M in, $25/M out, cached $0.50/M -- 10% of input, matching this table's existing cache-discount ratio) and made it the default anthropic model (was claude-opus-4-8). Fable 5 remains the reasoning-upgrade/super-mode target, unchanged. 2026-08-04: added Alibaba Cloud Qwen provider (OpenAI-compatible, DashScope international endpoint https://dashscope-intl.aliyuncs.com/compatible-mode/v1) with default model qwen3.8-max -- confirmed business pricing $2.00/M input ($0.002/1k), $6.00/M output ($0.006/1k), cache estimated at 10% of input ($0.0002/1k, not independently confirmed). Also added qwen3.7-max as a fallback tier at $1.25/M in ($0.00125/1k) / $3.75/M out ($0.00375/1k). Note: the model string initially configured (qwen-v3.8 / qwen-v3.7) does not exist on DashScope and 404s -- live-probed the real key against qwen3.8-max, qwen3.7-max, qwen3-max, qwen-max, and qwen-plus, all of which work; corrected the default and pricing keys to qwen3.8-max/qwen3.7-max and fixed QWEN_MODEL in both .env files to match. Qwen's Max tier is a standard instruct model, not reasoning-class -- normal temperature/max_tokens handling."
|
|
16
16
|
},
|
|
17
17
|
anthropic: {
|
|
18
18
|
"claude-opus-4-5": {
|
|
@@ -200,6 +200,18 @@ var pricing_default = {
|
|
|
200
200
|
cached_per_1k: 3e-4
|
|
201
201
|
}
|
|
202
202
|
},
|
|
203
|
+
qwen: {
|
|
204
|
+
"qwen3.8-max": {
|
|
205
|
+
prompt_per_1k: 2e-3,
|
|
206
|
+
completion_per_1k: 6e-3,
|
|
207
|
+
cached_per_1k: 2e-4
|
|
208
|
+
},
|
|
209
|
+
"qwen3.7-max": {
|
|
210
|
+
prompt_per_1k: 125e-5,
|
|
211
|
+
completion_per_1k: 375e-5,
|
|
212
|
+
cached_per_1k: 125e-6
|
|
213
|
+
}
|
|
214
|
+
},
|
|
203
215
|
_tiers: {
|
|
204
216
|
low: {
|
|
205
217
|
description: "Cheap, fast models for simple subtasks (extraction, formatting, short summaries).",
|
|
@@ -397,7 +409,11 @@ var PROVIDER_CAPS = {
|
|
|
397
409
|
system_role: "separate",
|
|
398
410
|
supports_temperature: true,
|
|
399
411
|
reasoning_prefixes: ["gemini-3.1-pro", "gemini-3.5-flash", "gemini-2.5-pro"]
|
|
400
|
-
}
|
|
412
|
+
},
|
|
413
|
+
// Qwen's "Max" tier (qwen-v3.8, qwen-v3.7) is a standard instruct/chat
|
|
414
|
+
// model, not a visible-reasoning model like QwQ — normal temperature and
|
|
415
|
+
// max_tokens handling, same as xai/mistral/groq/deepseek.
|
|
416
|
+
qwen: { family: "openai_chat", system_role: "inline", supports_temperature: true }
|
|
401
417
|
};
|
|
402
418
|
function isReasoningModel(provider, model) {
|
|
403
419
|
const caps = PROVIDER_CAPS[provider];
|
|
@@ -968,7 +984,12 @@ var OPENAI_COMPAT_DEFAULT_URLS = {
|
|
|
968
984
|
groq: "https://api.groq.com/openai/v1/chat/completions",
|
|
969
985
|
deepseek: "https://api.deepseek.com/v1/chat/completions",
|
|
970
986
|
// Moonshot AI (Kimi). OpenAI-compatible chat completions endpoint.
|
|
971
|
-
kimi: "https://api.moonshot.ai/v1/chat/completions"
|
|
987
|
+
kimi: "https://api.moonshot.ai/v1/chat/completions",
|
|
988
|
+
// Alibaba Cloud DashScope (Qwen). International endpoint — the mainland
|
|
989
|
+
// China endpoint (dashscope.aliyuncs.com) 403s/404s from outside China in
|
|
990
|
+
// a way that can look like a bad key; this is the deliberate default for
|
|
991
|
+
// a US-based deployment.
|
|
992
|
+
qwen: "https://dashscope-intl.aliyuncs.com/compatible-mode/v1/chat/completions"
|
|
972
993
|
};
|
|
973
994
|
var OPENAI_COMPAT_NATIVE_STRUCTURED = /* @__PURE__ */ new Set([
|
|
974
995
|
"openai"
|
|
@@ -1137,7 +1158,8 @@ var DEFAULT_MODELS = {
|
|
|
1137
1158
|
groq: "llama-3.3-70b-versatile",
|
|
1138
1159
|
deepseek: "deepseek-chat",
|
|
1139
1160
|
gemini: "gemini-3.1-pro-preview",
|
|
1140
|
-
kimi: "kimi-k3"
|
|
1161
|
+
kimi: "kimi-k3",
|
|
1162
|
+
qwen: "qwen3.8-max"
|
|
1141
1163
|
};
|
|
1142
1164
|
function buildProviders(opts) {
|
|
1143
1165
|
const out = {};
|
|
@@ -1152,7 +1174,8 @@ function buildProviders(opts) {
|
|
|
1152
1174
|
{ name: "mistral", keyEnv: "MISTRAL_API_KEY", modelEnv: "MISTRAL_MODEL", defaultModel: DEFAULT_MODELS["mistral"] },
|
|
1153
1175
|
{ name: "groq", keyEnv: "GROQ_API_KEY", modelEnv: "GROQ_MODEL", defaultModel: DEFAULT_MODELS["groq"] },
|
|
1154
1176
|
{ name: "deepseek", keyEnv: "DEEPSEEK_API_KEY", modelEnv: "DEEPSEEK_MODEL", defaultModel: DEFAULT_MODELS["deepseek"] },
|
|
1155
|
-
{ name: "kimi", keyEnv: "KIMI_API_KEY", modelEnv: "KIMI_MODEL", defaultModel: DEFAULT_MODELS["kimi"] }
|
|
1177
|
+
{ name: "kimi", keyEnv: "KIMI_API_KEY", modelEnv: "KIMI_MODEL", defaultModel: DEFAULT_MODELS["kimi"] },
|
|
1178
|
+
{ name: "qwen", keyEnv: "QWEN_API_KEY", modelEnv: "QWEN_MODEL", defaultModel: DEFAULT_MODELS["qwen"] }
|
|
1156
1179
|
];
|
|
1157
1180
|
for (const s of openAiCompatSpec) {
|
|
1158
1181
|
const apiKey = opts.env[s.keyEnv];
|
|
@@ -1224,7 +1247,7 @@ import { z } from "zod";
|
|
|
1224
1247
|
|
|
1225
1248
|
// src/server-meta.ts
|
|
1226
1249
|
var SERVER_NAME = "crosscheck-agent";
|
|
1227
|
-
var SERVER_VERSION = true ? "0.2.
|
|
1250
|
+
var SERVER_VERSION = true ? "0.2.11" : "0.0.0-dev";
|
|
1228
1251
|
|
|
1229
1252
|
// src/tools/audit.ts
|
|
1230
1253
|
import { readdirSync, readFileSync as readFileSync3, statSync } from "fs";
|
|
@@ -4073,7 +4096,9 @@ var SUPER_MODELS = {
|
|
|
4073
4096
|
// already the default — no-op retarget
|
|
4074
4097
|
gemini: DEFAULT_MODELS["gemini"],
|
|
4075
4098
|
// already the default — no-op retarget
|
|
4076
|
-
kimi: DEFAULT_MODELS["kimi"]
|
|
4099
|
+
kimi: DEFAULT_MODELS["kimi"],
|
|
4100
|
+
// already the default — no-op retarget
|
|
4101
|
+
qwen: DEFAULT_MODELS["qwen"]
|
|
4077
4102
|
// already the default — no-op retarget
|
|
4078
4103
|
};
|
|
4079
4104
|
var SUPER_PROVIDER_NAMES = Object.keys(SUPER_MODELS);
|
|
@@ -9876,7 +9901,8 @@ var KNOWN_PROVIDERS6 = [
|
|
|
9876
9901
|
"mistral",
|
|
9877
9902
|
"groq",
|
|
9878
9903
|
"deepseek",
|
|
9879
|
-
"kimi"
|
|
9904
|
+
"kimi",
|
|
9905
|
+
"qwen"
|
|
9880
9906
|
];
|
|
9881
9907
|
var USAGE_HINT = "Pass a 'providers' array to confer/debate/plan/review to pick an ad-hoc subset, e.g. providers=['openai','gemini']. Omit the field to use the configured active set.";
|
|
9882
9908
|
function runListProviders(args, opts) {
|
|
@@ -11448,7 +11474,7 @@ var CHECK_INTERVAL_SECONDS = 3 * 24 * 60 * 60;
|
|
|
11448
11474
|
var DEFAULT_PACKAGE = "crosscheck-cli";
|
|
11449
11475
|
var FETCH_TIMEOUT_MS = 3e3;
|
|
11450
11476
|
function engineVersion() {
|
|
11451
|
-
return true ? "0.2.
|
|
11477
|
+
return true ? "0.2.11" : "0.0.0-dev";
|
|
11452
11478
|
}
|
|
11453
11479
|
function defaultUpdateCachePath() {
|
|
11454
11480
|
const base = process.env["CROSSCHECK_DATA_DIR"] || path9.join(os.homedir() || os.tmpdir(), ".crosscheck");
|
|
@@ -13069,7 +13095,8 @@ var KEY_ENV = {
|
|
|
13069
13095
|
groq: "GROQ_API_KEY",
|
|
13070
13096
|
deepseek: "DEEPSEEK_API_KEY",
|
|
13071
13097
|
mistral: "MISTRAL_API_KEY",
|
|
13072
|
-
kimi: "KIMI_API_KEY"
|
|
13098
|
+
kimi: "KIMI_API_KEY",
|
|
13099
|
+
qwen: "QWEN_API_KEY"
|
|
13073
13100
|
};
|
|
13074
13101
|
function createCrosscheck(opts) {
|
|
13075
13102
|
const env = {};
|