@chessceo/mcp 0.49.6 → 0.49.8
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/analysis/auto.js +4 -3
- package/dist/analysis/deep.js +3 -1
- package/dist/analysis/file_handle.js +11 -0
- package/dist/index.js +2 -3
- package/dist/tools.js +13 -12
- package/docs/engine-usage.md +19 -0
- package/package.json +1 -1
package/dist/analysis/auto.js
CHANGED
|
@@ -24,6 +24,7 @@
|
|
|
24
24
|
// a mid-run crash / cancellation doesn't lose the whole walk.
|
|
25
25
|
import { authedRequest, fetchGame } from "../http.js";
|
|
26
26
|
import { analysisToStoredEval } from "./response.js";
|
|
27
|
+
import { analyseRouting } from "./file_handle.js";
|
|
27
28
|
import { parsePGN } from "../pgn/parser.js";
|
|
28
29
|
import { buildIdIndex, positionKey, resolveNodeId, ROOT_ID } from "../pgn/paths.js";
|
|
29
30
|
const evalJobs = new Map();
|
|
@@ -146,7 +147,7 @@ export async function autoEvaluate(args, applyBatchMutations) {
|
|
|
146
147
|
// Unawaited — runs concurrently with the tool response. Any thrown
|
|
147
148
|
// error gets recorded on the job so the LLM's status poll surfaces
|
|
148
149
|
// it instead of the process seeing an unhandled rejection.
|
|
149
|
-
void runEvalJob(job, id, targets, movetimeMs, applyBatchMutations).catch(err => {
|
|
150
|
+
void runEvalJob(job, id, targets, movetimeMs, applyBatchMutations, analyseRouting(args)).catch(err => {
|
|
150
151
|
job.status = "error";
|
|
151
152
|
job.error = err instanceof Error ? err.message : String(err);
|
|
152
153
|
job.finishedAt = Date.now();
|
|
@@ -167,7 +168,7 @@ export async function autoEvaluate(args, applyBatchMutations) {
|
|
|
167
168
|
// against the per-combo semaphore in the backend anyway), checkpoints
|
|
168
169
|
// every SAVE_EVERY_N successfully-evaluated nodes so partial progress
|
|
169
170
|
// is durable, and re-anchors the version after each save.
|
|
170
|
-
async function runEvalJob(job, fileId, targets, movetimeMs, applyBatchMutations) {
|
|
171
|
+
async function runEvalJob(job, fileId, targets, movetimeMs, applyBatchMutations, routing) {
|
|
171
172
|
const pending = [];
|
|
172
173
|
const flush = async () => {
|
|
173
174
|
if (pending.length === 0)
|
|
@@ -196,7 +197,7 @@ async function runEvalJob(job, fileId, targets, movetimeMs, applyBatchMutations)
|
|
|
196
197
|
if (job.cancelled || aborted)
|
|
197
198
|
break;
|
|
198
199
|
try {
|
|
199
|
-
const analysis = await authedRequest("POST", "/api/agent/cloud-engines/analyse", { fen: t.fen, movetime_ms: movetimeMs, multipv: 1 });
|
|
200
|
+
const analysis = await authedRequest("POST", "/api/agent/cloud-engines/analyse", { fen: t.fen, movetime_ms: movetimeMs, multipv: 1, ...routing });
|
|
200
201
|
const ev = analysisToStoredEval(analysis);
|
|
201
202
|
if (ev) {
|
|
202
203
|
pending.push({ op: "set_ceo_eval", node_id: t.nodeId, ceoEval: ev });
|
package/dist/analysis/deep.js
CHANGED
|
@@ -16,7 +16,7 @@
|
|
|
16
16
|
// Extracted from index.ts in v0.44 as part of the file split.
|
|
17
17
|
import { authedRequest } from "../http.js";
|
|
18
18
|
import { analysisToStoredEval, convertCloudSnapshotResponse } from "./response.js";
|
|
19
|
-
import { resolveFromNodeOrFen, storeEvalOnNode, } from "./file_handle.js";
|
|
19
|
+
import { analyseRouting, resolveFromNodeOrFen, storeEvalOnNode, } from "./file_handle.js";
|
|
20
20
|
const deepJobs = new Map();
|
|
21
21
|
const DEEP_JOB_TTL_MS = 15 * 60 * 1000;
|
|
22
22
|
function newDeepJobId() {
|
|
@@ -43,6 +43,7 @@ export async function deepAnalyseStart(args) {
|
|
|
43
43
|
const jobId = newDeepJobId();
|
|
44
44
|
const job = {
|
|
45
45
|
id: jobId,
|
|
46
|
+
routing: analyseRouting(args),
|
|
46
47
|
status: "running",
|
|
47
48
|
fileHandle: resolved.file,
|
|
48
49
|
fen,
|
|
@@ -73,6 +74,7 @@ async function runDeepJob(job) {
|
|
|
73
74
|
movetime_ms: job.movetimeMs,
|
|
74
75
|
stockfish_multipv: job.multipv,
|
|
75
76
|
engines: ["stockfish"],
|
|
77
|
+
...job.routing,
|
|
76
78
|
};
|
|
77
79
|
let raw;
|
|
78
80
|
try {
|
|
@@ -154,3 +154,14 @@ export async function storeEvalOnNode(handle, ev) {
|
|
|
154
154
|
return [];
|
|
155
155
|
}
|
|
156
156
|
}
|
|
157
|
+
// Routing fields for POST /api/agent/cloud-engines/analyse: which rental to
|
|
158
|
+
// run on and which engine legs to run. Omitted fields let the backend pick.
|
|
159
|
+
export function analyseRouting(args) {
|
|
160
|
+
const out = {};
|
|
161
|
+
if (typeof args.contract_id === "string" && args.contract_id.trim()) {
|
|
162
|
+
out.contract_id = args.contract_id.trim();
|
|
163
|
+
}
|
|
164
|
+
if (Array.isArray(args.engines))
|
|
165
|
+
out.engines = args.engines;
|
|
166
|
+
return out;
|
|
167
|
+
}
|
package/dist/index.js
CHANGED
|
@@ -25,7 +25,7 @@ import { authContext, authedRequest, deleteGame, fetchGame, get, makeFileId, res
|
|
|
25
25
|
import { analysisToStoredEval, capPvsInResponse, convertCloudSnapshotResponse, fetchCompactEval, } from "./analysis/response.js";
|
|
26
26
|
import { autoEvaluate, autoEvaluateCancel, autoEvaluateStatus, } from "./analysis/auto.js";
|
|
27
27
|
import { deepAnalyseCancel, deepAnalyseStart, deepAnalyseStatus, } from "./analysis/deep.js";
|
|
28
|
-
import { getNodeByPath, resolveFromNodeOrFen, storeEvalOnNode, } from "./analysis/file_handle.js";
|
|
28
|
+
import { analyseRouting, getNodeByPath, resolveFromNodeOrFen, storeEvalOnNode, } from "./analysis/file_handle.js";
|
|
29
29
|
import { findPositionInCourses, readCourseAtPosition, runSfEval } from "./courses.js";
|
|
30
30
|
import { applyBatchMutations, applyMutation, argNodeId, } from "./prep/mutations.js";
|
|
31
31
|
import { listNodes, listTranspositions, readPrepFile, } from "./prep/read.js";
|
|
@@ -412,8 +412,7 @@ async function callToolInner(name, args) {
|
|
|
412
412
|
body.lc0_multipv = args.lc0_multipv;
|
|
413
413
|
if (typeof args.contempt === "number")
|
|
414
414
|
body.contempt = args.contempt;
|
|
415
|
-
|
|
416
|
-
body.engines = args.engines;
|
|
415
|
+
Object.assign(body, analyseRouting(args));
|
|
417
416
|
const raw = await authedRequest("POST", "/api/agent/cloud-engines/analyse", body);
|
|
418
417
|
const converted = convertCloudSnapshotResponse(raw, fen);
|
|
419
418
|
// PV cap: engine PVs beyond ~6 plies are speculative (the tail is
|
package/dist/tools.js
CHANGED
|
@@ -75,7 +75,7 @@ export const TOOLS = [
|
|
|
75
75
|
name: "get_prep_position",
|
|
76
76
|
description: "Query one position within a prep session created by `prepare_opponent`. Returns move statistics (frequency + win rate + last-played date per move) plus the actual games played from that position, in one call.\n\n" +
|
|
77
77
|
"Position input: prefer `file_id`+`node_id` when inside a prep file (server derives FEN from the tree). Otherwise pass `fen`.\n\n" +
|
|
78
|
-
"AUTO-EVAL: if a cloud
|
|
78
|
+
"AUTO-EVAL: if a cloud engine instance is running, the response includes `.eval` (whichever engines that rental provides — Stockfish, Lc0, or both) so you don't need a separate cloud_analyse.\n\n" +
|
|
79
79
|
"Reading the response — CRITICAL:\n" +
|
|
80
80
|
"• Win % is one weight, not a verdict. Sample size matters (3 games at 66% is noise; 300 at 55% is signal).\n" +
|
|
81
81
|
"• Prep is symmetric information — both sides see the same history. Assume the opponent knows the weakness you spotted.\n" +
|
|
@@ -118,7 +118,7 @@ export const TOOLS = [
|
|
|
118
118
|
"- `gm-classical` — GM classical games (both players ≥2500, real thinking-time). BEST for opening prep — every move is signal, avgElo ~2600 across all listed moves.\n" +
|
|
119
119
|
"- `main` — the whole 11.7M-game DB. Widest coverage but noisiest (includes 1000-Elo blunder-fests in the move stats). Use as fallback when gm-classical's totalCount is too small to be informative.\n\n" +
|
|
120
120
|
"Game movetext is trimmed to the moves AFTER the queried position (using each game's plyNumber). Saves ~70% of the bytes vs full movetext.\n\n" +
|
|
121
|
-
"AUTO-EVAL: if a cloud
|
|
121
|
+
"AUTO-EVAL: if a cloud engine instance is running, the response includes `.eval` with a compact read from the engine(s) it provides and the corresponding NAG. Do NOT fire cloud_analyse separately for the same FEN. When called with `file_id`+`node_id`, the eval is also auto-stored on that node's `ceoEval` — later readable via quote_engine_eval.",
|
|
122
122
|
inputSchema: {
|
|
123
123
|
type: "object",
|
|
124
124
|
properties: {
|
|
@@ -284,7 +284,7 @@ export const TOOLS = [
|
|
|
284
284
|
},
|
|
285
285
|
{
|
|
286
286
|
name: "list_cloud_engines",
|
|
287
|
-
description: "List the user's currently running cloud engines. Use before starting a new one, or to find the contract_id for stop_cloud_engine. `cloud_analyse`
|
|
287
|
+
description: "List the user's currently running cloud engines. Use before starting a new one, or to find the contract_id for stop_cloud_engine. `cloud_analyse` uses the only running rental that can serve the request unless you pass `contract_id`, so listing is only necessary when there are several.",
|
|
288
288
|
inputSchema: { type: "object", properties: {} },
|
|
289
289
|
},
|
|
290
290
|
{
|
|
@@ -303,9 +303,9 @@ export const TOOLS = [
|
|
|
303
303
|
},
|
|
304
304
|
{
|
|
305
305
|
name: "cloud_analyse",
|
|
306
|
-
description: "Runs a synchronous ~2s analysis on the user's running
|
|
306
|
+
description: "Runs a synchronous ~2s analysis on one of the user's running engine instances and returns the requested engines' final read for the FEN — depth, top-N candidate moves with scores (**scoreCp is White-POV centipawns**: +20 = White is +0.20 pawns better regardless of whose turn it is; matches the sign convention used everywhere else in this MCP, including the stored ceoEval). Mate is White-POV plies-to-mate (+5 = White mates in 5). Also returns each engine's principal variation.\n\n" +
|
|
307
307
|
"GROUNDING: every claim you make about a position must trace back to actual engine output from a call in THIS session. Don't invent evaluations, don't name 'best moves' you haven't seen the engine list, don't fabricate variations that 'look plausible.' Compute is cheap — call this 5-10 times while walking a tree rather than pattern-matching from your training data. When you don't have data for the position, either run the tool or say so; don't fill the gap with chess prose the user can't distinguish from measured output.\n\n" +
|
|
308
|
-
"
|
|
308
|
+
"Runs on any of the caller's running rentals — combo, stockfish-only, or lc0-only. Pass `contract_id` to choose one; without it the only rental that can serve the requested engines is used, and the error lists the candidates if there are several.\n\n" +
|
|
309
309
|
"How to read the response:\n" +
|
|
310
310
|
"• Stockfish is objective truth — trust it for 'does this line hold?' 'is there a tactic?' 'is this endgame drawn?' A Stockfish 0.00 means 'objectively equal', NOT 'trivial draw' — one side can still be much harder to play in practice.\n" +
|
|
311
311
|
"• Lc0 is practical eval — trust it for 'which side is easier?' 'which candidate is best when Stockfish shows several as equal?' Lc0 sees long-term positional factors Stockfish's fixed search can miss.\n" +
|
|
@@ -350,10 +350,11 @@ export const TOOLS = [
|
|
|
350
350
|
maximum: 100,
|
|
351
351
|
description: "Lc0 contempt bias. Signed 0-100 strength (same scale as the web UI's ContemptStrength slider — server multiplies by 8 to get the internal cp bias). 0 = objective (default). Positive favours White, negative favours Black. Typical: ±15 light nudge, ±30-60 real fighting play, ±80-100 maximum steer. Not applied to Stockfish. See engine_usage_primer for when to use.",
|
|
352
352
|
},
|
|
353
|
+
contract_id: { type: "string", description: "Which of your running rentals to use, from list_cloud_engines. Any engine shape works (combo, stockfish-only, lc0-only). Omit it when only one rental can serve the request; if several can, the error lists their contract ids." },
|
|
353
354
|
engines: {
|
|
354
355
|
type: "array",
|
|
355
356
|
items: { type: "string", enum: ["stockfish", "lc0"] },
|
|
356
|
-
description: "Which engines to run. Default = both. Use `[\"lc0\"]` to skip Stockfish (e.g. while a deep_analyse job is holding the SF slot on the same
|
|
357
|
+
description: "Which engines to run. Default = both. Use `[\"lc0\"]` to skip Stockfish (e.g. while a deep_analyse job is holding the SF slot on the same rental). Use `[\"stockfish\"]` when only the objective read matters. The skipped engine's field is omitted from the response.",
|
|
357
358
|
},
|
|
358
359
|
pv_max_plies: {
|
|
359
360
|
type: "integer",
|
|
@@ -721,14 +722,14 @@ export const TOOLS = [
|
|
|
721
722
|
},
|
|
722
723
|
{
|
|
723
724
|
name: "auto_evaluate",
|
|
724
|
-
description: "Walk the tree from `node_id` (default `'r'` = whole file) and populate the persistent `ceoEval` on every descendant via cloud_analyse. Requires a running cloud
|
|
725
|
+
description: "Walk the tree from `node_id` (default `'r'` = whole file) and populate the persistent `ceoEval` on every descendant via cloud_analyse. Requires a running cloud engine instance (any shape; pass contract_id to choose).\n\n" +
|
|
725
726
|
"**Async job — returns immediately.** Response: `{ job_id, target_count, status: 'running', estimated_seconds }`. Then poll `auto_evaluate_status(job_id)` until `done: true`. Cancel a run with `auto_evaluate_cancel(job_id)` — partial progress is preserved. Do useful other work between polls (write more of the tree, walk the opponent's repertoire) — the engine runs in the background.\n\n" +
|
|
726
727
|
"Progress is checkpointed to the prep file every 8 successfully-evaluated nodes, so a cancel / crash / MCP restart mid-run leaves the tree partially populated rather than losing everything. On MCP restart the job record disappears; re-run auto_evaluate and `only_missing=true` naturally skips what was already saved.\n\n" +
|
|
727
728
|
"**Does NOT set visible NAGs.** NAG placement is your call, not the engine's — an opening tree full of 0.00 positions doesn't need a `$10` (=) glyph on every move. Use quote_engine_eval on individual nodes before writing prose that references engine numbers.\n\n" +
|
|
728
|
-
"Costs real money — one cloud_analyse per node. A 200-node walk at default movetime is ~5 min of engine time (calls serialise on the per-
|
|
729
|
+
"Costs real money — one cloud_analyse per node. A 200-node walk at default movetime is ~5 min of engine time (calls serialise on the per-engine semaphore in the backend).",
|
|
729
730
|
inputSchema: {
|
|
730
731
|
type: "object",
|
|
731
|
-
properties: {
|
|
732
|
+
properties: { contract_id: { type: "string", description: "Which of your running rentals to use, from list_cloud_engines. Any engine shape works (combo, stockfish-only, lc0-only). Omit it when only one rental can serve the request; if several can, the error lists their contract ids." },
|
|
732
733
|
id: { type: "string" },
|
|
733
734
|
node_id: { type: "string", description: "Subtree root (default 'r' = whole file)." },
|
|
734
735
|
only_missing: { type: "boolean", description: "Skip nodes that already carry a stored ceoEval (default true)." },
|
|
@@ -762,12 +763,12 @@ export const TOOLS = [
|
|
|
762
763
|
},
|
|
763
764
|
{
|
|
764
765
|
name: "deep_analyse",
|
|
765
|
-
description: "Start a long Stockfish think on a single position (up to 5 min movetime). Returns a `job_id` immediately; poll `deep_analyse_status(job_id)` for the result, cancel with `deep_analyse_cancel(job_id)`. Runs SF only — Lc0 doesn't benefit from long thinks past a handful of seconds — and **holds only the SF engine slot on the
|
|
766
|
+
description: "Start a long Stockfish think on a single position (up to 5 min movetime). Returns a `job_id` immediately; poll `deep_analyse_status(job_id)` for the result, cancel with `deep_analyse_cancel(job_id)`. Runs SF only — Lc0 doesn't benefit from long thinks past a handful of seconds — and **holds only the SF engine slot on the rental, so `cloud_analyse(..., engines: [\"lc0\"])` stays available for other work in parallel**.\n\n" +
|
|
766
767
|
"Use this when a specific critical position deserves depth — a novelty candidate, a hairy tactical shot, a difficult endgame — and you want Stockfish at depth 35+ rather than the ~depth 22 you get from a 2s cloud_analyse. Movetime is in ms; typical: 30_000-60_000 for 'careful check', 120_000-300_000 for 'find the truth'.\n\n" +
|
|
767
768
|
"Result shape when done matches cloud_analyse's Stockfish leg (depth, top-N candidates with scoreCp/mate, best move, PV). Auto-stores the eval on `file_id`+`node_id` when both are supplied, same as cloud_analyse.",
|
|
768
769
|
inputSchema: {
|
|
769
770
|
type: "object",
|
|
770
|
-
properties: {
|
|
771
|
+
properties: { contract_id: { type: "string", description: "Which of your running rentals to use, from list_cloud_engines. Any engine shape works (combo, stockfish-only, lc0-only). Omit it when only one rental can serve the request; if several can, the error lists their contract ids." },
|
|
771
772
|
file_id: { type: "string", description: "Prep file id. Combine with `node_id` to derive FEN from the tree AND persist the result on the node's ceoEval." },
|
|
772
773
|
node_id: { type: "string", description: "Node id inside `file_id`. Root is 'r'. When set, overrides `fen`/`moves`." },
|
|
773
774
|
fen: { type: "string", description: "Position as FEN. Only used if `node_id` is not set." },
|
|
@@ -926,7 +927,7 @@ export const TOOLS = [
|
|
|
926
927
|
{
|
|
927
928
|
name: "prep_snapshot",
|
|
928
929
|
description: "One call, three parallel fetches at the same position: opponent's stats on their side, your stats on your side, and the 11.7M-game general database at that position. Use this while walking the opening tree — one round trip instead of three separate calls, and you can compare the three views directly (e.g. opponent has 2 games here but the general DB has 8k → prep candidate).\n\n" +
|
|
929
|
-
"AUTO-EVAL: if a cloud
|
|
930
|
+
"AUTO-EVAL: if a cloud engine instance is running, the response includes a top-level `.eval` (the engines that rental provides, read at the shared position) so you get four signals in one call. Do NOT fire cloud_analyse separately for the same FEN.",
|
|
930
931
|
inputSchema: {
|
|
931
932
|
type: "object",
|
|
932
933
|
properties: {
|
package/docs/engine-usage.md
CHANGED
|
@@ -192,6 +192,25 @@ Override the defaults with `stockfish_multipv` / `lc0_multipv` when a specific p
|
|
|
192
192
|
|
|
193
193
|
The skipped engine's field is omitted from the response (not present as an empty object).
|
|
194
194
|
|
|
195
|
+
## Renting and choosing engines
|
|
196
|
+
|
|
197
|
+
Analysis needs a running rental. Start one with `start_cloud_engine`, using a `machine_type` SKU from `list_cloud_machine_options`. Each SKU provides one of three shapes, shown in its `engine` field:
|
|
198
|
+
|
|
199
|
+
- **`stockfish`**: Stockfish only. Fits when the question is objective (is this tactic real, does this defense hold).
|
|
200
|
+
- **`lc0`**: Lc0 only. Fits the practical, human-feel read, or running alongside a `deep_analyse` that holds the Stockfish side.
|
|
201
|
+
- **`combo`**: both engines in one container. The default, and the right choice for prep decisions, where both reads are needed on the same position.
|
|
202
|
+
|
|
203
|
+
Rule of thumb: start a `combo` rental unless the task only needs one engine. A single-engine rental gives one read per position, so a prep walk that needs both reads on a single-engine rental takes two passes.
|
|
204
|
+
|
|
205
|
+
`list_cloud_machine_options` returns the price for each SKU. Show the user the price and get their confirmation before starting a rental; every rental bills per second until `stop_cloud_engine`.
|
|
206
|
+
|
|
207
|
+
Routing to a rental:
|
|
208
|
+
|
|
209
|
+
- `cloud_analyse`, `auto_evaluate`, and `deep_analyse` each take an optional `contract_id` (from `list_cloud_engines`).
|
|
210
|
+
- Without `contract_id`, the call uses the only running rental that can serve the request. If several can, the error lists their contract ids, and you pick one.
|
|
211
|
+
- Any shape works. With `engines` set, the rental must provide those engines: asking for `engines: ["lc0"]` on a stockfish-only rental fails.
|
|
212
|
+
- Do not guess contract ids. Call `list_cloud_engines` first if more than one rental is running.
|
|
213
|
+
|
|
195
214
|
## Worked example
|
|
196
215
|
|
|
197
216
|
User is preparing Black against a 2600 opponent who plays 1.e4 c5 2.Nf3 d6 3.d4 cxd4 4.Nxd4 Nf6 5.Nc3 a6 6.Be3 e5. You want to know if 7.Nb3 or 7.Nf3 is more testing.
|
package/package.json
CHANGED