axstack 0.20.29 → 0.20.31
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +6 -4
- package/bin/axstack.js +1 -0
- package/docs/installation.md +20 -9
- package/docs/workflows.md +42 -17
- package/package.json +1 -1
- package/profiles/presets/claude-only.json +81 -45
- package/profiles/presets/codex-only.json +78 -42
- package/profiles/presets/mixed.json +83 -47
- package/skills/axstack/references/autopilot.md +108 -0
- package/skills/axstack/references/candidate-publication.md +9 -0
- package/skills/axstack/references/contracts.md +5 -1
- package/skills/axstack/references/diligence.md +23 -0
- package/skills/axstack/references/lifecycle.md +2 -2
- package/skills/axstack/references/orca-runtime.md +25 -6
- package/skills/axstack/references/role-roster.md +37 -0
- package/skills/axstack/references/routing.md +19 -31
- package/skills/axstack/references/run-record.md +2 -0
- package/skills/axstack/scripts/resolve-models.js +46 -0
- package/skills/axstack-align/SKILL.md +8 -4
- package/skills/axstack-audit/SKILL.md +15 -2
- package/skills/axstack-implement/SKILL.md +30 -11
- package/skills/axstack-improve/SKILL.md +5 -0
- package/skills/axstack-relay/SKILL.md +9 -2
- package/skills/axstack-research/SKILL.md +14 -2
- package/skills/axstack-review/SKILL.md +29 -13
- package/skills/axstack-spec/SKILL.md +8 -1
- package/skills/axstack-tickets/SKILL.md +9 -3
- package/skills/axstack-watch/SKILL.md +27 -13
- package/skills/axstack-watch/references/watch-runtime.md +11 -4
- package/src/installer.js +8 -0
- package/src/roles.js +39 -11
|
@@ -5,10 +5,10 @@
|
|
|
5
5
|
"id": "axstack-advisor-astra",
|
|
6
6
|
"name": "Axstack Astra adviser",
|
|
7
7
|
"provider": "codex",
|
|
8
|
-
"model": "gpt-6-astra",
|
|
9
8
|
"modeId": "full-access",
|
|
10
9
|
"thinkingOptionId": "high",
|
|
11
|
-
"notes": "Independent Astra adviser for Align, Spec, and unresolved consequential decisions. Receives the same bounded evidence and question as Opus; reuse only unchanged receipts."
|
|
10
|
+
"notes": "Independent Astra adviser for Align, Spec, and unresolved consequential decisions. Receives the same bounded evidence and question as Opus; reuse only unchanged receipts.",
|
|
11
|
+
"modelClass": "astra"
|
|
12
12
|
},
|
|
13
13
|
{
|
|
14
14
|
"id": "axstack-advisor-opus",
|
|
@@ -23,73 +23,91 @@
|
|
|
23
23
|
"id": "axstack-owner",
|
|
24
24
|
"name": "Axstack PR owner",
|
|
25
25
|
"provider": "codex",
|
|
26
|
-
"model": "gpt-6-sol",
|
|
27
26
|
"modeId": "full-access",
|
|
28
27
|
"thinkingOptionId": "high",
|
|
29
|
-
"notes": "Persistent PR owner: one owner per PR, accountable for candidate, fixes, verification evidence, and monitoring. May delegate coding but never edits a worker-owned candidate concurrently. Launches eligible independent reviewers."
|
|
28
|
+
"notes": "Persistent PR owner: one owner per PR, accountable for candidate, fixes, verification evidence, and monitoring. May delegate coding but never edits a worker-owned candidate concurrently. Launches eligible independent reviewers.",
|
|
29
|
+
"modelClass": "sol"
|
|
30
30
|
},
|
|
31
31
|
{
|
|
32
32
|
"id": "axstack-author",
|
|
33
33
|
"name": "Axstack author",
|
|
34
34
|
"provider": "codex",
|
|
35
|
-
"model": "gpt-6-sol",
|
|
36
35
|
"modeId": "full-access",
|
|
37
36
|
"thinkingOptionId": "high",
|
|
38
|
-
"notes": "Ordinary implementation and repairs. Uses strict red-green-refactor and remains the exclusive writer for a candidate.
|
|
37
|
+
"notes": "Ordinary implementation and repairs. Uses strict red-green-refactor and remains the exclusive writer for a candidate. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
38
|
+
"modelClass": "sol"
|
|
39
39
|
},
|
|
40
40
|
{
|
|
41
41
|
"id": "axstack-reviewer-primary",
|
|
42
42
|
"name": "Axstack reviewer (primary)",
|
|
43
43
|
"provider": "codex",
|
|
44
|
-
"model": "gpt-6-sol",
|
|
45
44
|
"modeId": "full-access",
|
|
46
45
|
"thinkingOptionId": "high",
|
|
47
|
-
"notes": "Primary reviewer in the ordered codex-only peer pair: Sol high followed by Luna xhigh. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated."
|
|
46
|
+
"notes": "Primary reviewer in the ordered codex-only peer pair: Sol high followed by Luna xhigh. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated.",
|
|
47
|
+
"modelClass": "sol"
|
|
48
48
|
},
|
|
49
49
|
{
|
|
50
50
|
"id": "axstack-reviewer-secondary",
|
|
51
51
|
"name": "Axstack reviewer (secondary)",
|
|
52
52
|
"provider": "codex",
|
|
53
|
-
"model": "gpt-6-luna",
|
|
54
53
|
"modeId": "full-access",
|
|
55
54
|
"thinkingOptionId": "xhigh",
|
|
56
|
-
"notes": "Secondary reviewer in the ordered codex-only peer pair: Sol high followed by Luna xhigh. Eligible authored reviewer for a Sol-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated."
|
|
55
|
+
"notes": "Secondary reviewer in the ordered codex-only peer pair: Sol high followed by Luna xhigh. Eligible authored reviewer for a Sol-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated.",
|
|
56
|
+
"modelClass": "luna"
|
|
57
|
+
},
|
|
58
|
+
{
|
|
59
|
+
"id": "axstack-diligence",
|
|
60
|
+
"name": "Axstack diligence checker",
|
|
61
|
+
"provider": "codex",
|
|
62
|
+
"modeId": "full-access",
|
|
63
|
+
"thinkingOptionId": "high",
|
|
64
|
+
"notes": "Read-only diligence for exact-revision PRs and bounded research, spec, ticket, receipt, and release claims. Returns PASS or FINDINGS with evidence; never authors or edits.",
|
|
65
|
+
"modelClass": "sol"
|
|
57
66
|
},
|
|
58
67
|
{
|
|
59
68
|
"id": "axstack-checker",
|
|
60
69
|
"name": "Axstack tracker checker",
|
|
61
70
|
"provider": "codex",
|
|
62
|
-
"model": "gpt-6-luna",
|
|
63
71
|
"modeId": "full-access",
|
|
64
72
|
"thinkingOptionId": "low",
|
|
65
|
-
"notes": "Report-only discrepancy checker for the selected external tracker (Linear or GitHub Issues). Never mutates the tracker; the driver independently verifies evidence before applying updates. A null model means explicit user selection is required before dispatch and must never launch a provider default."
|
|
73
|
+
"notes": "Report-only discrepancy checker for the selected external tracker (Linear or GitHub Issues). Never mutates the tracker; the driver independently verifies evidence before applying updates. A null model means explicit user selection is required before dispatch and must never launch a provider default.",
|
|
74
|
+
"modelClass": "luna"
|
|
66
75
|
},
|
|
67
76
|
{
|
|
68
77
|
"id": "axstack-research-requirements",
|
|
69
78
|
"name": "Axstack research requirements",
|
|
70
79
|
"provider": "codex",
|
|
71
|
-
"model": "gpt-6-astra",
|
|
72
80
|
"modeId": "full-access",
|
|
73
81
|
"thinkingOptionId": "medium",
|
|
74
|
-
"notes": "Research requirements analyst: scopes bounded questions and acceptance for a research task.
|
|
82
|
+
"notes": "Research requirements analyst: scopes bounded questions and acceptance for a research task. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
83
|
+
"modelClass": "astra"
|
|
75
84
|
},
|
|
76
85
|
{
|
|
77
86
|
"id": "axstack-research-code",
|
|
78
87
|
"name": "Axstack research code",
|
|
79
88
|
"provider": "codex",
|
|
80
|
-
"model": "gpt-6-sol",
|
|
81
89
|
"modeId": "full-access",
|
|
82
90
|
"thinkingOptionId": "high",
|
|
83
|
-
"notes": "Research code investigator: verifies behavior against inspected code and executable evidence.
|
|
91
|
+
"notes": "Research code investigator: verifies behavior against inspected code and executable evidence. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
92
|
+
"modelClass": "sol"
|
|
93
|
+
},
|
|
94
|
+
{
|
|
95
|
+
"id": "axstack-research-code-sol",
|
|
96
|
+
"name": "Axstack research code Sol",
|
|
97
|
+
"provider": "codex",
|
|
98
|
+
"modeId": "full-access",
|
|
99
|
+
"thinkingOptionId": "high",
|
|
100
|
+
"notes": "Independent Sol high research code investigator: verifies code behavior and executable evidence on the same bounded brief without cross-reading. Reports findings for driver reconciliation.",
|
|
101
|
+
"modelClass": "sol"
|
|
84
102
|
},
|
|
85
103
|
{
|
|
86
104
|
"id": "axstack-research-web",
|
|
87
105
|
"name": "Axstack research web",
|
|
88
106
|
"provider": "codex",
|
|
89
|
-
"model": "gpt-6-sol",
|
|
90
107
|
"modeId": "full-access",
|
|
91
108
|
"thinkingOptionId": "low",
|
|
92
|
-
"notes": "Research web reader: gathers primary-source facts efficiently.
|
|
109
|
+
"notes": "Research web reader: gathers primary-source facts efficiently. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
110
|
+
"modelClass": "sol"
|
|
93
111
|
},
|
|
94
112
|
{
|
|
95
113
|
"id": "axstack-research-web-google",
|
|
@@ -113,109 +131,127 @@
|
|
|
113
131
|
"id": "axstack-explainer",
|
|
114
132
|
"name": "Axstack explainer",
|
|
115
133
|
"provider": "codex",
|
|
116
|
-
"model": "gpt-6-sol",
|
|
117
134
|
"modeId": "full-access",
|
|
118
135
|
"thinkingOptionId": "high",
|
|
119
|
-
"notes": "Complex visual explanation author: traces systems, changes, and implementation gaps in requested artifacts and verifies rendered behavior where applicable.
|
|
136
|
+
"notes": "Complex visual explanation author: traces systems, changes, and implementation gaps in requested artifacts and verifies rendered behavior where applicable. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
137
|
+
"modelClass": "sol"
|
|
120
138
|
},
|
|
121
139
|
{
|
|
122
140
|
"id": "axstack-explainer-review",
|
|
123
141
|
"name": "Axstack explainer reviewer",
|
|
124
142
|
"provider": "codex",
|
|
125
|
-
"model": "gpt-6-luna",
|
|
126
143
|
"modeId": "full-access",
|
|
127
144
|
"thinkingOptionId": "xhigh",
|
|
128
|
-
"notes": "Independent visual explanation reviewer: checks the exact artifact for text and source fidelity. The rendered pass belongs to axstack-ui-verifier. Any artifact change invalidates its review.
|
|
145
|
+
"notes": "Independent visual explanation reviewer: checks the exact artifact for text and source fidelity. The rendered pass belongs to axstack-ui-verifier. Any artifact change invalidates its review. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
146
|
+
"modelClass": "luna"
|
|
129
147
|
},
|
|
130
148
|
{
|
|
131
149
|
"id": "axstack-ui-verifier",
|
|
132
150
|
"name": "Axstack UI verifier",
|
|
133
151
|
"provider": "codex",
|
|
134
|
-
"model": "gpt-6-sol",
|
|
135
152
|
"modeId": "full-access",
|
|
136
153
|
"thinkingOptionId": "medium",
|
|
137
|
-
"notes": "Read-only UI verifier: runs Playwright or the browser against the given build, URL, or artifact; captures screenshots, interactions, accessibility, desktop/mobile, and reduced-motion evidence in the dispatch evidence folder; returns a verdict with evidence paths. Never edits source.
|
|
154
|
+
"notes": "Read-only UI verifier: runs Playwright or the browser against the given build, URL, or artifact; captures screenshots, interactions, accessibility, desktop/mobile, and reduced-motion evidence in the dispatch evidence folder; returns a verdict with evidence paths. Never edits source. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
155
|
+
"modelClass": "sol"
|
|
138
156
|
},
|
|
139
157
|
{
|
|
140
158
|
"id": "axstack-explore-codebase",
|
|
141
159
|
"name": "Axstack codebase explorer",
|
|
142
160
|
"provider": "codex",
|
|
143
|
-
"model": "gpt-6-sol",
|
|
144
161
|
"modeId": "full-access",
|
|
145
162
|
"thinkingOptionId": "high",
|
|
146
|
-
"notes": "Codebase mapper: explores repository structure and interfaces for research and handoff context.
|
|
163
|
+
"notes": "Codebase mapper: explores repository structure and interfaces for research and handoff context. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
164
|
+
"modelClass": "sol"
|
|
147
165
|
},
|
|
148
166
|
{
|
|
149
167
|
"id": "axstack-explore-execution",
|
|
150
168
|
"name": "Axstack execution explorer",
|
|
151
169
|
"provider": "codex",
|
|
152
|
-
"model": "gpt-6-sol",
|
|
153
170
|
"modeId": "full-access",
|
|
154
171
|
"thinkingOptionId": "high",
|
|
155
|
-
"notes": "Execution explorer: runs bounded checks of runtime behavior where authorized.
|
|
172
|
+
"notes": "Execution explorer: runs bounded checks of runtime behavior where authorized. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
173
|
+
"modelClass": "sol"
|
|
174
|
+
},
|
|
175
|
+
{
|
|
176
|
+
"id": "axstack-explore-execution-sol",
|
|
177
|
+
"name": "Axstack execution explorer Sol",
|
|
178
|
+
"provider": "codex",
|
|
179
|
+
"modeId": "full-access",
|
|
180
|
+
"thinkingOptionId": "high",
|
|
181
|
+
"notes": "Independent Sol high execution explorer: runs authorized bounded runtime checks on the same brief without cross-reading. Reports findings for driver reconciliation.",
|
|
182
|
+
"modelClass": "sol"
|
|
156
183
|
},
|
|
157
184
|
{
|
|
158
185
|
"id": "axstack-monitor",
|
|
159
186
|
"name": "Axstack monitor",
|
|
160
187
|
"provider": "codex",
|
|
161
|
-
"model": "gpt-6-sol",
|
|
162
188
|
"modeId": "full-access",
|
|
163
189
|
"thinkingOptionId": "low",
|
|
164
|
-
"notes": "Optional independent read-only observer for a standalone PR watch. Reads GitHub, feedback, and checks, persists event IDs, and wakes the owner only for a new actionable event. Never sends, authors, reviews, replies, or acts as either reusable PR manager. Healthy snapshots stay quiet. Chat-run mode: one same-host native read-only observer per Run; fresh finite passes report precise deltas internally to the original Run/driver and may disable/read back only their own automation at verified stop. No repair, dispatch, public notification, or replacement coordinator. Effective scheduled model/effort and wake require live proof."
|
|
190
|
+
"notes": "Optional independent read-only observer for a standalone PR watch. Reads GitHub, feedback, and checks, persists event IDs, and wakes the owner only for a new actionable event. Never sends, authors, reviews, replies, or acts as either reusable PR manager. Healthy snapshots stay quiet. Chat-run mode: one same-host native read-only observer per Run; fresh finite passes report precise deltas internally to the original Run/driver and may disable/read back only their own automation at verified stop. No repair, dispatch, public notification, or replacement coordinator. Effective scheduled model/effort and wake require live proof.",
|
|
191
|
+
"modelClass": "sol"
|
|
165
192
|
},
|
|
166
193
|
{
|
|
167
194
|
"id": "axstack-auditor",
|
|
168
195
|
"name": "Axstack auditor",
|
|
169
196
|
"provider": "codex",
|
|
170
|
-
"model": "gpt-6-luna",
|
|
171
197
|
"modeId": "full-access",
|
|
172
198
|
"thinkingOptionId": "xhigh",
|
|
173
|
-
"notes": "Read-only end-of-run and checkpoint auditor. Collects scope and outcome evidence with counts and denominators and reports PASS, FAIL, or UNKNOWN without inventing numbers. Never edits, merges, activates, or audits itself."
|
|
199
|
+
"notes": "Read-only end-of-run and checkpoint auditor. Collects scope and outcome evidence with counts and denominators and reports PASS, FAIL, or UNKNOWN without inventing numbers. Never edits, merges, activates, or audits itself.",
|
|
200
|
+
"modelClass": "luna"
|
|
201
|
+
},
|
|
202
|
+
{
|
|
203
|
+
"id": "axstack-auditor-sol",
|
|
204
|
+
"name": "Axstack auditor Sol",
|
|
205
|
+
"provider": "codex",
|
|
206
|
+
"modeId": "full-access",
|
|
207
|
+
"thinkingOptionId": "high",
|
|
208
|
+
"notes": "Independent Sol high read-only auditor: checks the same bounded run evidence without cross-reading. Reports findings for driver reconciliation; never edits, merges, or activates.",
|
|
209
|
+
"modelClass": "sol"
|
|
174
210
|
},
|
|
175
211
|
{
|
|
176
212
|
"id": "axstack-debug-investigator-1",
|
|
177
213
|
"name": "Axstack debug investigator 1",
|
|
178
214
|
"provider": "codex",
|
|
179
|
-
"model": "gpt-6-sol",
|
|
180
215
|
"modeId": "full-access",
|
|
181
216
|
"thinkingOptionId": "high",
|
|
182
|
-
"notes": "Debug investigator seat 1. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. This preset repeats
|
|
217
|
+
"notes": "Debug investigator seat 1. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. This preset repeats the Sol class at high effort because it has fewer model families.",
|
|
218
|
+
"modelClass": "sol"
|
|
183
219
|
},
|
|
184
220
|
{
|
|
185
221
|
"id": "axstack-debug-investigator-2",
|
|
186
222
|
"name": "Axstack debug investigator 2",
|
|
187
223
|
"provider": "codex",
|
|
188
|
-
"model": "gpt-6-sol",
|
|
189
224
|
"modeId": "full-access",
|
|
190
225
|
"thinkingOptionId": "high",
|
|
191
|
-
"notes": "Debug investigator seat 2. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. This preset repeats
|
|
226
|
+
"notes": "Debug investigator seat 2. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. This preset repeats the Sol class at high effort because it has fewer model families.",
|
|
227
|
+
"modelClass": "sol"
|
|
192
228
|
},
|
|
193
229
|
{
|
|
194
230
|
"id": "axstack-debug-investigator-3",
|
|
195
231
|
"name": "Axstack debug investigator 3",
|
|
196
232
|
"provider": "codex",
|
|
197
|
-
"model": "gpt-6-sol",
|
|
198
233
|
"modeId": "full-access",
|
|
199
234
|
"thinkingOptionId": "high",
|
|
200
|
-
"notes": "Debug investigator seat 3. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. This preset repeats
|
|
235
|
+
"notes": "Debug investigator seat 3. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. This preset repeats the Sol class at high effort because it has fewer model families.",
|
|
236
|
+
"modelClass": "sol"
|
|
201
237
|
},
|
|
202
238
|
{
|
|
203
239
|
"id": "axstack-debug-investigator-4",
|
|
204
240
|
"name": "Axstack debug investigator 4",
|
|
205
241
|
"provider": "codex",
|
|
206
|
-
"model": "gpt-6-sol",
|
|
207
242
|
"modeId": "full-access",
|
|
208
243
|
"thinkingOptionId": "high",
|
|
209
|
-
"notes": "Debug investigator seat 4. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. This preset repeats
|
|
244
|
+
"notes": "Debug investigator seat 4. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. This preset repeats the Sol class at high effort because it has fewer model families.",
|
|
245
|
+
"modelClass": "sol"
|
|
210
246
|
},
|
|
211
247
|
{
|
|
212
248
|
"id": "axstack-arena-judge-astra",
|
|
213
249
|
"name": "Axstack arena judge Astra",
|
|
214
250
|
"provider": "codex",
|
|
215
|
-
"model": "gpt-6-astra",
|
|
216
251
|
"modeId": "full-access",
|
|
217
252
|
"thinkingOptionId": "xhigh",
|
|
218
|
-
"notes": "Arena judge (Astra seat). Read-only cross-judge for an arena-grade Align decision: receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads the other judge, never mutates. Disagreement between judges is surfaced by the driver, never averaged."
|
|
253
|
+
"notes": "Arena judge (Astra seat). Read-only cross-judge for an arena-grade Align decision: receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads the other judge, never mutates. Disagreement between judges is surfaced by the driver, never averaged.",
|
|
254
|
+
"modelClass": "astra"
|
|
219
255
|
},
|
|
220
256
|
{
|
|
221
257
|
"id": "axstack-escalation-fable",
|
|
@@ -5,55 +5,64 @@
|
|
|
5
5
|
"id": "axstack-advisor-astra",
|
|
6
6
|
"name": "Axstack Astra adviser",
|
|
7
7
|
"provider": "codex",
|
|
8
|
-
"model": "gpt-6-astra",
|
|
9
8
|
"modeId": "full-access",
|
|
10
9
|
"thinkingOptionId": "high",
|
|
11
|
-
"notes": "Independent Astra adviser for Align, Spec, and unresolved consequential decisions. Receives the same bounded evidence and question as Opus; reuse only unchanged receipts."
|
|
10
|
+
"notes": "Independent Astra adviser for Align, Spec, and unresolved consequential decisions. Receives the same bounded evidence and question as Opus; reuse only unchanged receipts.",
|
|
11
|
+
"modelClass": "astra"
|
|
12
12
|
},
|
|
13
13
|
{
|
|
14
14
|
"id": "axstack-advisor-opus",
|
|
15
15
|
"name": "Axstack Opus adviser",
|
|
16
16
|
"provider": "claude",
|
|
17
|
-
"model": "claude-opus-5-5",
|
|
18
17
|
"modeId": "bypassPermissions",
|
|
19
18
|
"thinkingOptionId": "xhigh",
|
|
20
|
-
"notes": "Independent Opus adviser for Align, Spec, and debug L1; authors the Claude arena candidate. Same bounded evidence and question as Astra; reuse only unchanged receipts."
|
|
19
|
+
"notes": "Independent Opus adviser for Align, Spec, and debug L1; authors the Claude arena candidate. Same bounded evidence and question as Astra; reuse only unchanged receipts. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
20
|
+
"modelClass": "opus"
|
|
21
21
|
},
|
|
22
22
|
{
|
|
23
23
|
"id": "axstack-owner",
|
|
24
24
|
"name": "Axstack PR owner",
|
|
25
25
|
"provider": "claude",
|
|
26
|
-
"model": "claude-opus-5-5",
|
|
27
26
|
"modeId": "bypassPermissions",
|
|
28
27
|
"thinkingOptionId": "medium",
|
|
29
|
-
"notes": "Persistent PR owner: one owner per PR, accountable for candidate, fixes, verification evidence, and monitoring. May delegate coding but never edits a worker-owned candidate concurrently. Launches eligible independent reviewers."
|
|
28
|
+
"notes": "Persistent PR owner: one owner per PR, accountable for candidate, fixes, verification evidence, and monitoring. May delegate coding but never edits a worker-owned candidate concurrently. Launches eligible independent reviewers. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
29
|
+
"modelClass": "opus"
|
|
30
30
|
},
|
|
31
31
|
{
|
|
32
32
|
"id": "axstack-author",
|
|
33
33
|
"name": "Axstack author",
|
|
34
34
|
"provider": "codex",
|
|
35
|
-
"model": "gpt-6-sol",
|
|
36
35
|
"modeId": "full-access",
|
|
37
36
|
"thinkingOptionId": "high",
|
|
38
|
-
"notes": "Ordinary implementation and repairs. Uses strict red-green-refactor and remains the exclusive writer for a candidate.
|
|
37
|
+
"notes": "Ordinary implementation and repairs. Uses strict red-green-refactor and remains the exclusive writer for a candidate. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
38
|
+
"modelClass": "sol"
|
|
39
39
|
},
|
|
40
40
|
{
|
|
41
41
|
"id": "axstack-reviewer-primary",
|
|
42
42
|
"name": "Axstack reviewer (primary)",
|
|
43
43
|
"provider": "codex",
|
|
44
|
-
"model": "gpt-6-sol",
|
|
45
44
|
"modeId": "full-access",
|
|
46
45
|
"thinkingOptionId": "high",
|
|
47
|
-
"notes": "Primary reviewer in the ordered mixed peer pair: Sol high followed by Opus medium. Eligible authored reviewer for an Opus-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated."
|
|
46
|
+
"notes": "Primary reviewer in the ordered mixed peer pair: Sol high followed by Opus medium. Eligible authored reviewer for an Opus-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated.",
|
|
47
|
+
"modelClass": "sol"
|
|
48
48
|
},
|
|
49
49
|
{
|
|
50
50
|
"id": "axstack-reviewer-secondary",
|
|
51
51
|
"name": "Axstack reviewer (secondary)",
|
|
52
52
|
"provider": "claude",
|
|
53
|
-
"model": "claude-opus-5-5",
|
|
54
53
|
"modeId": "bypassPermissions",
|
|
55
54
|
"thinkingOptionId": "medium",
|
|
56
|
-
"notes": "Secondary reviewer in the ordered mixed peer pair: Sol high followed by Opus medium. Eligible authored reviewer for a Sol-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated."
|
|
55
|
+
"notes": "Secondary reviewer in the ordered mixed peer pair: Sol high followed by Opus medium. Eligible authored reviewer for a Sol-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
56
|
+
"modelClass": "opus"
|
|
57
|
+
},
|
|
58
|
+
{
|
|
59
|
+
"id": "axstack-diligence",
|
|
60
|
+
"name": "Axstack diligence checker",
|
|
61
|
+
"provider": "claude",
|
|
62
|
+
"modeId": "bypassPermissions",
|
|
63
|
+
"thinkingOptionId": "high",
|
|
64
|
+
"notes": "Read-only diligence for exact-revision PRs and bounded research, spec, ticket, receipt, and release claims. Returns PASS or FINDINGS with evidence; never authors or edits. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
65
|
+
"modelClass": "sonnet"
|
|
57
66
|
},
|
|
58
67
|
{
|
|
59
68
|
"id": "axstack-checker",
|
|
@@ -68,28 +77,37 @@
|
|
|
68
77
|
"id": "axstack-research-requirements",
|
|
69
78
|
"name": "Axstack research requirements",
|
|
70
79
|
"provider": "claude",
|
|
71
|
-
"model": "claude-sonnet-5-5",
|
|
72
80
|
"modeId": "bypassPermissions",
|
|
73
81
|
"thinkingOptionId": "high",
|
|
74
|
-
"notes": "Sonnet high research requirements analyst: scopes bounded questions and acceptance for a research task.
|
|
82
|
+
"notes": "Sonnet high research requirements analyst: scopes bounded questions and acceptance for a research task. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
83
|
+
"modelClass": "sonnet"
|
|
75
84
|
},
|
|
76
85
|
{
|
|
77
86
|
"id": "axstack-research-code",
|
|
78
87
|
"name": "Axstack research code",
|
|
88
|
+
"provider": "claude",
|
|
89
|
+
"modeId": "bypassPermissions",
|
|
90
|
+
"thinkingOptionId": "high",
|
|
91
|
+
"notes": "Sonnet high research code investigator: verifies behavior against inspected code and executable evidence. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
92
|
+
"modelClass": "sonnet"
|
|
93
|
+
},
|
|
94
|
+
{
|
|
95
|
+
"id": "axstack-research-code-sol",
|
|
96
|
+
"name": "Axstack research code Sol",
|
|
79
97
|
"provider": "codex",
|
|
80
|
-
"model": "gpt-6-sol",
|
|
81
98
|
"modeId": "full-access",
|
|
82
99
|
"thinkingOptionId": "high",
|
|
83
|
-
"notes": "
|
|
100
|
+
"notes": "Independent Sol high research code investigator: verifies code behavior and executable evidence on the same bounded brief without cross-reading. Reports findings for driver reconciliation.",
|
|
101
|
+
"modelClass": "sol"
|
|
84
102
|
},
|
|
85
103
|
{
|
|
86
104
|
"id": "axstack-research-web",
|
|
87
105
|
"name": "Axstack research web",
|
|
88
106
|
"provider": "claude",
|
|
89
|
-
"model": "claude-sonnet-5-5",
|
|
90
107
|
"modeId": "bypassPermissions",
|
|
91
108
|
"thinkingOptionId": "high",
|
|
92
|
-
"notes": "Sonnet high research web reader: gathers primary-source facts efficiently.
|
|
109
|
+
"notes": "Sonnet high research web reader: gathers primary-source facts efficiently. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
110
|
+
"modelClass": "sonnet"
|
|
93
111
|
},
|
|
94
112
|
{
|
|
95
113
|
"id": "axstack-research-web-google",
|
|
@@ -113,127 +131,145 @@
|
|
|
113
131
|
"id": "axstack-explainer",
|
|
114
132
|
"name": "Axstack explainer",
|
|
115
133
|
"provider": "claude",
|
|
116
|
-
"model": "claude-sonnet-5-5",
|
|
117
134
|
"modeId": "bypassPermissions",
|
|
118
135
|
"thinkingOptionId": "high",
|
|
119
|
-
"notes": "Complex visual explanation author: traces systems, changes, and implementation gaps in requested artifacts and verifies rendered behavior where applicable.
|
|
136
|
+
"notes": "Complex visual explanation author: traces systems, changes, and implementation gaps in requested artifacts and verifies rendered behavior where applicable. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
137
|
+
"modelClass": "sonnet"
|
|
120
138
|
},
|
|
121
139
|
{
|
|
122
140
|
"id": "axstack-explainer-review",
|
|
123
141
|
"name": "Axstack explainer reviewer",
|
|
124
142
|
"provider": "codex",
|
|
125
|
-
"model": "gpt-6-luna",
|
|
126
143
|
"modeId": "full-access",
|
|
127
144
|
"thinkingOptionId": "xhigh",
|
|
128
|
-
"notes": "Independent visual explanation reviewer: checks the exact artifact for text and source fidelity. The rendered pass belongs to axstack-ui-verifier. Any artifact change invalidates its review.
|
|
145
|
+
"notes": "Independent visual explanation reviewer: checks the exact artifact for text and source fidelity. The rendered pass belongs to axstack-ui-verifier. Any artifact change invalidates its review. Resolve the class at run start; hold unavailable models except explicit pre-turn Codex rejection within class.",
|
|
146
|
+
"modelClass": "luna"
|
|
129
147
|
},
|
|
130
148
|
{
|
|
131
149
|
"id": "axstack-ui-verifier",
|
|
132
150
|
"name": "Axstack UI verifier",
|
|
133
151
|
"provider": "claude",
|
|
134
|
-
"model": "claude-sonnet-5-5",
|
|
135
152
|
"modeId": "bypassPermissions",
|
|
136
153
|
"thinkingOptionId": "high",
|
|
137
|
-
"notes": "Read-only UI verifier: runs Playwright or the browser against the given build, URL, or artifact; captures screenshots, interactions, accessibility, desktop/mobile, and reduced-motion evidence in the dispatch evidence folder; returns a verdict with evidence paths. Never edits source.
|
|
154
|
+
"notes": "Read-only UI verifier: runs Playwright or the browser against the given build, URL, or artifact; captures screenshots, interactions, accessibility, desktop/mobile, and reduced-motion evidence in the dispatch evidence folder; returns a verdict with evidence paths. Never edits source. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
155
|
+
"modelClass": "sonnet"
|
|
138
156
|
},
|
|
139
157
|
{
|
|
140
158
|
"id": "axstack-explore-codebase",
|
|
141
159
|
"name": "Axstack codebase explorer",
|
|
142
160
|
"provider": "claude",
|
|
143
|
-
"model": "claude-sonnet-5-5",
|
|
144
161
|
"modeId": "bypassPermissions",
|
|
145
162
|
"thinkingOptionId": "high",
|
|
146
|
-
"notes": "Codebase mapper: explores repository structure and interfaces for research and handoff context.
|
|
163
|
+
"notes": "Codebase mapper: explores repository structure and interfaces for research and handoff context. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
164
|
+
"modelClass": "sonnet"
|
|
147
165
|
},
|
|
148
166
|
{
|
|
149
167
|
"id": "axstack-explore-execution",
|
|
150
168
|
"name": "Axstack execution explorer",
|
|
169
|
+
"provider": "claude",
|
|
170
|
+
"modeId": "bypassPermissions",
|
|
171
|
+
"thinkingOptionId": "high",
|
|
172
|
+
"notes": "Sonnet high execution explorer: runs bounded checks of runtime behavior where authorized. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
173
|
+
"modelClass": "sonnet"
|
|
174
|
+
},
|
|
175
|
+
{
|
|
176
|
+
"id": "axstack-explore-execution-sol",
|
|
177
|
+
"name": "Axstack execution explorer Sol",
|
|
151
178
|
"provider": "codex",
|
|
152
|
-
"model": "gpt-6-sol",
|
|
153
179
|
"modeId": "full-access",
|
|
154
180
|
"thinkingOptionId": "high",
|
|
155
|
-
"notes": "
|
|
181
|
+
"notes": "Independent Sol high execution explorer: runs authorized bounded runtime checks on the same brief without cross-reading. Reports findings for driver reconciliation.",
|
|
182
|
+
"modelClass": "sol"
|
|
156
183
|
},
|
|
157
184
|
{
|
|
158
185
|
"id": "axstack-monitor",
|
|
159
186
|
"name": "Axstack monitor",
|
|
160
187
|
"provider": "claude",
|
|
161
|
-
"model": "claude-sonnet-5-5",
|
|
162
188
|
"modeId": "bypassPermissions",
|
|
163
189
|
"thinkingOptionId": "high",
|
|
164
|
-
"notes": "Optional Sonnet high independent read-only observer for a standalone PR watch. Reads GitHub, feedback, and checks, persists event IDs, and wakes the owner only for a new actionable event. Never sends, authors, reviews, replies, or acts as either reusable PR manager. Healthy snapshots stay quiet. Chat-run mode: one same-host native read-only observer per Run; fresh finite passes report precise deltas internally to the original Run/driver and may disable/read back only their own automation at verified stop. No repair, dispatch, public notification, or replacement coordinator. Effective scheduled model/effort and wake require live proof."
|
|
190
|
+
"notes": "Optional Sonnet high independent read-only observer for a standalone PR watch. Reads GitHub, feedback, and checks, persists event IDs, and wakes the owner only for a new actionable event. Never sends, authors, reviews, replies, or acts as either reusable PR manager. Healthy snapshots stay quiet. Chat-run mode: one same-host native read-only observer per Run; fresh finite passes report precise deltas internally to the original Run/driver and may disable/read back only their own automation at verified stop. No repair, dispatch, public notification, or replacement coordinator. Effective scheduled model/effort and wake require live proof. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
191
|
+
"modelClass": "sonnet"
|
|
165
192
|
},
|
|
166
193
|
{
|
|
167
194
|
"id": "axstack-auditor",
|
|
168
195
|
"name": "Axstack auditor",
|
|
196
|
+
"provider": "claude",
|
|
197
|
+
"modeId": "bypassPermissions",
|
|
198
|
+
"thinkingOptionId": "high",
|
|
199
|
+
"notes": "Sonnet high read-only end-of-run and checkpoint auditor. Collects scope and outcome evidence with counts and denominators and reports PASS, FAIL, or UNKNOWN without inventing numbers. Never edits, merges, activates, or audits itself. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
200
|
+
"modelClass": "sonnet"
|
|
201
|
+
},
|
|
202
|
+
{
|
|
203
|
+
"id": "axstack-auditor-sol",
|
|
204
|
+
"name": "Axstack auditor Sol",
|
|
169
205
|
"provider": "codex",
|
|
170
|
-
"model": "gpt-6-luna",
|
|
171
206
|
"modeId": "full-access",
|
|
172
|
-
"thinkingOptionId": "
|
|
173
|
-
"notes": "
|
|
207
|
+
"thinkingOptionId": "high",
|
|
208
|
+
"notes": "Independent Sol high read-only auditor: checks the same bounded run evidence without cross-reading. Reports findings for driver reconciliation; never edits, merges, or activates.",
|
|
209
|
+
"modelClass": "sol"
|
|
174
210
|
},
|
|
175
211
|
{
|
|
176
212
|
"id": "axstack-debug-investigator-1",
|
|
177
213
|
"name": "Axstack debug investigator 1",
|
|
178
214
|
"provider": "claude",
|
|
179
|
-
"model": "claude-opus-5-5",
|
|
180
215
|
"modeId": "bypassPermissions",
|
|
181
216
|
"thinkingOptionId": "medium",
|
|
182
|
-
"notes": "Debug investigator seat 1. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity."
|
|
217
|
+
"notes": "Debug investigator seat 1. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
218
|
+
"modelClass": "opus"
|
|
183
219
|
},
|
|
184
220
|
{
|
|
185
221
|
"id": "axstack-debug-investigator-2",
|
|
186
222
|
"name": "Axstack debug investigator 2",
|
|
187
223
|
"provider": "codex",
|
|
188
|
-
"model": "gpt-6-sol",
|
|
189
224
|
"modeId": "full-access",
|
|
190
225
|
"thinkingOptionId": "high",
|
|
191
|
-
"notes": "Debug investigator seat 2. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity."
|
|
226
|
+
"notes": "Debug investigator seat 2. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity.",
|
|
227
|
+
"modelClass": "sol"
|
|
192
228
|
},
|
|
193
229
|
{
|
|
194
230
|
"id": "axstack-debug-investigator-3",
|
|
195
231
|
"name": "Axstack debug investigator 3",
|
|
196
232
|
"provider": "claude",
|
|
197
|
-
"model": "claude-sonnet-5-5",
|
|
198
233
|
"modeId": "bypassPermissions",
|
|
199
234
|
"thinkingOptionId": "high",
|
|
200
|
-
"notes": "Debug investigator seat 3. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity."
|
|
235
|
+
"notes": "Debug investigator seat 3. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
236
|
+
"modelClass": "sonnet"
|
|
201
237
|
},
|
|
202
238
|
{
|
|
203
239
|
"id": "axstack-debug-investigator-4",
|
|
204
240
|
"name": "Axstack debug investigator 4",
|
|
205
241
|
"provider": "codex",
|
|
206
|
-
"model": "gpt-6-sol",
|
|
207
242
|
"modeId": "full-access",
|
|
208
243
|
"thinkingOptionId": "high",
|
|
209
|
-
"notes": "Debug investigator seat 4. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity."
|
|
244
|
+
"notes": "Debug investigator seat 4. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity.",
|
|
245
|
+
"modelClass": "sol"
|
|
210
246
|
},
|
|
211
247
|
{
|
|
212
248
|
"id": "axstack-arena-judge-astra",
|
|
213
249
|
"name": "Axstack arena judge Astra",
|
|
214
250
|
"provider": "codex",
|
|
215
|
-
"model": "gpt-6-astra",
|
|
216
251
|
"modeId": "full-access",
|
|
217
252
|
"thinkingOptionId": "xhigh",
|
|
218
|
-
"notes": "Arena judge (Astra seat). Read-only cross-judge for an arena-grade Align decision: receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads the other judge, never mutates. Disagreement between judges is surfaced by the driver, never averaged."
|
|
253
|
+
"notes": "Arena judge (Astra seat). Read-only cross-judge for an arena-grade Align decision: receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads the other judge, never mutates. Disagreement between judges is surfaced by the driver, never averaged.",
|
|
254
|
+
"modelClass": "astra"
|
|
219
255
|
},
|
|
220
256
|
{
|
|
221
257
|
"id": "axstack-escalation-fable",
|
|
222
258
|
"name": "Axstack escalation Fable",
|
|
223
259
|
"provider": "claude",
|
|
224
|
-
"model": "claude-fable-5-1",
|
|
225
260
|
"modeId": "bypassPermissions",
|
|
226
261
|
"thinkingOptionId": "xhigh",
|
|
227
|
-
"notes": "Fable escalation seat for arena round 2, high-stakes plain AGREE, or the bounded escalation trigger. Fresh session per use; never reuses adviser or candidate context. Read-only round 2 judge: scores every candidate by label; never authors a candidate."
|
|
262
|
+
"notes": "Fable escalation seat for arena round 2, high-stakes plain AGREE, or the bounded escalation trigger. Fresh session per use; never reuses adviser or candidate context. Read-only round 2 judge: scores every candidate by label; never authors a candidate. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
263
|
+
"modelClass": "fable"
|
|
228
264
|
},
|
|
229
265
|
{
|
|
230
266
|
"id": "axstack-arena-judge-opus",
|
|
231
267
|
"name": "Axstack arena judge Opus",
|
|
232
268
|
"provider": "claude",
|
|
233
|
-
"model": "claude-opus-5-5",
|
|
234
269
|
"modeId": "bypassPermissions",
|
|
235
270
|
"thinkingOptionId": "xhigh",
|
|
236
|
-
"notes": "Read-only round 1 arena judge. Receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads another judge, never mutates; the driver compares its verdict without averaging."
|
|
271
|
+
"notes": "Read-only round 1 arena judge. Receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads another judge, never mutates; the driver compares its verdict without averaging. Claude alias resolves at first launch; a Claude rejection holds.",
|
|
272
|
+
"modelClass": "opus"
|
|
237
273
|
},
|
|
238
274
|
{
|
|
239
275
|
"id": "axstack-arena-candidate-grok",
|