axstack 0.20.30 → 0.21.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +24 -23
- package/bin/axstack.js +18 -5
- package/docs/installation.md +101 -46
- package/docs/workflows.md +179 -117
- package/package.json +3 -3
- package/profiles/presets/claude-only.json +46 -46
- package/profiles/presets/codex-only.json +50 -50
- package/profiles/presets/mixed.json +59 -59
- package/skills/axstack/references/automations.md +127 -137
- package/skills/axstack/references/autopilot.md +121 -0
- package/skills/axstack/references/candidate-publication.md +13 -8
- package/skills/axstack/references/contracts.md +10 -4
- package/skills/axstack/references/diligence.md +3 -1
- package/skills/axstack/references/evidence-archive.md +38 -33
- package/skills/axstack/references/lifecycle.md +64 -50
- package/skills/axstack/references/review-manager-prompt.md +13 -11
- package/skills/axstack/references/role-roster.md +19 -9
- package/skills/axstack/references/routing.md +33 -18
- package/skills/axstack/references/run-record.md +36 -15
- package/skills/axstack/references/t3-runtime.md +234 -0
- package/skills/axstack/references/test-audit-weekly.md +62 -0
- package/skills/axstack/references/test-value.md +120 -0
- package/skills/axstack/references/ui-verification.md +5 -1
- package/skills/axstack/references/workspace-hygiene.md +102 -156
- package/skills/axstack/scripts/pr-digest.js +120 -0
- package/skills/axstack/scripts/resolve-models.js +102 -0
- package/skills/axstack-align/SKILL.md +25 -11
- package/skills/axstack-audit/SKILL.md +22 -5
- package/skills/axstack-audit/references/record.md +1 -1
- package/skills/axstack-cleanup/SKILL.md +69 -87
- package/skills/axstack-debug/SKILL.md +1 -1
- package/skills/axstack-explain/SKILL.md +1 -1
- package/skills/axstack-explain/references/visual-qa.md +2 -0
- package/skills/axstack-implement/SKILL.md +76 -26
- package/skills/axstack-improve/SKILL.md +24 -4
- package/skills/axstack-relay/SKILL.md +16 -7
- package/skills/axstack-research/SKILL.md +11 -4
- package/skills/axstack-review/SKILL.md +42 -32
- package/skills/axstack-spec/SKILL.md +23 -14
- package/skills/axstack-tickets/SKILL.md +13 -11
- package/skills/axstack-watch/SKILL.md +117 -34
- package/skills/axstack-watch/references/watch-runtime.md +61 -69
- package/src/capabilities.js +33 -69
- package/src/installer.js +9 -1
- package/src/instructions.js +9 -4
- package/src/roles.js +38 -10
- package/skills/axstack/references/orca-runtime.md +0 -183
- package/skills/axstack/scripts/trust-path.js +0 -123
|
@@ -5,64 +5,64 @@
|
|
|
5
5
|
"id": "axstack-advisor-astra",
|
|
6
6
|
"name": "Axstack Astra adviser",
|
|
7
7
|
"provider": "codex",
|
|
8
|
-
"model": "gpt-6-astra",
|
|
9
8
|
"modeId": "full-access",
|
|
10
9
|
"thinkingOptionId": "high",
|
|
11
|
-
"notes": "Independent Astra adviser for Align, Spec, and unresolved consequential decisions. Receives the same bounded evidence and question as Opus; reuse only unchanged receipts."
|
|
10
|
+
"notes": "Independent Astra adviser for Align, Spec, and unresolved consequential decisions. Receives the same bounded evidence and question as Opus; reuse only unchanged receipts.",
|
|
11
|
+
"modelClass": "astra"
|
|
12
12
|
},
|
|
13
13
|
{
|
|
14
14
|
"id": "axstack-advisor-opus",
|
|
15
15
|
"name": "Axstack Opus adviser",
|
|
16
16
|
"provider": "claude",
|
|
17
|
-
"model": "claude-opus-5-5",
|
|
18
17
|
"modeId": "bypassPermissions",
|
|
19
18
|
"thinkingOptionId": "xhigh",
|
|
20
|
-
"notes": "Independent Opus adviser for Align, Spec, and debug L1; authors the Claude arena candidate. Same bounded evidence and question as Astra; reuse only unchanged receipts."
|
|
19
|
+
"notes": "Independent Opus adviser for Align, Spec, and debug L1; authors the Claude arena candidate. Same bounded evidence and question as Astra; reuse only unchanged receipts. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
20
|
+
"modelClass": "opus"
|
|
21
21
|
},
|
|
22
22
|
{
|
|
23
23
|
"id": "axstack-owner",
|
|
24
24
|
"name": "Axstack PR owner",
|
|
25
25
|
"provider": "claude",
|
|
26
|
-
"model": "claude-opus-5-5",
|
|
27
26
|
"modeId": "bypassPermissions",
|
|
28
27
|
"thinkingOptionId": "medium",
|
|
29
|
-
"notes": "Persistent PR owner: one owner per PR, accountable for candidate, fixes, verification evidence, and monitoring. May delegate coding but never edits a worker-owned candidate concurrently. Launches eligible independent reviewers."
|
|
28
|
+
"notes": "Persistent PR owner: one owner per PR, accountable for candidate, fixes, verification evidence, and monitoring. May delegate coding but never edits a worker-owned candidate concurrently. Launches eligible independent reviewers. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
29
|
+
"modelClass": "opus"
|
|
30
30
|
},
|
|
31
31
|
{
|
|
32
32
|
"id": "axstack-author",
|
|
33
33
|
"name": "Axstack author",
|
|
34
34
|
"provider": "codex",
|
|
35
|
-
"model": "gpt-6-sol",
|
|
36
35
|
"modeId": "full-access",
|
|
37
36
|
"thinkingOptionId": "high",
|
|
38
|
-
"notes": "Ordinary implementation and repairs. Uses strict red-green-refactor and remains the exclusive writer for a candidate.
|
|
37
|
+
"notes": "Ordinary implementation and repairs. Uses strict red-green-refactor and remains the exclusive writer for a candidate. Resolve the class from saved T3 capabilities at run start; rejection and unavailable settings hold without substitution.",
|
|
38
|
+
"modelClass": "sol"
|
|
39
39
|
},
|
|
40
40
|
{
|
|
41
41
|
"id": "axstack-reviewer-primary",
|
|
42
42
|
"name": "Axstack reviewer (primary)",
|
|
43
43
|
"provider": "codex",
|
|
44
|
-
"model": "gpt-6-sol",
|
|
45
44
|
"modeId": "full-access",
|
|
46
45
|
"thinkingOptionId": "high",
|
|
47
|
-
"notes": "Primary reviewer in the ordered mixed peer pair: Sol high followed by Opus medium. Eligible authored reviewer for an Opus-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated."
|
|
46
|
+
"notes": "Primary reviewer in the ordered mixed peer pair: Sol high followed by Opus medium. Eligible authored reviewer for an Opus-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated.",
|
|
47
|
+
"modelClass": "sol"
|
|
48
48
|
},
|
|
49
49
|
{
|
|
50
50
|
"id": "axstack-reviewer-secondary",
|
|
51
51
|
"name": "Axstack reviewer (secondary)",
|
|
52
52
|
"provider": "claude",
|
|
53
|
-
"model": "claude-opus-5-5",
|
|
54
53
|
"modeId": "bypassPermissions",
|
|
55
54
|
"thinkingOptionId": "medium",
|
|
56
|
-
"notes": "Secondary reviewer in the ordered mixed peer pair: Sol high followed by Opus medium. Eligible authored reviewer for a Sol-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated."
|
|
55
|
+
"notes": "Secondary reviewer in the ordered mixed peer pair: Sol high followed by Opus medium. Eligible authored reviewer for a Sol-authored candidate. Never author or owner; review exact SHA and base across all six angles and acceptance. Peer first pass stays isolated. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
56
|
+
"modelClass": "opus"
|
|
57
57
|
},
|
|
58
58
|
{
|
|
59
59
|
"id": "axstack-diligence",
|
|
60
60
|
"name": "Axstack diligence checker",
|
|
61
61
|
"provider": "claude",
|
|
62
|
-
"model": "claude-sonnet-5-5",
|
|
63
62
|
"modeId": "bypassPermissions",
|
|
64
63
|
"thinkingOptionId": "high",
|
|
65
|
-
"notes": "Read-only diligence for exact-revision PRs and bounded research, spec, ticket, receipt, and release claims. Returns PASS or FINDINGS with evidence; never authors or edits."
|
|
64
|
+
"notes": "Read-only diligence for exact-revision PRs and bounded research, spec, ticket, receipt, and release claims. Returns PASS or FINDINGS with evidence; never authors or edits. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
65
|
+
"modelClass": "sonnet"
|
|
66
66
|
},
|
|
67
67
|
{
|
|
68
68
|
"id": "axstack-checker",
|
|
@@ -71,43 +71,43 @@
|
|
|
71
71
|
"model": null,
|
|
72
72
|
"modeId": "full-access",
|
|
73
73
|
"thinkingOptionId": "low",
|
|
74
|
-
"notes": "Report-only discrepancy checker for the selected external tracker (Linear or GitHub Issues). Never mutates the tracker; the driver independently verifies evidence before applying updates. Cheap report-only seat on Gemini Flash;
|
|
74
|
+
"notes": "Report-only discrepancy checker for the selected external tracker (Linear or GitHub Issues). Never mutates the tracker; the driver independently verifies evidence before applying updates. Cheap report-only seat on Gemini Flash; dispatch through the T3 Antigravity ACP provider."
|
|
75
75
|
},
|
|
76
76
|
{
|
|
77
77
|
"id": "axstack-research-requirements",
|
|
78
78
|
"name": "Axstack research requirements",
|
|
79
79
|
"provider": "claude",
|
|
80
|
-
"model": "claude-sonnet-5-5",
|
|
81
80
|
"modeId": "bypassPermissions",
|
|
82
81
|
"thinkingOptionId": "high",
|
|
83
|
-
"notes": "Sonnet high research requirements analyst: scopes bounded questions and acceptance for a research task.
|
|
82
|
+
"notes": "Sonnet high research requirements analyst: scopes bounded questions and acceptance for a research task. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
83
|
+
"modelClass": "sonnet"
|
|
84
84
|
},
|
|
85
85
|
{
|
|
86
86
|
"id": "axstack-research-code",
|
|
87
87
|
"name": "Axstack research code",
|
|
88
88
|
"provider": "claude",
|
|
89
|
-
"model": "claude-sonnet-5-5",
|
|
90
89
|
"modeId": "bypassPermissions",
|
|
91
90
|
"thinkingOptionId": "high",
|
|
92
|
-
"notes": "Sonnet high research code investigator: verifies behavior against inspected code and executable evidence.
|
|
91
|
+
"notes": "Sonnet high research code investigator: verifies behavior against inspected code and executable evidence. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
92
|
+
"modelClass": "sonnet"
|
|
93
93
|
},
|
|
94
94
|
{
|
|
95
95
|
"id": "axstack-research-code-sol",
|
|
96
96
|
"name": "Axstack research code Sol",
|
|
97
97
|
"provider": "codex",
|
|
98
|
-
"model": "gpt-6-sol",
|
|
99
98
|
"modeId": "full-access",
|
|
100
99
|
"thinkingOptionId": "high",
|
|
101
|
-
"notes": "Independent Sol high research code investigator: verifies code behavior and executable evidence on the same bounded brief without cross-reading. Reports findings for driver reconciliation."
|
|
100
|
+
"notes": "Independent Sol high research code investigator: verifies code behavior and executable evidence on the same bounded brief without cross-reading. Reports findings for driver reconciliation.",
|
|
101
|
+
"modelClass": "sol"
|
|
102
102
|
},
|
|
103
103
|
{
|
|
104
104
|
"id": "axstack-research-web",
|
|
105
105
|
"name": "Axstack research web",
|
|
106
106
|
"provider": "claude",
|
|
107
|
-
"model": "claude-sonnet-5-5",
|
|
108
107
|
"modeId": "bypassPermissions",
|
|
109
108
|
"thinkingOptionId": "high",
|
|
110
|
-
"notes": "Sonnet high research web reader: gathers primary-source facts efficiently.
|
|
109
|
+
"notes": "Sonnet high research web reader: gathers primary-source facts efficiently. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
110
|
+
"modelClass": "sonnet"
|
|
111
111
|
},
|
|
112
112
|
{
|
|
113
113
|
"id": "axstack-research-web-google",
|
|
@@ -116,7 +116,7 @@
|
|
|
116
116
|
"model": null,
|
|
117
117
|
"modeId": "full-access",
|
|
118
118
|
"thinkingOptionId": "high",
|
|
119
|
-
"notes": "Google-Search-grounded web research branch on Gemini via the Antigravity
|
|
119
|
+
"notes": "Google-Search-grounded web research branch on Gemini via the T3 Antigravity ACP provider; resolve the model from saved T3 capabilities; cites URL + access date per claim and re-opens sources because the provider's search citation format is undocumented; report-only."
|
|
120
120
|
},
|
|
121
121
|
{
|
|
122
122
|
"id": "axstack-research-x",
|
|
@@ -125,151 +125,151 @@
|
|
|
125
125
|
"model": null,
|
|
126
126
|
"modeId": "full-access",
|
|
127
127
|
"thinkingOptionId": "high",
|
|
128
|
-
"notes": "X/Twitter-only research branch: Grok can read X; use for questions where posts, threads, announcements, or sentiment on X are answer-changing evidence. Cites post URLs and dates; never the sole source for a verified claim; report-only, no writes.
|
|
128
|
+
"notes": "X/Twitter-only research branch: Grok can read X; use for questions where posts, threads, announcements, or sentiment on X are answer-changing evidence. Cites post URLs and dates; never the sole source for a verified claim; report-only, no writes. Use the T3 Grok provider and resolve the model from saved capabilities."
|
|
129
129
|
},
|
|
130
130
|
{
|
|
131
131
|
"id": "axstack-explainer",
|
|
132
132
|
"name": "Axstack explainer",
|
|
133
133
|
"provider": "claude",
|
|
134
|
-
"model": "claude-sonnet-5-5",
|
|
135
134
|
"modeId": "bypassPermissions",
|
|
136
135
|
"thinkingOptionId": "high",
|
|
137
|
-
"notes": "Complex visual explanation author: traces systems, changes, and implementation gaps in requested artifacts and verifies rendered behavior where applicable.
|
|
136
|
+
"notes": "Complex visual explanation author: traces systems, changes, and implementation gaps in requested artifacts and verifies rendered behavior where applicable. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
137
|
+
"modelClass": "sonnet"
|
|
138
138
|
},
|
|
139
139
|
{
|
|
140
140
|
"id": "axstack-explainer-review",
|
|
141
141
|
"name": "Axstack explainer reviewer",
|
|
142
142
|
"provider": "codex",
|
|
143
|
-
"model": "gpt-6-luna",
|
|
144
143
|
"modeId": "full-access",
|
|
145
144
|
"thinkingOptionId": "xhigh",
|
|
146
|
-
"notes": "Independent visual explanation reviewer: checks the exact artifact for text and source fidelity. The rendered pass belongs to axstack-ui-verifier. Any artifact change invalidates its review.
|
|
145
|
+
"notes": "Independent visual explanation reviewer: checks the exact artifact for text and source fidelity. The rendered pass belongs to axstack-ui-verifier. Any artifact change invalidates its review. Resolve the class from saved T3 capabilities at run start; rejection and unavailable settings hold without substitution.",
|
|
146
|
+
"modelClass": "luna"
|
|
147
147
|
},
|
|
148
148
|
{
|
|
149
149
|
"id": "axstack-ui-verifier",
|
|
150
150
|
"name": "Axstack UI verifier",
|
|
151
151
|
"provider": "claude",
|
|
152
|
-
"model": "claude-sonnet-5-5",
|
|
153
152
|
"modeId": "bypassPermissions",
|
|
154
153
|
"thinkingOptionId": "high",
|
|
155
|
-
"notes": "Read-only UI verifier:
|
|
154
|
+
"notes": "Read-only UI verifier: uses T3 preview_* tools against the given build, URL, or artifact; captures screenshots, interactions, accessibility, desktop/mobile, and reduced-motion evidence in the dispatch evidence folder; returns a verdict with evidence paths. Never edits source. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
155
|
+
"modelClass": "sonnet"
|
|
156
156
|
},
|
|
157
157
|
{
|
|
158
158
|
"id": "axstack-explore-codebase",
|
|
159
159
|
"name": "Axstack codebase explorer",
|
|
160
160
|
"provider": "claude",
|
|
161
|
-
"model": "claude-sonnet-5-5",
|
|
162
161
|
"modeId": "bypassPermissions",
|
|
163
162
|
"thinkingOptionId": "high",
|
|
164
|
-
"notes": "Codebase mapper: explores repository structure and interfaces for research and handoff context.
|
|
163
|
+
"notes": "Codebase mapper: explores repository structure and interfaces for research and handoff context. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
164
|
+
"modelClass": "sonnet"
|
|
165
165
|
},
|
|
166
166
|
{
|
|
167
167
|
"id": "axstack-explore-execution",
|
|
168
168
|
"name": "Axstack execution explorer",
|
|
169
169
|
"provider": "claude",
|
|
170
|
-
"model": "claude-sonnet-5-5",
|
|
171
170
|
"modeId": "bypassPermissions",
|
|
172
171
|
"thinkingOptionId": "high",
|
|
173
|
-
"notes": "Sonnet high execution explorer: runs bounded checks of runtime behavior where authorized.
|
|
172
|
+
"notes": "Sonnet high execution explorer: runs bounded checks of runtime behavior where authorized. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
173
|
+
"modelClass": "sonnet"
|
|
174
174
|
},
|
|
175
175
|
{
|
|
176
176
|
"id": "axstack-explore-execution-sol",
|
|
177
177
|
"name": "Axstack execution explorer Sol",
|
|
178
178
|
"provider": "codex",
|
|
179
|
-
"model": "gpt-6-sol",
|
|
180
179
|
"modeId": "full-access",
|
|
181
180
|
"thinkingOptionId": "high",
|
|
182
|
-
"notes": "Independent Sol high execution explorer: runs authorized bounded runtime checks on the same brief without cross-reading. Reports findings for driver reconciliation."
|
|
181
|
+
"notes": "Independent Sol high execution explorer: runs authorized bounded runtime checks on the same brief without cross-reading. Reports findings for driver reconciliation.",
|
|
182
|
+
"modelClass": "sol"
|
|
183
183
|
},
|
|
184
184
|
{
|
|
185
185
|
"id": "axstack-monitor",
|
|
186
186
|
"name": "Axstack monitor",
|
|
187
187
|
"provider": "claude",
|
|
188
|
-
"model": "claude-sonnet-5-5",
|
|
189
188
|
"modeId": "bypassPermissions",
|
|
190
189
|
"thinkingOptionId": "high",
|
|
191
|
-
"notes": "Optional Sonnet high independent read-only observer for a standalone PR watch. Reads GitHub, feedback, and checks, persists event IDs, and wakes the owner only for a new actionable event. Never sends, authors, reviews, replies, or acts as either reusable PR manager. Healthy snapshots stay quiet. Chat-run mode: one same-host native read-only observer per Run; fresh finite passes report precise deltas internally to the original Run/driver and may disable/read back only their own automation at verified stop. No repair, dispatch, public notification, or replacement coordinator. Effective scheduled model/effort and wake require live proof."
|
|
190
|
+
"notes": "Optional Sonnet high independent read-only observer for a standalone PR watch. Reads GitHub, feedback, and checks, persists event IDs, and wakes the owner only for a new actionable event. Never sends, authors, reviews, replies, or acts as either reusable PR manager. Healthy snapshots stay quiet. Chat-run mode: one same-host native read-only observer per Run; fresh finite passes report precise deltas internally to the original Run/driver and may disable/read back only their own automation at verified stop. No repair, dispatch, public notification, or replacement coordinator. Effective scheduled model/effort and wake require live proof. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
191
|
+
"modelClass": "sonnet"
|
|
192
192
|
},
|
|
193
193
|
{
|
|
194
194
|
"id": "axstack-auditor",
|
|
195
195
|
"name": "Axstack auditor",
|
|
196
196
|
"provider": "claude",
|
|
197
|
-
"model": "claude-sonnet-5-5",
|
|
198
197
|
"modeId": "bypassPermissions",
|
|
199
198
|
"thinkingOptionId": "high",
|
|
200
|
-
"notes": "Sonnet high read-only end-of-run and checkpoint auditor. Collects scope and outcome evidence with counts and denominators and reports PASS, FAIL, or UNKNOWN without inventing numbers. Never edits, merges, activates, or audits itself."
|
|
199
|
+
"notes": "Sonnet high read-only end-of-run and checkpoint auditor. Collects scope and outcome evidence with counts and denominators and reports PASS, FAIL, or UNKNOWN without inventing numbers. Never edits, merges, activates, or audits itself. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
200
|
+
"modelClass": "sonnet"
|
|
201
201
|
},
|
|
202
202
|
{
|
|
203
203
|
"id": "axstack-auditor-sol",
|
|
204
204
|
"name": "Axstack auditor Sol",
|
|
205
205
|
"provider": "codex",
|
|
206
|
-
"model": "gpt-6-sol",
|
|
207
206
|
"modeId": "full-access",
|
|
208
207
|
"thinkingOptionId": "high",
|
|
209
|
-
"notes": "Independent Sol high read-only auditor: checks the same bounded run evidence without cross-reading. Reports findings for driver reconciliation; never edits, merges, or activates."
|
|
208
|
+
"notes": "Independent Sol high read-only auditor: checks the same bounded run evidence without cross-reading. Reports findings for driver reconciliation; never edits, merges, or activates.",
|
|
209
|
+
"modelClass": "sol"
|
|
210
210
|
},
|
|
211
211
|
{
|
|
212
212
|
"id": "axstack-debug-investigator-1",
|
|
213
213
|
"name": "Axstack debug investigator 1",
|
|
214
214
|
"provider": "claude",
|
|
215
|
-
"model": "claude-opus-5-5",
|
|
216
215
|
"modeId": "bypassPermissions",
|
|
217
216
|
"thinkingOptionId": "medium",
|
|
218
|
-
"notes": "Debug investigator seat 1. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity."
|
|
217
|
+
"notes": "Debug investigator seat 1. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
218
|
+
"modelClass": "opus"
|
|
219
219
|
},
|
|
220
220
|
{
|
|
221
221
|
"id": "axstack-debug-investigator-2",
|
|
222
222
|
"name": "Axstack debug investigator 2",
|
|
223
223
|
"provider": "codex",
|
|
224
|
-
"model": "gpt-6-sol",
|
|
225
224
|
"modeId": "full-access",
|
|
226
225
|
"thinkingOptionId": "high",
|
|
227
|
-
"notes": "Debug investigator seat 2. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity."
|
|
226
|
+
"notes": "Debug investigator seat 2. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity.",
|
|
227
|
+
"modelClass": "sol"
|
|
228
228
|
},
|
|
229
229
|
{
|
|
230
230
|
"id": "axstack-debug-investigator-3",
|
|
231
231
|
"name": "Axstack debug investigator 3",
|
|
232
232
|
"provider": "claude",
|
|
233
|
-
"model": "claude-sonnet-5-5",
|
|
234
233
|
"modeId": "bypassPermissions",
|
|
235
234
|
"thinkingOptionId": "high",
|
|
236
|
-
"notes": "Debug investigator seat 3. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity."
|
|
235
|
+
"notes": "Debug investigator seat 3. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
236
|
+
"modelClass": "sonnet"
|
|
237
237
|
},
|
|
238
238
|
{
|
|
239
239
|
"id": "axstack-debug-investigator-4",
|
|
240
240
|
"name": "Axstack debug investigator 4",
|
|
241
241
|
"provider": "codex",
|
|
242
|
-
"model": "gpt-6-sol",
|
|
243
242
|
"modeId": "full-access",
|
|
244
243
|
"thinkingOptionId": "high",
|
|
245
|
-
"notes": "Debug investigator seat 4. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity."
|
|
244
|
+
"notes": "Debug investigator seat 4. Dispatched only by axstack-debug at L1 with the shared evidence packet and one distinct brief; never reads another investigator's output. Works in its own disposable worktree at the pinned revision plus the recorded dirty patch; may instrument there for probes; never commits, pushes, publishes, or creates children. Returns one receipt per brief. Independence comes from brief isolation, not model diversity.",
|
|
245
|
+
"modelClass": "sol"
|
|
246
246
|
},
|
|
247
247
|
{
|
|
248
248
|
"id": "axstack-arena-judge-astra",
|
|
249
249
|
"name": "Axstack arena judge Astra",
|
|
250
250
|
"provider": "codex",
|
|
251
|
-
"model": "gpt-6-astra",
|
|
252
251
|
"modeId": "full-access",
|
|
253
252
|
"thinkingOptionId": "xhigh",
|
|
254
|
-
"notes": "Arena judge (Astra seat). Read-only cross-judge for an arena-grade Align decision: receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads the other judge, never mutates. Disagreement between judges is surfaced by the driver, never averaged."
|
|
253
|
+
"notes": "Arena judge (Astra seat). Read-only cross-judge for an arena-grade Align decision: receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads the other judge, never mutates. Disagreement between judges is surfaced by the driver, never averaged.",
|
|
254
|
+
"modelClass": "astra"
|
|
255
255
|
},
|
|
256
256
|
{
|
|
257
257
|
"id": "axstack-escalation-fable",
|
|
258
258
|
"name": "Axstack escalation Fable",
|
|
259
259
|
"provider": "claude",
|
|
260
|
-
"model": "claude-fable-5-1",
|
|
261
260
|
"modeId": "bypassPermissions",
|
|
262
261
|
"thinkingOptionId": "xhigh",
|
|
263
|
-
"notes": "Fable escalation seat for arena round 2, high-stakes plain AGREE, or the bounded escalation trigger. Fresh session per use; never reuses adviser or candidate context. Read-only round 2 judge: scores every candidate by label; never authors a candidate."
|
|
262
|
+
"notes": "Fable escalation seat for arena round 2, high-stakes plain AGREE, or the bounded escalation trigger. Fresh session per use; never reuses adviser or candidate context. Read-only round 2 judge: scores every candidate by label; never authors a candidate. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
263
|
+
"modelClass": "fable"
|
|
264
264
|
},
|
|
265
265
|
{
|
|
266
266
|
"id": "axstack-arena-judge-opus",
|
|
267
267
|
"name": "Axstack arena judge Opus",
|
|
268
268
|
"provider": "claude",
|
|
269
|
-
"model": "claude-opus-5-5",
|
|
270
269
|
"modeId": "bypassPermissions",
|
|
271
270
|
"thinkingOptionId": "xhigh",
|
|
272
|
-
"notes": "Read-only round 1 arena judge. Receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads another judge, never mutates; the driver compares its verdict without averaging."
|
|
271
|
+
"notes": "Read-only round 1 arena judge. Receives the rubric and every candidate by label, scores each criterion, and recommends a base with rationale. Never authors a candidate, never cross-reads another judge, never mutates; the driver compares its verdict without averaging. T3 resolves Claude classes from saved capabilities; a Claude rejection holds.",
|
|
272
|
+
"modelClass": "opus"
|
|
273
273
|
},
|
|
274
274
|
{
|
|
275
275
|
"id": "axstack-arena-candidate-grok",
|
|
@@ -278,7 +278,7 @@
|
|
|
278
278
|
"model": null,
|
|
279
279
|
"modeId": "full-access",
|
|
280
280
|
"thinkingOptionId": "high",
|
|
281
|
-
"notes": "Independent Grok Rung 2 design candidate.
|
|
281
|
+
"notes": "Independent Grok Rung 2 design candidate. Use the T3 Grok provider and resolve the model from saved capabilities. Receives the same brief without cross-reading; returns design, rationale, and rejected alternatives."
|
|
282
282
|
},
|
|
283
283
|
{
|
|
284
284
|
"id": "axstack-arena-candidate-antigravity",
|
|
@@ -287,7 +287,7 @@
|
|
|
287
287
|
"model": null,
|
|
288
288
|
"modeId": "full-access",
|
|
289
289
|
"thinkingOptionId": "high",
|
|
290
|
-
"notes": "Independent Antigravity Rung 2 design candidate.
|
|
290
|
+
"notes": "Independent Antigravity Rung 2 design candidate. Use the T3 Antigravity ACP provider and resolve the model from saved capabilities. Receives the same brief without cross-reading; returns design, rationale, and rejected alternatives."
|
|
291
291
|
}
|
|
292
292
|
]
|
|
293
293
|
}
|