dsh-aris-panel 0.1.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (442) hide show
  1. package/LICENSE +21 -0
  2. package/README.md +98 -0
  3. package/README_CN.md +87 -0
  4. package/dsh/checkout.patch.yml +38 -0
  5. package/dsh/client.js +634 -0
  6. package/dsh/cordis.patch.yml +44 -0
  7. package/dsh/index.mjs +76 -0
  8. package/dsh/run-status.mjs +182 -0
  9. package/dsh/scope-limits.mjs +50 -0
  10. package/dsh/workbench.mjs +291 -0
  11. package/mcp-servers/claude-review/README.md +93 -0
  12. package/mcp-servers/claude-review/run_with_claude_aws.sh +49 -0
  13. package/mcp-servers/claude-review/server.py +718 -0
  14. package/mcp-servers/codex-image2/README.md +65 -0
  15. package/mcp-servers/codex-image2/server.py +893 -0
  16. package/mcp-servers/feishu-bridge/requirements.txt +1 -0
  17. package/mcp-servers/feishu-bridge/server.py +240 -0
  18. package/mcp-servers/gemini-review/README.md +171 -0
  19. package/mcp-servers/gemini-review/server.py +1856 -0
  20. package/mcp-servers/llm-chat/requirements.txt +1 -0
  21. package/mcp-servers/llm-chat/server.py +664 -0
  22. package/mcp-servers/manual-review/README.md +133 -0
  23. package/mcp-servers/manual-review/server.py +910 -0
  24. package/mcp-servers/manual-review/ui.html +279 -0
  25. package/mcp-servers/minimax-chat/requirements.txt +1 -0
  26. package/mcp-servers/minimax-chat/server.py +381 -0
  27. package/package.json +51 -0
  28. package/skills/ablation-planner/SKILL.md +123 -0
  29. package/skills/alphaxiv/SKILL.md +196 -0
  30. package/skills/analyze-results/SKILL.md +46 -0
  31. package/skills/arxiv/SKILL.md +248 -0
  32. package/skills/auto-paper-improvement-loop/SKILL.md +651 -0
  33. package/skills/auto-review-loop/SKILL.md +1137 -0
  34. package/skills/auto-review-loop-llm/SKILL.md +259 -0
  35. package/skills/auto-review-loop-minimax/SKILL.md +302 -0
  36. package/skills/citation-audit/SKILL.md +502 -0
  37. package/skills/claims-drafting/SKILL.md +227 -0
  38. package/skills/comm-lit-review/SKILL.md +297 -0
  39. package/skills/deepxiv/SKILL.md +263 -0
  40. package/skills/dse-loop/SKILL.md +296 -0
  41. package/skills/embodiment-description/SKILL.md +129 -0
  42. package/skills/exa-search/SKILL.md +205 -0
  43. package/skills/experiment-audit/SKILL.md +311 -0
  44. package/skills/experiment-bridge/SKILL.md +376 -0
  45. package/skills/experiment-plan/SKILL.md +249 -0
  46. package/skills/experiment-queue/SKILL.md +431 -0
  47. package/skills/experiment-queue/scripts/build_manifest.py +142 -0
  48. package/skills/experiment-queue/scripts/queue_manager.py +433 -0
  49. package/skills/feishu-notify/SKILL.md +156 -0
  50. package/skills/figure-description/SKILL.md +138 -0
  51. package/skills/figure-spec/SKILL.md +262 -0
  52. package/skills/figure-spec/scripts/figure_renderer.py +799 -0
  53. package/skills/formula-derivation/SKILL.md +280 -0
  54. package/skills/gemini-search/SKILL.md +231 -0
  55. package/skills/grant-proposal/SKILL.md +698 -0
  56. package/skills/idea-creator/SKILL.md +542 -0
  57. package/skills/idea-discovery/SKILL.md +521 -0
  58. package/skills/idea-discovery-robot/SKILL.md +363 -0
  59. package/skills/integrity-forensics/SKILL.md +284 -0
  60. package/skills/interview-cheatsheet/SKILL.md +245 -0
  61. package/skills/invention-structuring/SKILL.md +188 -0
  62. package/skills/jurisdiction-format/SKILL.md +192 -0
  63. package/skills/kill-argument/SKILL.md +437 -0
  64. package/skills/mermaid-diagram/SKILL.md +419 -0
  65. package/skills/meta-apply/SKILL.md +141 -0
  66. package/skills/meta-optimize/SKILL.md +437 -0
  67. package/skills/monitor-experiment/SKILL.md +140 -0
  68. package/skills/novelty-check/SKILL.md +101 -0
  69. package/skills/openalex/SKILL.md +237 -0
  70. package/skills/overleaf-sync/SKILL.md +220 -0
  71. package/skills/paper-claim-audit/SKILL.md +348 -0
  72. package/skills/paper-compile/SKILL.md +266 -0
  73. package/skills/paper-figure/SKILL.md +312 -0
  74. package/skills/paper-illustration/SKILL.md +736 -0
  75. package/skills/paper-illustration-image2/SKILL.md +391 -0
  76. package/skills/paper-illustration-image2/scripts/paper_illustration_image2.py +255 -0
  77. package/skills/paper-plan/SKILL.md +386 -0
  78. package/skills/paper-poster/SKILL.md +19 -0
  79. package/skills/paper-poster-html/DESIGN_FINAL.md +176 -0
  80. package/skills/paper-poster-html/IMPLEMENTATION_CONVENTIONS.md +161 -0
  81. package/skills/paper-poster-html/LICENSES/posterly-MIT.txt +21 -0
  82. package/skills/paper-poster-html/NOTICE.md +57 -0
  83. package/skills/paper-poster-html/SKILL.md +323 -0
  84. package/skills/paper-poster-html/scripts/_posterly/__init__.py +0 -0
  85. package/skills/paper-poster-html/scripts/_posterly/canvas.py +200 -0
  86. package/skills/paper-poster-html/scripts/_posterly/measure.py +588 -0
  87. package/skills/paper-poster-html/scripts/_posterly/polish.py +498 -0
  88. package/skills/paper-poster-html/scripts/_posterly/preflight.py +489 -0
  89. package/skills/paper-poster-html/scripts/_posterly/render.py +215 -0
  90. package/skills/paper-poster-html/scripts/_posterly/textutil.py +16 -0
  91. package/skills/paper-poster-html/scripts/_posterly/verify_final.py +171 -0
  92. package/skills/paper-poster-html/scripts/asset_check.py +897 -0
  93. package/skills/paper-poster-html/scripts/extract_pdf_figures.py +666 -0
  94. package/skills/paper-poster-html/scripts/poster_check.py +251 -0
  95. package/skills/paper-poster-html/scripts/preprocess_figures.py +238 -0
  96. package/skills/paper-poster-html/scripts/render_preview.py +217 -0
  97. package/skills/paper-poster-html/scripts/run_gates.py +556 -0
  98. package/skills/paper-poster-html/scripts/style_check.py +1324 -0
  99. package/skills/paper-poster-html/templates/COMPONENTS.md +462 -0
  100. package/skills/paper-poster-html/templates/README.md +170 -0
  101. package/skills/paper-poster-html/templates/landscape_4col.html +1032 -0
  102. package/skills/paper-poster-html/templates/landscape_hero.html +1046 -0
  103. package/skills/paper-poster-html/templates/portrait_2col.html +947 -0
  104. package/skills/paper-poster-html/templates/tokens/acl.json +9 -0
  105. package/skills/paper-poster-html/templates/tokens/cvpr.json +9 -0
  106. package/skills/paper-poster-html/templates/tokens/generic.json +9 -0
  107. package/skills/paper-poster-html/templates/tokens/iclr.json +9 -0
  108. package/skills/paper-poster-html/templates/tokens/icml.json +9 -0
  109. package/skills/paper-poster-html/templates/tokens/neurips.json +9 -0
  110. package/skills/paper-slides/SKILL.md +635 -0
  111. package/skills/paper-talk/SKILL.md +381 -0
  112. package/skills/paper-write/SKILL.md +604 -0
  113. package/skills/paper-write/templates/IEEEtran.bst +2409 -0
  114. package/skills/paper-write/templates/IEEEtran.cls +6347 -0
  115. package/skills/paper-write/templates/iclr2026.tex +84 -0
  116. package/skills/paper-write/templates/icml2025.tex +87 -0
  117. package/skills/paper-write/templates/ieee_conference.tex +89 -0
  118. package/skills/paper-write/templates/ieee_journal.tex +93 -0
  119. package/skills/paper-write/templates/math_commands.tex +48 -0
  120. package/skills/paper-write/templates/neurips2025.tex +80 -0
  121. package/skills/paper-writing/SKILL.md +916 -0
  122. package/skills/patent-novelty-check/SKILL.md +153 -0
  123. package/skills/patent-pipeline/SKILL.md +344 -0
  124. package/skills/patent-review/SKILL.md +203 -0
  125. package/skills/pixel-art/SKILL.md +137 -0
  126. package/skills/prior-art-search/SKILL.md +146 -0
  127. package/skills/proof-checker/SKILL.md +866 -0
  128. package/skills/proof-orchestrator/NOTICE.md +24 -0
  129. package/skills/proof-orchestrator/SKILL.md +254 -0
  130. package/skills/proof-orchestrator/references/audit-output-contract.md +126 -0
  131. package/skills/proof-orchestrator/references/deepseek-routing.md +74 -0
  132. package/skills/proof-orchestrator/references/dispatch-prompts.md +227 -0
  133. package/skills/proof-orchestrator/references/notation-audit.md +135 -0
  134. package/skills/proof-orchestrator/references/proof-audit-rubric.md +70 -0
  135. package/skills/proof-orchestrator/references/stress-tests.md +38 -0
  136. package/skills/proof-writer/SKILL.md +223 -0
  137. package/skills/qzcli/SKILL.md +324 -0
  138. package/skills/rebuttal/SKILL.md +376 -0
  139. package/skills/render-html/SKILL.md +316 -0
  140. package/skills/render-html/scripts/render_html.py +1006 -0
  141. package/skills/render-html/scripts/templates/academic.html +703 -0
  142. package/skills/render-html/scripts/templates/dashboard.html +333 -0
  143. package/skills/research-lit/SKILL.md +756 -0
  144. package/skills/research-pipeline/SKILL.md +384 -0
  145. package/skills/research-refine/SKILL.md +770 -0
  146. package/skills/research-refine-pipeline/SKILL.md +186 -0
  147. package/skills/research-review/SKILL.md +198 -0
  148. package/skills/research-wiki/SKILL.md +461 -0
  149. package/skills/resubmit-pipeline/SKILL.md +447 -0
  150. package/skills/result-to-claim/SKILL.md +311 -0
  151. package/skills/run-experiment/SKILL.md +313 -0
  152. package/skills/semantic-scholar/SKILL.md +236 -0
  153. package/skills/serverless-modal/SKILL.md +335 -0
  154. package/skills/shared-references/acceptance-gate.md +324 -0
  155. package/skills/shared-references/assurance-contract.md +248 -0
  156. package/skills/shared-references/capture-antipatterns.md +78 -0
  157. package/skills/shared-references/citation-discipline.md +583 -0
  158. package/skills/shared-references/compute-env-contract.md +163 -0
  159. package/skills/shared-references/effort-contract.md +183 -0
  160. package/skills/shared-references/evidence-precheck.md +65 -0
  161. package/skills/shared-references/experiment-integrity.md +49 -0
  162. package/skills/shared-references/external-cadence.md +326 -0
  163. package/skills/shared-references/fan-out-pattern.md +366 -0
  164. package/skills/shared-references/injection-hygiene.md +127 -0
  165. package/skills/shared-references/integration-contract.md +461 -0
  166. package/skills/shared-references/output-composition.md +93 -0
  167. package/skills/shared-references/output-language.md +45 -0
  168. package/skills/shared-references/output-manifest.md +49 -0
  169. package/skills/shared-references/output-versioning.md +111 -0
  170. package/skills/shared-references/patent-format-cn.md +199 -0
  171. package/skills/shared-references/patent-format-ep.md +173 -0
  172. package/skills/shared-references/patent-format-us.md +161 -0
  173. package/skills/shared-references/patent-writing-principles.md +197 -0
  174. package/skills/shared-references/prior-art-databases.md +141 -0
  175. package/skills/shared-references/resumable-runs.md +109 -0
  176. package/skills/shared-references/review-scope-limits.md +81 -0
  177. package/skills/shared-references/review-tracing.md +391 -0
  178. package/skills/shared-references/reviewer-independence.md +79 -0
  179. package/skills/shared-references/reviewer-routing.md +852 -0
  180. package/skills/shared-references/skill-governance.md +104 -0
  181. package/skills/shared-references/taste-calibration.md +85 -0
  182. package/skills/shared-references/venue-checklists.md +114 -0
  183. package/skills/shared-references/wiki-helper-resolution.md +134 -0
  184. package/skills/shared-references/writing-principles.md +525 -0
  185. package/skills/skills-codex/README.md +102 -0
  186. package/skills/skills-codex/README_CN.md +100 -0
  187. package/skills/skills-codex/ablation-planner/SKILL.md +126 -0
  188. package/skills/skills-codex/alphaxiv/SKILL.md +186 -0
  189. package/skills/skills-codex/analyze-results/SKILL.md +45 -0
  190. package/skills/skills-codex/arxiv/SKILL.md +210 -0
  191. package/skills/skills-codex/auto-paper-improvement-loop/SKILL.md +574 -0
  192. package/skills/skills-codex/auto-review-loop/SKILL.md +500 -0
  193. package/skills/skills-codex/auto-review-loop-llm/SKILL.md +247 -0
  194. package/skills/skills-codex/auto-review-loop-minimax/SKILL.md +290 -0
  195. package/skills/skills-codex/citation-audit/SKILL.md +504 -0
  196. package/skills/skills-codex/claims-drafting/SKILL.md +239 -0
  197. package/skills/skills-codex/comm-lit-review/SKILL.md +299 -0
  198. package/skills/skills-codex/comm-lit-review/references/domain-taxonomy.md +57 -0
  199. package/skills/skills-codex/comm-lit-review/references/output-template.md +37 -0
  200. package/skills/skills-codex/comm-lit-review/references/source-policy.md +99 -0
  201. package/skills/skills-codex/comm-lit-review/references/venue-tiering.md +112 -0
  202. package/skills/skills-codex/deepxiv/SKILL.md +142 -0
  203. package/skills/skills-codex/dse-loop/SKILL.md +285 -0
  204. package/skills/skills-codex/embodiment-description/SKILL.md +129 -0
  205. package/skills/skills-codex/exa-search/SKILL.md +192 -0
  206. package/skills/skills-codex/experiment-audit/SKILL.md +286 -0
  207. package/skills/skills-codex/experiment-bridge/SKILL.md +356 -0
  208. package/skills/skills-codex/experiment-plan/SKILL.md +249 -0
  209. package/skills/skills-codex/experiment-queue/SKILL.md +401 -0
  210. package/skills/skills-codex/feishu-notify/SKILL.md +155 -0
  211. package/skills/skills-codex/figure-description/SKILL.md +138 -0
  212. package/skills/skills-codex/figure-spec/SKILL.md +252 -0
  213. package/skills/skills-codex/formula-derivation/SKILL.md +280 -0
  214. package/skills/skills-codex/gemini-search/SKILL.md +205 -0
  215. package/skills/skills-codex/grant-proposal/SKILL.md +626 -0
  216. package/skills/skills-codex/idea-creator/SKILL.md +405 -0
  217. package/skills/skills-codex/idea-discovery/SKILL.md +475 -0
  218. package/skills/skills-codex/idea-discovery-robot/SKILL.md +362 -0
  219. package/skills/skills-codex/integrity-forensics/SKILL.md +106 -0
  220. package/skills/skills-codex/interview-cheatsheet/SKILL.md +245 -0
  221. package/skills/skills-codex/invention-structuring/SKILL.md +188 -0
  222. package/skills/skills-codex/jurisdiction-format/SKILL.md +192 -0
  223. package/skills/skills-codex/kill-argument/SKILL.md +403 -0
  224. package/skills/skills-codex/mermaid-diagram/SKILL.md +379 -0
  225. package/skills/skills-codex/meta-apply/SKILL.md +154 -0
  226. package/skills/skills-codex/meta-optimize/SKILL.md +348 -0
  227. package/skills/skills-codex/monitor-experiment/SKILL.md +98 -0
  228. package/skills/skills-codex/novelty-check/SKILL.md +89 -0
  229. package/skills/skills-codex/openalex/SKILL.md +228 -0
  230. package/skills/skills-codex/overleaf-sync/SKILL.md +220 -0
  231. package/skills/skills-codex/paper-claim-audit/SKILL.md +350 -0
  232. package/skills/skills-codex/paper-compile/SKILL.md +253 -0
  233. package/skills/skills-codex/paper-figure/SKILL.md +311 -0
  234. package/skills/skills-codex/paper-illustration/SKILL.md +690 -0
  235. package/skills/skills-codex/paper-illustration-image2/SKILL.md +383 -0
  236. package/skills/skills-codex/paper-illustration-image2/scripts/paper_illustration_image2.py +255 -0
  237. package/skills/skills-codex/paper-plan/SKILL.md +278 -0
  238. package/skills/skills-codex/paper-poster/SKILL.md +19 -0
  239. package/skills/skills-codex/paper-poster-html/SKILL.md +377 -0
  240. package/skills/skills-codex/paper-slides/SKILL.md +571 -0
  241. package/skills/skills-codex/paper-talk/SKILL.md +381 -0
  242. package/skills/skills-codex/paper-write/SKILL.md +411 -0
  243. package/skills/skills-codex/paper-write/templates/IEEEtran.bst +2409 -0
  244. package/skills/skills-codex/paper-write/templates/IEEEtran.cls +6347 -0
  245. package/skills/skills-codex/paper-write/templates/aaai2026.bst +1493 -0
  246. package/skills/skills-codex/paper-write/templates/aaai2026.sty +315 -0
  247. package/skills/skills-codex/paper-write/templates/aaai2026.tex +952 -0
  248. package/skills/skills-codex/paper-write/templates/acl.sty +312 -0
  249. package/skills/skills-codex/paper-write/templates/acl2026.tex +377 -0
  250. package/skills/skills-codex/paper-write/templates/acl_natbib.bst +1940 -0
  251. package/skills/skills-codex/paper-write/templates/acm.bst +3081 -0
  252. package/skills/skills-codex/paper-write/templates/acm_mm2026.tex +204 -0
  253. package/skills/skills-codex/paper-write/templates/acmart.cls +3520 -0
  254. package/skills/skills-codex/paper-write/templates/cvpr.bst +1448 -0
  255. package/skills/skills-codex/paper-write/templates/cvpr.sty +508 -0
  256. package/skills/skills-codex/paper-write/templates/cvpr2026.tex +63 -0
  257. package/skills/skills-codex/paper-write/templates/iclr2026.tex +84 -0
  258. package/skills/skills-codex/paper-write/templates/iclr2026_conference.bst +1440 -0
  259. package/skills/skills-codex/paper-write/templates/iclr2026_conference.sty +246 -0
  260. package/skills/skills-codex/paper-write/templates/icml2026.sty +767 -0
  261. package/skills/skills-codex/paper-write/templates/icml2026.tex +662 -0
  262. package/skills/skills-codex/paper-write/templates/ieee_conference.tex +89 -0
  263. package/skills/skills-codex/paper-write/templates/ieee_journal.tex +93 -0
  264. package/skills/skills-codex/paper-write/templates/math_commands.tex +48 -0
  265. package/skills/skills-codex/paper-write/templates/neurips2026.tex +493 -0
  266. package/skills/skills-codex/paper-write/templates/neurips_2026.sty +437 -0
  267. package/skills/skills-codex/paper-writing/SKILL.md +731 -0
  268. package/skills/skills-codex/patent-novelty-check/SKILL.md +153 -0
  269. package/skills/skills-codex/patent-pipeline/SKILL.md +344 -0
  270. package/skills/skills-codex/patent-review/SKILL.md +202 -0
  271. package/skills/skills-codex/pixel-art/SKILL.md +139 -0
  272. package/skills/skills-codex/prior-art-search/SKILL.md +146 -0
  273. package/skills/skills-codex/proof-checker/SKILL.md +554 -0
  274. package/skills/skills-codex/proof-orchestrator/SKILL.md +260 -0
  275. package/skills/skills-codex/proof-orchestrator/references/audit-output-contract.md +126 -0
  276. package/skills/skills-codex/proof-orchestrator/references/deepseek-routing.md +76 -0
  277. package/skills/skills-codex/proof-orchestrator/references/dispatch-prompts.md +227 -0
  278. package/skills/skills-codex/proof-orchestrator/references/notation-audit.md +135 -0
  279. package/skills/skills-codex/proof-orchestrator/references/proof-audit-rubric.md +70 -0
  280. package/skills/skills-codex/proof-orchestrator/references/stress-tests.md +38 -0
  281. package/skills/skills-codex/proof-writer/SKILL.md +222 -0
  282. package/skills/skills-codex/qzcli/SKILL.md +324 -0
  283. package/skills/skills-codex/rebuttal/SKILL.md +305 -0
  284. package/skills/skills-codex/render-html/SKILL.md +305 -0
  285. package/skills/skills-codex/render-html/scripts/__pycache__/render_html.cpython-314.pyc +0 -0
  286. package/skills/skills-codex/render-html/scripts/render_html.py +909 -0
  287. package/skills/skills-codex/render-html/scripts/templates/academic.html +342 -0
  288. package/skills/skills-codex/render-html/scripts/templates/dashboard.html +333 -0
  289. package/skills/skills-codex/research-lit/SKILL.md +464 -0
  290. package/skills/skills-codex/research-pipeline/SKILL.md +340 -0
  291. package/skills/skills-codex/research-refine/SKILL.md +721 -0
  292. package/skills/skills-codex/research-refine-pipeline/SKILL.md +186 -0
  293. package/skills/skills-codex/research-review/SKILL.md +135 -0
  294. package/skills/skills-codex/research-wiki/SKILL.md +421 -0
  295. package/skills/skills-codex/resubmit-pipeline/SKILL.md +444 -0
  296. package/skills/skills-codex/result-to-claim/SKILL.md +246 -0
  297. package/skills/skills-codex/run-experiment/SKILL.md +236 -0
  298. package/skills/skills-codex/semantic-scholar/SKILL.md +219 -0
  299. package/skills/skills-codex/serverless-modal/SKILL.md +335 -0
  300. package/skills/skills-codex/shared-references/acceptance-gate.md +336 -0
  301. package/skills/skills-codex/shared-references/assurance-contract.md +139 -0
  302. package/skills/skills-codex/shared-references/capture-antipatterns.md +84 -0
  303. package/skills/skills-codex/shared-references/citation-discipline.md +452 -0
  304. package/skills/skills-codex/shared-references/compute-env-contract.md +163 -0
  305. package/skills/skills-codex/shared-references/effort-contract.md +143 -0
  306. package/skills/skills-codex/shared-references/evidence-precheck.md +73 -0
  307. package/skills/skills-codex/shared-references/experiment-integrity.md +49 -0
  308. package/skills/skills-codex/shared-references/external-cadence.md +334 -0
  309. package/skills/skills-codex/shared-references/fan-out-pattern.md +375 -0
  310. package/skills/skills-codex/shared-references/injection-hygiene.md +132 -0
  311. package/skills/skills-codex/shared-references/integration-contract.md +372 -0
  312. package/skills/skills-codex/shared-references/output-composition.md +98 -0
  313. package/skills/skills-codex/shared-references/output-language.md +45 -0
  314. package/skills/skills-codex/shared-references/output-manifest.md +40 -0
  315. package/skills/skills-codex/shared-references/output-versioning.md +111 -0
  316. package/skills/skills-codex/shared-references/patent-format-cn.md +199 -0
  317. package/skills/skills-codex/shared-references/patent-format-ep.md +173 -0
  318. package/skills/skills-codex/shared-references/patent-format-us.md +161 -0
  319. package/skills/skills-codex/shared-references/patent-writing-principles.md +197 -0
  320. package/skills/skills-codex/shared-references/prior-art-databases.md +141 -0
  321. package/skills/skills-codex/shared-references/resumable-runs.md +125 -0
  322. package/skills/skills-codex/shared-references/review-scope-limits.md +81 -0
  323. package/skills/skills-codex/shared-references/review-tracing.md +144 -0
  324. package/skills/skills-codex/shared-references/reviewer-independence.md +66 -0
  325. package/skills/skills-codex/shared-references/reviewer-routing.md +128 -0
  326. package/skills/skills-codex/shared-references/skill-governance.md +119 -0
  327. package/skills/skills-codex/shared-references/taste-calibration.md +90 -0
  328. package/skills/skills-codex/shared-references/venue-checklists.md +73 -0
  329. package/skills/skills-codex/shared-references/wiki-helper-resolution.md +69 -0
  330. package/skills/skills-codex/shared-references/writing-principles.md +525 -0
  331. package/skills/skills-codex/slides-polish/SKILL.md +563 -0
  332. package/skills/skills-codex/specification-writing/SKILL.md +211 -0
  333. package/skills/skills-codex/system-profile/SKILL.md +103 -0
  334. package/skills/skills-codex/training-check/SKILL.md +83 -0
  335. package/skills/skills-codex/vast-gpu/SKILL.md +394 -0
  336. package/skills/skills-codex/web-debug-search/SKILL.md +334 -0
  337. package/skills/skills-codex/wiki-enrich/SKILL.md +255 -0
  338. package/skills/skills-codex/writing-systems-papers/SKILL.md +184 -0
  339. package/skills/skills-codex-claude-review/README.md +79 -0
  340. package/skills/skills-codex-claude-review/README_CN.md +78 -0
  341. package/skills/skills-codex-claude-review/auto-paper-improvement-loop/SKILL.md +581 -0
  342. package/skills/skills-codex-claude-review/auto-review-loop/SKILL.md +510 -0
  343. package/skills/skills-codex-claude-review/novelty-check/SKILL.md +102 -0
  344. package/skills/skills-codex-claude-review/paper-figure/SKILL.md +319 -0
  345. package/skills/skills-codex-claude-review/paper-plan/SKILL.md +287 -0
  346. package/skills/skills-codex-claude-review/paper-write/SKILL.md +420 -0
  347. package/skills/skills-codex-claude-review/research-refine/SKILL.md +732 -0
  348. package/skills/skills-codex-claude-review/research-review/SKILL.md +149 -0
  349. package/skills/skills-codex-gemini-review/README.md +176 -0
  350. package/skills/skills-codex-gemini-review/README_CN.md +175 -0
  351. package/skills/skills-codex-gemini-review/auto-paper-improvement-loop/SKILL.md +331 -0
  352. package/skills/skills-codex-gemini-review/auto-review-loop/SKILL.md +304 -0
  353. package/skills/skills-codex-gemini-review/grant-proposal/SKILL.md +630 -0
  354. package/skills/skills-codex-gemini-review/idea-creator/SKILL.md +263 -0
  355. package/skills/skills-codex-gemini-review/idea-discovery/SKILL.md +275 -0
  356. package/skills/skills-codex-gemini-review/idea-discovery-robot/SKILL.md +365 -0
  357. package/skills/skills-codex-gemini-review/novelty-check/SKILL.md +92 -0
  358. package/skills/skills-codex-gemini-review/paper-figure/SKILL.md +289 -0
  359. package/skills/skills-codex-gemini-review/paper-plan/SKILL.md +265 -0
  360. package/skills/skills-codex-gemini-review/paper-poster-html/SKILL.md +102 -0
  361. package/skills/skills-codex-gemini-review/paper-slides/SKILL.md +582 -0
  362. package/skills/skills-codex-gemini-review/paper-write/SKILL.md +346 -0
  363. package/skills/skills-codex-gemini-review/paper-writing/SKILL.md +312 -0
  364. package/skills/skills-codex-gemini-review/research-refine/SKILL.md +674 -0
  365. package/skills/skills-codex-gemini-review/research-review/SKILL.md +112 -0
  366. package/skills/slides-polish/SKILL.md +565 -0
  367. package/skills/specification-writing/SKILL.md +211 -0
  368. package/skills/system-profile/SKILL.md +103 -0
  369. package/skills/training-check/SKILL.md +132 -0
  370. package/skills/vast-gpu/SKILL.md +394 -0
  371. package/skills/web-debug-search/SKILL.md +334 -0
  372. package/skills/wiki-enrich/SKILL.md +257 -0
  373. package/skills/writing-systems-papers/SKILL.md +184 -0
  374. package/templates/CLAUDE_MD_TEMPLATE.md +29 -0
  375. package/templates/EXPERIMENT_LOG_TEMPLATE.md +47 -0
  376. package/templates/EXPERIMENT_PLAN_TEMPLATE.md +51 -0
  377. package/templates/EXPERIMENT_PLAN_TEMPLATE_CN.md +53 -0
  378. package/templates/FINDINGS_TEMPLATE.md +52 -0
  379. package/templates/IDEA_CANDIDATES_TEMPLATE.md +47 -0
  380. package/templates/IDEA_CANDIDATES_TEMPLATE_CN.md +47 -0
  381. package/templates/INVENTION_BRIEF_TEMPLATE.md +87 -0
  382. package/templates/MANIFEST_TEMPLATE.md +7 -0
  383. package/templates/NARRATIVE_REPORT_TEMPLATE.md +49 -0
  384. package/templates/PAPER_PLAN_TEMPLATE.md +47 -0
  385. package/templates/PATENT_CLAIMS_TEMPLATE.md +78 -0
  386. package/templates/PATENT_SPECIFICATION_TEMPLATE.md +68 -0
  387. package/templates/README.md +57 -0
  388. package/templates/RESEARCH_BRIEF_TEMPLATE.md +35 -0
  389. package/templates/RESEARCH_BRIEF_TEMPLATE_CN.md +41 -0
  390. package/templates/RESEARCH_CONTRACT_TEMPLATE.md +60 -0
  391. package/templates/claude-hooks/corpus_write_guard.json +16 -0
  392. package/templates/claude-hooks/corpus_write_guard.py +85 -0
  393. package/templates/claude-hooks/meta_logging.json +74 -0
  394. package/templates/gitignore-trace.txt +3 -0
  395. package/tools/__pycache__/check_skills_inventory.cpython-314.pyc +0 -0
  396. package/tools/arxiv_fetch.py +311 -0
  397. package/tools/capture_filter.py +126 -0
  398. package/tools/check_skills_inventory.py +273 -0
  399. package/tools/convert_skills_to_llm_chat.py +282 -0
  400. package/tools/copilot_native_evidence.py +818 -0
  401. package/tools/deepxiv_fetch.py +213 -0
  402. package/tools/evidence_check.py +212 -0
  403. package/tools/exa_search.py +425 -0
  404. package/tools/experiment_queue/README.md +118 -0
  405. package/tools/experiment_queue/build_manifest.py +44 -0
  406. package/tools/experiment_queue/queue_manager.py +44 -0
  407. package/tools/extract_paper_style.py +560 -0
  408. package/tools/figure_renderer.py +69 -0
  409. package/tools/forensics_gate.py +669 -0
  410. package/tools/generate_codex_claude_review_overrides.py +299 -0
  411. package/tools/idea_discovery_gate.py +256 -0
  412. package/tools/install_aris.ps1 +1372 -0
  413. package/tools/install_aris.sh +1370 -0
  414. package/tools/install_aris_codex.sh +1023 -0
  415. package/tools/install_aris_copilot.sh +1052 -0
  416. package/tools/iteration_log.py +143 -0
  417. package/tools/lint_skills_helpers.sh +84 -0
  418. package/tools/meta_opt/check_ready.sh +80 -0
  419. package/tools/meta_opt/log_event.sh +91 -0
  420. package/tools/meta_opt/trigger_eval.py +280 -0
  421. package/tools/meta_opt/trigger_evals.sample.json +28 -0
  422. package/tools/openalex_fetch.py +326 -0
  423. package/tools/overleaf_audit.sh +104 -0
  424. package/tools/overleaf_setup.sh +150 -0
  425. package/tools/paper_illustration_image2.py +62 -0
  426. package/tools/provenance.py +294 -0
  427. package/tools/research_wiki.py +1720 -0
  428. package/tools/review_gate.py +502 -0
  429. package/tools/run_state.py +399 -0
  430. package/tools/save_trace.sh +477 -0
  431. package/tools/semantic_scholar_fetch.py +438 -0
  432. package/tools/skill-groups.tsv +116 -0
  433. package/tools/skill_picker.py +238 -0
  434. package/tools/smart_update.ps1 +521 -0
  435. package/tools/smart_update.sh +591 -0
  436. package/tools/smart_update_codex.sh +419 -0
  437. package/tools/smart_update_copilot.sh +605 -0
  438. package/tools/threat_scan.py +222 -0
  439. package/tools/verify_paper_audits.sh +487 -0
  440. package/tools/verify_papers.py +613 -0
  441. package/tools/verify_wiki_coverage.sh +176 -0
  442. package/tools/watchdog.py +485 -0
@@ -0,0 +1,247 @@
1
+ ---
2
+ name: "auto-review-loop-llm"
3
+ description: "Autonomous research review loop using any OpenAI-compatible LLM API. Configure via llm-chat MCP server or environment variables. Trigger with \"auto review loop llm\" or \"llm review\"."
4
+ ---
5
+
6
+ # Auto Review Loop (Generic LLM): Autonomous Research Improvement
7
+
8
+ Autonomously iterate: review → implement fixes → re-review, until the external reviewer gives a positive assessment or MAX_ROUNDS is reached.
9
+
10
+ ## Context: $ARGUMENTS
11
+
12
+ ## Constants
13
+
14
+ - MAX_ROUNDS = 4
15
+ - POSITIVE_THRESHOLD: score >= 6/10 AND verdict ∈ {"ready", "almost"} — both must hold, matching the operative STOP CONDITION below. Verdict vocabulary is {"ready", "almost", "not ready"}. (Earlier wording used "or" + a stale verdict set; the AND form is authoritative.)
16
+ - REVIEW_DOC: `review-stage/AUTO_REVIEW.md` (cumulative log) *(fall back to `./AUTO_REVIEW.md` for legacy projects)*
17
+
18
+ ## LLM Configuration
19
+
20
+ This skill uses **any OpenAI-compatible API** for external review via the `llm-chat` MCP server.
21
+
22
+ ### Configuration via MCP Server (Recommended)
23
+
24
+ Add to `~/.codex/settings.json`:
25
+
26
+ ```json
27
+ {
28
+ "mcpServers": {
29
+ "llm-chat": {
30
+ "command": "/usr/bin/python3",
31
+ "args": ["/Users/yourname/.codex/mcp-servers/llm-chat/server.py"],
32
+ "env": {
33
+ "LLM_API_KEY": "your-api-key",
34
+ "LLM_BASE_URL": "https://api.deepseek.com/v1",
35
+ "LLM_MODEL": "deepseek-chat"
36
+ }
37
+ }
38
+ }
39
+ }
40
+ ```
41
+
42
+ ### Supported Providers
43
+
44
+ | Provider | LLM_BASE_URL | LLM_MODEL |
45
+ |----------|--------------|-----------|
46
+ | **OpenAI** | `https://api.openai.com/v1` | `gpt-4o`, `o3` |
47
+ | **DeepSeek** | `https://api.deepseek.com/v1` | `deepseek-chat`, `deepseek-reasoner` |
48
+ | **MiniMax** | `https://api.minimax.io/v1` | `MiniMax-M3` |
49
+ | **Kimi (Moonshot)** | `https://api.moonshot.cn/v1` | `moonshot-v1-8k`, `moonshot-v1-32k` |
50
+ | **ZhiPu (GLM)** | `https://open.bigmodel.cn/api/paas/v4` | `glm-4`, `glm-4-plus` |
51
+ | **SiliconFlow** | `https://api.siliconflow.cn/v1` | `Qwen/Qwen2.5-72B-Instruct` |
52
+ | **阿里云百炼** | `https://dashscope.aliyuncs.com/compatible-mode/v1` | `qwen-max` |
53
+ | **零一万物** | `https://api.lingyiwanwu.com/v1` | `yi-large` |
54
+
55
+ ## API Call Method
56
+
57
+ **Primary: MCP Tool**
58
+
59
+ ```
60
+ mcp__llm-chat__chat:
61
+ message: |
62
+ [Review prompt content]
63
+ model: "deepseek-chat"
64
+ system: "You are a senior ML reviewer..."
65
+ ```
66
+
67
+ **Fallback: curl**
68
+
69
+ ```bash
70
+ curl -s "${LLM_BASE_URL}/chat/completions" \
71
+ -H "Content-Type: application/json" \
72
+ -H "Authorization: Bearer ${LLM_API_KEY}" \
73
+ -d '{
74
+ "model": "${LLM_MODEL}",
75
+ "messages": [
76
+ {"role": "system", "content": "You are a senior ML reviewer..."},
77
+ {"role": "user", "content": "[review prompt]"}
78
+ ],
79
+ "max_tokens": 4096
80
+ }'
81
+ ```
82
+
83
+ ## State Persistence (Compact Recovery)
84
+
85
+ Persist state to `review-stage/REVIEW_STATE.json` after each round:
86
+
87
+ ```json
88
+ {
89
+ "round": 2,
90
+ "status": "in_progress",
91
+ "last_score": 5.0,
92
+ "last_verdict": "not ready",
93
+ "pending_experiments": [],
94
+ "timestamp": "2026-03-15T10:00:00"
95
+ }
96
+ ```
97
+
98
+ **Write this file at the end of every Phase E** (after documenting the round).
99
+
100
+ **On completion**, set `"status": "completed"`.
101
+
102
+ ## Workflow
103
+
104
+ ### Initialization
105
+
106
+ 1. **Check `review-stage/REVIEW_STATE.json`** for recovery *(fall back to `./REVIEW_STATE.json` if not found — legacy path)*
107
+ 2. Read project context and prior reviews
108
+ 3. Initialize round counter
109
+
110
+ ### Loop (up to MAX_ROUNDS)
111
+
112
+ #### Phase A: Review
113
+
114
+ **If MCP available:**
115
+ ```
116
+ mcp__llm-chat__chat:
117
+ system: "You are a senior ML reviewer (NeurIPS/ICML level)."
118
+ message: |
119
+ [Round N/MAX_ROUNDS of autonomous review loop]
120
+
121
+ [Full research context: claims, methods, results, known weaknesses]
122
+ [Changes since last round, if any]
123
+
124
+ 1. Score this work 1-10 for a top venue
125
+ 2. List remaining critical weaknesses (ranked by severity)
126
+ 3. For each weakness, specify the MINIMUM fix
127
+ 4. State clearly: is this READY for submission? Yes/No/Almost
128
+
129
+ Be brutally honest. If the work is ready, say so clearly.
130
+ ```
131
+
132
+ **If MCP NOT available:**
133
+ ```bash
134
+ curl -s "${LLM_BASE_URL}/chat/completions" \
135
+ -H "Content-Type: application/json" \
136
+ -H "Authorization: Bearer ${LLM_API_KEY}" \
137
+ -d '{
138
+ "model": "${LLM_MODEL}",
139
+ "messages": [
140
+ {"role": "system", "content": "You are a senior ML reviewer (NeurIPS/ICML level)."},
141
+ {"role": "user", "content": "[Full review prompt]"}
142
+ ],
143
+ "max_tokens": 4096
144
+ }'
145
+ ```
146
+
147
+ #### Phase B: Parse Assessment
148
+
149
+ **CRITICAL: Save the FULL raw response** verbatim. Then extract:
150
+ - **Score** (numeric 1-10)
151
+ - **Verdict** ("ready" / "almost" / "not ready")
152
+ - **Action items** (ranked list of fixes)
153
+
154
+ **STOP**: If score >= 6 AND verdict ∈ {"ready", "almost"} (exact — "not ready" does NOT qualify)
155
+
156
+ #### Phase C: Implement Fixes
157
+
158
+ Priority: metric additions > reframing > new experiments
159
+
160
+ #### Phase D: Wait for Results
161
+
162
+ Monitor remote experiments
163
+
164
+ #### Phase E: Document Round
165
+
166
+ Append to `review-stage/AUTO_REVIEW.md`:
167
+
168
+ ```markdown
169
+ ## Round N (timestamp)
170
+
171
+ ### Assessment (Summary)
172
+ - Score: X/10
173
+ - Verdict: [ready/almost/not ready]
174
+ - Key criticisms: [bullet list]
175
+
176
+ ### Reviewer Raw Response
177
+
178
+ <details>
179
+ <summary>Click to expand full reviewer response</summary>
180
+
181
+ [Paste the COMPLETE raw response here — verbatim, unedited.]
182
+
183
+ </details>
184
+
185
+ ### Actions Taken
186
+ - [what was implemented/changed]
187
+
188
+ ### Results
189
+ - [experiment outcomes, if any]
190
+
191
+ ### Status
192
+ - [continuing to round N+1 / stopping]
193
+ ```
194
+
195
+ **Write `review-stage/REVIEW_STATE.json`** with current state.
196
+
197
+ ### Termination
198
+
199
+ 1. Set `review-stage/REVIEW_STATE.json` status to "completed"
200
+ 2. Write final summary
201
+
202
+ ## Key Rules
203
+
204
+ - **Large file handling**: If the Write tool fails due to file size, immediately retry using Bash (`cat << 'EOF' > file`) to write in chunks. Do NOT ask the user for permission — just do it silently.
205
+
206
+ - Be honest about weaknesses
207
+ - Implement fixes BEFORE re-reviewing
208
+ - Document everything
209
+ - Include previous context in round 2+ prompts
210
+ - Prefer MCP tool over curl when available
211
+
212
+ ## Prompt Template for Round 2+
213
+
214
+ ```
215
+ mcp__llm-chat__chat:
216
+ system: "You are a senior ML reviewer (NeurIPS/ICML level)."
217
+ message: |
218
+ [Round N/MAX_ROUNDS of autonomous review loop]
219
+
220
+ ## Previous Review Summary (Round N-1)
221
+ - Previous Score: X/10
222
+ - Previous Verdict: [ready/almost/not ready]
223
+ - Previous Key Weaknesses: [list]
224
+
225
+ ## Changes Since Last Review
226
+ 1. [Action 1]: [result]
227
+ 2. [Action 2]: [result]
228
+
229
+ ## Updated Results
230
+ [paste updated metrics/tables]
231
+
232
+ Please re-score and re-assess:
233
+ 1. Score this work 1-10 for a top venue
234
+ 2. List remaining critical weaknesses (ranked by severity)
235
+ 3. For each weakness, specify the MINIMUM fix
236
+ 4. State clearly: is this READY for submission? Yes/No/Almost
237
+
238
+ Be brutally honest. If the work is ready, say so clearly.
239
+ ```
240
+
241
+
242
+ ## Output Protocols
243
+
244
+ > Follow these shared protocols for all output files:
245
+ > - **[Output Versioning Protocol](../../shared-references/output-versioning.md)** — write timestamped file first, then copy to fixed name
246
+ > - **[Output Manifest Protocol](../../shared-references/output-manifest.md)** — log every output to MANIFEST.md
247
+ > - **[Output Language Protocol](../../shared-references/output-language.md)** — respect the project's language setting
@@ -0,0 +1,290 @@
1
+ ---
2
+ name: "auto-review-loop-minimax"
3
+ description: "Autonomous multi-round research review loop using MiniMax API. Use when you want to use MiniMax instead of Codex MCP for external review. Trigger with \"auto review loop minimax\" or \"minimax review\"."
4
+ ---
5
+
6
+ # Auto Review Loop (MiniMax Version): Autonomous Research Improvement
7
+
8
+ Autonomously iterate: review → implement fixes → re-review, until the external reviewer gives a positive assessment or MAX_ROUNDS is reached.
9
+
10
+ ## Context: $ARGUMENTS
11
+
12
+ ## Constants
13
+
14
+ - MAX_ROUNDS = 4
15
+ - POSITIVE_THRESHOLD: score >= 6/10 AND verdict ∈ {"ready", "almost"} — both must hold, matching the operative STOP CONDITION below. Verdict vocabulary is {"ready", "almost", "not ready"}. (Earlier wording used "or" + a stale verdict set; the AND form is authoritative.)
16
+ - REVIEW_DOC: `review-stage/AUTO_REVIEW.md` (cumulative log) *(fall back to `./AUTO_REVIEW.md` for legacy projects)*
17
+ - REVIEWER_MODEL = `MiniMax-M3` — Model used via MiniMax API
18
+
19
+ ## API Configuration
20
+
21
+ This skill uses MiniMax API for external review. Two methods are supported:
22
+
23
+ ### Method 1: MCP Tool (Primary)
24
+
25
+ If `mcp__minimax-chat__minimax_chat` is available, use it:
26
+
27
+ ```
28
+ mcp__minimax-chat__minimax_chat:
29
+ message: |
30
+ [Review prompt content]
31
+ model: "MiniMax-M3"
32
+ system: "You are a senior machine learning researcher..."
33
+ ```
34
+
35
+ ### Method 2: curl (Fallback)
36
+
37
+ If MCP is not available, use curl directly:
38
+
39
+ ```bash
40
+ curl -s "https://api.minimax.io/v1/chat/completions" \
41
+ -H "Content-Type: application/json" \
42
+ -H "Authorization: Bearer $MINIMAX_API_KEY" \
43
+ -d '{
44
+ "model": "MiniMax-M3",
45
+ "messages": [
46
+ {"role": "system", "content": "You are a senior ML researcher..."},
47
+ {"role": "user", "content": "[Review prompt]"}
48
+ ],
49
+ "max_tokens": 4096
50
+ }'
51
+ ```
52
+
53
+ **API Key**: Read from `~/.codex/settings.json` under `env.MINIMAX_API_KEY`, or from environment variable.
54
+
55
+ **Why MiniMax instead of a secondary Codex agent?** Codex CLI uses OpenAI's Responses API (`/v1/responses`) which is not supported by third-party providers. See: https://github.com/openai/codex/discussions/7782
56
+
57
+ ## State Persistence (Compact Recovery)
58
+
59
+ Long-running loops may hit the context window limit, triggering automatic compaction. To survive this, persist state to `review-stage/REVIEW_STATE.json` after each round:
60
+
61
+ ```json
62
+ {
63
+ "round": 2,
64
+ "status": "in_progress",
65
+ "last_score": 5.0,
66
+ "last_verdict": "not ready",
67
+ "pending_experiments": ["screen_name_1"],
68
+ "timestamp": "2026-03-13T21:00:00"
69
+ }
70
+ ```
71
+
72
+ **Write this file at the end of every Phase E** (after documenting the round). Overwrite each time — only the latest state matters.
73
+
74
+ **On completion** (positive assessment or max rounds), set `"status": "completed"` so future invocations don't accidentally resume a finished loop.
75
+
76
+ ## Workflow
77
+
78
+ ### Initialization
79
+
80
+ 1. **Check for `review-stage/REVIEW_STATE.json`** *(fall back to `./REVIEW_STATE.json` if not found — legacy path)*:
81
+ - If neither path exists: **fresh start** (normal case)
82
+ - If it exists AND `status` is `"completed"`: **fresh start** (previous loop finished normally)
83
+ - If it exists AND `status` is `"in_progress"` AND `timestamp` is older than 24 hours: **fresh start** (stale state from a killed/abandoned run — delete the file and start over)
84
+ - If it exists AND `status` is `"in_progress"` AND `timestamp` is within 24 hours: **resume**
85
+ - Read the state file to recover `round`, `last_score`, `pending_experiments`
86
+ - Read `review-stage/AUTO_REVIEW.md` to restore full context of prior rounds *(fall back to `./AUTO_REVIEW.md`)*
87
+ - If `pending_experiments` is non-empty, check if they have completed (e.g., check screen sessions)
88
+ - Resume from the next round (round = saved round + 1)
89
+ - Log: "Recovered from context compaction. Resuming at Round N."
90
+ 2. Read project narrative documents, memory files, and any prior review documents
91
+ 3. Read recent experiment results (check output directories, logs)
92
+ 4. Identify current weaknesses and open TODOs from prior reviews
93
+ 5. Initialize round counter = 1 (unless recovered from state file)
94
+ 6. Create/update `review-stage/AUTO_REVIEW.md` with header and timestamp
95
+
96
+ ### Loop (repeat up to MAX_ROUNDS)
97
+
98
+ #### Phase A: Review
99
+
100
+ Send comprehensive context to the external reviewer.
101
+
102
+ **Check MCP availability first**, then use appropriate method:
103
+
104
+ **If MCP available (Primary):**
105
+ ```
106
+ Use mcp__minimax-chat__minimax_chat tool with:
107
+ - system: "You are a senior machine learning researcher serving as a reviewer for top-tier conferences like NeurIPS, ICML, and ICLR. Provide rigorous, constructive feedback."
108
+ - prompt: [Full review prompt with context]
109
+ - model: "MiniMax-M3"
110
+ ```
111
+
112
+ **If MCP NOT available (Fallback):**
113
+ ```bash
114
+ curl -s "https://api.minimax.io/v1/chat/completions" \
115
+ -H "Content-Type: application/json" \
116
+ -H "Authorization: Bearer $MINIMAX_API_KEY" \
117
+ -d '{
118
+ "model": "MiniMax-M3",
119
+ "messages": [
120
+ {
121
+ "role": "system",
122
+ "content": "You are a senior machine learning researcher serving as a reviewer for top-tier conferences like NeurIPS, ICML, and ICLR. Provide rigorous, constructive feedback."
123
+ },
124
+ {
125
+ "role": "user",
126
+ "content": "[Round N/MAX_ROUNDS of autonomous review loop]\n\n[Full research context: claims, methods, results, known weaknesses]\n[Changes since last round, if any]\n[For round 2+: Summary of previous review feedback and what was addressed]\n\nPlease act as a senior ML reviewer (NeurIPS/ICML level).\n\n1. Score this work 1-10 for a top venue\n2. List remaining critical weaknesses (ranked by severity)\n3. For each weakness, specify the MINIMUM fix (experiment, analysis, or reframing)\n4. State clearly: is this READY for submission? Yes/No/Almost\n\nBe brutally honest. If the work is ready, say so clearly."
127
+ }
128
+ ],
129
+ "max_tokens": 4096
130
+ }'
131
+ ```
132
+
133
+ **Note**: Each round is a standalone API call. For round 2+, include the summary of previous reviews and changes in the prompt itself.
134
+
135
+ #### Phase B: Parse Assessment
136
+
137
+ **CRITICAL: Save the FULL raw response** from the external reviewer verbatim (store in a variable for Phase E). Do NOT discard or summarize — the raw text is the primary record.
138
+
139
+ Then extract structured fields:
140
+ - **Score** (numeric 1-10)
141
+ - **Verdict** ("ready" / "almost" / "not ready")
142
+ - **Action items** (ranked list of fixes)
143
+
144
+ **STOP CONDITION**: If score >= 6 AND verdict ∈ {"ready", "almost"} (exact match — "not ready" does NOT qualify) → stop loop, document final state.
145
+
146
+ #### Phase C: Implement Fixes (if not stopping)
147
+
148
+ For each action item (highest priority first):
149
+
150
+ 1. **Code changes**: Write/modify experiment scripts, model code, analysis scripts
151
+ 2. **Run experiments**: Deploy to GPU server via SSH + screen/tmux
152
+ 3. **Analysis**: Run evaluation, collect results, update figures/tables
153
+ 4. **Documentation**: Update project notes and review document
154
+
155
+ Prioritization rules:
156
+ - Skip fixes requiring excessive compute (flag for manual follow-up)
157
+ - Skip fixes requiring external data/models not available
158
+ - Prefer reframing/analysis over new experiments when both address the concern
159
+ - Always implement metric additions (cheap, high impact)
160
+
161
+ #### Phase D: Wait for Results
162
+
163
+ If experiments were launched:
164
+ - Monitor remote sessions for completion
165
+ - Collect results from output files and logs
166
+
167
+ #### Phase E: Document Round
168
+
169
+ Append to `review-stage/AUTO_REVIEW.md`:
170
+
171
+ ```markdown
172
+ ## Round N (timestamp)
173
+
174
+ ### Assessment (Summary)
175
+ - Score: X/10
176
+ - Verdict: [ready/almost/not ready]
177
+ - Key criticisms: [bullet list]
178
+
179
+ ### Reviewer Raw Response
180
+
181
+ <details>
182
+ <summary>Click to expand full reviewer response</summary>
183
+
184
+ [Paste the COMPLETE raw response from the external reviewer here — verbatim, unedited.
185
+ This is the authoritative record. Do NOT truncate or paraphrase.]
186
+
187
+ </details>
188
+
189
+ ### Actions Taken
190
+ - [what was implemented/changed]
191
+
192
+ ### Results
193
+ - [experiment outcomes, if any]
194
+
195
+ ### Status
196
+ - [continuing to round N+1 / stopping]
197
+ ```
198
+
199
+ **Write `review-stage/REVIEW_STATE.json`** with current round, score, verdict, and any pending experiments.
200
+
201
+ Increment round counter → back to Phase A.
202
+
203
+ ### Termination
204
+
205
+ When loop ends (positive assessment or max rounds):
206
+
207
+ 1. Update `review-stage/REVIEW_STATE.json` with `"status": "completed"`
208
+ 2. Write final summary to `review-stage/AUTO_REVIEW.md`
209
+ 3. Update project notes with conclusions
210
+ 4. If stopped at max rounds without positive assessment:
211
+ - List remaining blockers
212
+ - Estimate effort needed for each
213
+ - Suggest whether to continue manually or pivot
214
+
215
+ ## Key Rules
216
+
217
+ - **Large file handling**: If the Write tool fails due to file size, immediately retry using Bash (`cat << 'EOF' > file`) to write in chunks. Do NOT ask the user for permission — just do it silently.
218
+
219
+ - Be honest — include negative results and failed experiments
220
+ - Do NOT hide weaknesses to game a positive score
221
+ - Implement fixes BEFORE re-reviewing (don't just promise to fix)
222
+ - If an experiment takes > 30 minutes, launch it and continue with other fixes while waiting
223
+ - Document EVERYTHING — the review log should be self-contained
224
+ - Update project notes after each round, not just at the end
225
+ - For round 2+, always include previous review context in the prompt
226
+ - Prefer MCP tool over curl when available (more reliable)
227
+
228
+ ## Prompt Template for Round 2+
229
+
230
+ **MCP Method (Primary):**
231
+ ```
232
+ mcp__minimax-chat__minimax_chat:
233
+ model: "MiniMax-M3"
234
+ system: "You are a senior machine learning researcher serving as a reviewer for top-tier conferences like NeurIPS, ICML, and ICLR. Provide rigorous, constructive feedback."
235
+ message: |
236
+ [Round N/MAX_ROUNDS of autonomous review loop]
237
+
238
+ ## Previous Review Summary (Round N-1)
239
+ - Previous Score: X/10
240
+ - Previous Verdict: [ready/almost/not ready]
241
+ - Previous Key Weaknesses: [list]
242
+
243
+ ## Changes Since Last Review
244
+ 1. [Action 1]: [result]
245
+ 2. [Action 2]: [result]
246
+ 3. [Action 3]: [result]
247
+
248
+ ## Updated Results
249
+ [paste updated metrics/tables]
250
+
251
+ ## Current Research Context
252
+ [brief summary of claims, methods, current state]
253
+
254
+ Please re-score and re-assess:
255
+ 1. Score this work 1-10 for a top venue
256
+ 2. List remaining critical weaknesses (ranked by severity)
257
+ 3. For each weakness, specify the MINIMUM fix
258
+ 4. State clearly: is this READY for submission? Yes/No/Almost
259
+
260
+ Be brutally honest. If the work is ready, say so clearly.
261
+ ```
262
+
263
+ **curl Fallback:**
264
+ ```bash
265
+ curl -s "https://api.minimax.io/v1/chat/completions" \
266
+ -H "Content-Type: application/json" \
267
+ -H "Authorization: Bearer $MINIMAX_API_KEY" \
268
+ -d '{
269
+ "model": "MiniMax-M3",
270
+ "messages": [
271
+ {
272
+ "role": "system",
273
+ "content": "You are a senior machine learning researcher serving as a reviewer for top-tier conferences like NeurIPS, ICML, and ICLR. Provide rigorous, constructive feedback."
274
+ },
275
+ {
276
+ "role": "user",
277
+ "content": "[Round N/MAX_ROUNDS of autonomous review loop]\n\n## Previous Review Summary (Round N-1)\n- Previous Score: X/10\n- Previous Verdict: [ready/almost/not ready]\n- Previous Key Weaknesses: [list]\n\n## Changes Since Last Review\n1. [Action 1]: [result]\n2. [Action 2]: [result]\n3. [Action 3]: [result]\n\n## Updated Results\n[paste updated metrics/tables]\n\n## Current Research Context\n[brief summary of claims, methods, current state]\n\nPlease re-score and re-assess:\n1. Score this work 1-10 for a top venue\n2. List remaining critical weaknesses (ranked by severity)\n3. For each weakness, specify the MINIMUM fix\n4. State clearly: is this READY for submission? Yes/No/Almost\n\nBe brutally honest. If the work is ready, say so clearly."
278
+ }
279
+ ],
280
+ "max_tokens": 4096
281
+ }'
282
+ ```
283
+
284
+
285
+ ## Output Protocols
286
+
287
+ > Follow these shared protocols for all output files:
288
+ > - **[Output Versioning Protocol](../../shared-references/output-versioning.md)** — write timestamped file first, then copy to fixed name
289
+ > - **[Output Manifest Protocol](../../shared-references/output-manifest.md)** — log every output to MANIFEST.md
290
+ > - **[Output Language Protocol](../../shared-references/output-language.md)** — respect the project's language setting