eduevidence 6.0.0 → 6.2.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (189) hide show
  1. package/CONTRIBUTING.md +105 -0
  2. package/README.md +93 -38
  3. package/README.zh-CN.md +26 -6
  4. package/SKILL.md +11 -2
  5. package/assets/readme/landing-tour.gif +0 -0
  6. package/assets/readme/studio-tour.gif +0 -0
  7. package/bin/eduevidence.js +2 -1
  8. package/docs/architecture.md +319 -43
  9. package/docs/demo-workplace-ai.md +1 -1
  10. package/docs/install-guide.md +1 -1
  11. package/docs/orchestration-role-model.md +1 -1
  12. package/docs/release-closeout/README.md +1 -1
  13. package/docs/sciverse-api.md +125 -0
  14. package/eduevidence_cli.py +10 -0
  15. package/engine/decision_policy.py +96 -0
  16. package/engine/evidence_graph.py +14 -10
  17. package/engine/gaps.py +42 -22
  18. package/engine/ids.py +2 -0
  19. package/engine/library.py +6 -2
  20. package/engine/living.py +34 -4
  21. package/engine/migration.py +88 -3
  22. package/engine/orchestration.py +5 -5
  23. package/engine/paths.py +2 -0
  24. package/engine/pilot.py +34 -32
  25. package/engine/taxonomy.py +211 -0
  26. package/engine/tribunal.py +43 -31
  27. package/engine/versions.py +1 -1
  28. package/examples/ai-coding-assistant-evidence/EduEvidence_Report.html +1360 -146
  29. package/examples/ai-coding-assistant-evidence/artifact_manifest.json +3 -3
  30. package/examples/ai-coding-assistant-evidence/citation_check.json +1 -1
  31. package/examples/ai-coding-assistant-evidence/final_verdict.json +107 -0
  32. package/examples/ai-coding-assistant-evidence/gate_report.json +101 -0
  33. package/examples/ai-coding-assistant-evidence/report_spec.json +23 -12
  34. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_academic.html +447 -127
  35. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_claude.html +447 -127
  36. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_datalab-dark.html +447 -127
  37. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_datalab.html +447 -127
  38. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_presentation.html +447 -127
  39. package/examples/ai-coding-assistant-evidence/reports-5themes/report_academic.html +1360 -146
  40. package/examples/ai-coding-assistant-evidence/reports-5themes/report_claude.html +1360 -146
  41. package/examples/ai-coding-assistant-evidence/reports-5themes/report_datalab-dark.html +1360 -146
  42. package/examples/ai-coding-assistant-evidence/reports-5themes/report_datalab.html +1360 -146
  43. package/examples/ai-coding-assistant-evidence/reports-5themes/report_presentation.html +1360 -146
  44. package/examples/ai-coding-assistant-evidence/result.json +13 -9
  45. package/examples/ai-coding-assistant-evidence/result.zh.json +45 -41
  46. package/examples/ai-coding-assistant-evidence/skeptic.json +72 -0
  47. package/examples/ai-coding-assistant-evidence/verdict.json +6 -2
  48. package/examples/spaced-retrieval-practice/applicability.json +14 -0
  49. package/examples/spaced-retrieval-practice/artifact_manifest.json +15 -0
  50. package/examples/spaced-retrieval-practice/claims.jsonl +3 -0
  51. package/examples/spaced-retrieval-practice/evidence.jsonl +6 -0
  52. package/examples/spaced-retrieval-practice/final_verdict.json +93 -0
  53. package/examples/spaced-retrieval-practice/frame.json +58 -0
  54. package/examples/spaced-retrieval-practice/gate_report.json +101 -0
  55. package/examples/spaced-retrieval-practice/methodology.json +78 -0
  56. package/examples/spaced-retrieval-practice/report_spec.json +212 -0
  57. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_academic.html +2728 -0
  58. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_claude.html +2728 -0
  59. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_datalab-dark.html +2728 -0
  60. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_datalab.html +2728 -0
  61. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_presentation.html +2728 -0
  62. package/examples/spaced-retrieval-practice/reports-5themes/report_academic.html +2728 -0
  63. package/examples/spaced-retrieval-practice/reports-5themes/report_claude.html +2728 -0
  64. package/examples/spaced-retrieval-practice/reports-5themes/report_datalab-dark.html +2728 -0
  65. package/examples/spaced-retrieval-practice/reports-5themes/report_datalab.html +2728 -0
  66. package/examples/spaced-retrieval-practice/reports-5themes/report_presentation.html +2728 -0
  67. package/examples/spaced-retrieval-practice/result.json +942 -0
  68. package/examples/spaced-retrieval-practice/result.zh.json +942 -0
  69. package/examples/spaced-retrieval-practice/skeptic.json +70 -0
  70. package/examples/spaced-retrieval-practice/sources.jsonl +7 -0
  71. package/examples/spaced-retrieval-practice/verdict.json +93 -0
  72. package/examples/workplace-ai-assistant/artifact_manifest.json +15 -0
  73. package/examples/workplace-ai-assistant/claims.jsonl +4 -4
  74. package/examples/workplace-ai-assistant/evidence.jsonl +4 -4
  75. package/examples/workplace-ai-assistant/evidence_graph.json +15 -15
  76. package/examples/workplace-ai-assistant/final_verdict.json +78 -0
  77. package/examples/workplace-ai-assistant/gate_report.json +101 -0
  78. package/examples/workplace-ai-assistant/report_spec.json +209 -40
  79. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_academic.html +435 -105
  80. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_claude.html +435 -105
  81. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_datalab-dark.html +435 -105
  82. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_datalab.html +435 -105
  83. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_presentation.html +435 -105
  84. package/examples/workplace-ai-assistant/reports-5themes/report_academic.html +2814 -0
  85. package/examples/workplace-ai-assistant/reports-5themes/report_claude.html +2814 -0
  86. package/examples/workplace-ai-assistant/reports-5themes/report_datalab-dark.html +2814 -0
  87. package/examples/workplace-ai-assistant/reports-5themes/report_datalab.html +2814 -0
  88. package/examples/workplace-ai-assistant/reports-5themes/report_presentation.html +2814 -0
  89. package/examples/workplace-ai-assistant/result.json +82 -20
  90. package/examples/workplace-ai-assistant/result.zh.json +82 -20
  91. package/examples/workplace-ai-assistant/skeptic.json +72 -0
  92. package/examples/workplace-ai-assistant/verdict.json +36 -10
  93. package/integrations/agent_mcp.py +2 -2
  94. package/package.json +12 -3
  95. package/pyproject.toml +4 -3
  96. package/references/report-copy-style.md +67 -0
  97. package/references/retrieval-compliance.md +75 -0
  98. package/references/retrieval-protocol.md +20 -0
  99. package/retrieval/audit.py +27 -3
  100. package/retrieval/fetch.py +96 -0
  101. package/retrieval/sciverse.py +398 -0
  102. package/retrieval/search.py +47 -7
  103. package/schemas/applicability.schema.json +94 -0
  104. package/schemas/chart-spec.schema.json +10 -3
  105. package/schemas/evidence.schema.json +316 -43
  106. package/schemas/fetch-result.schema.json +2 -1
  107. package/schemas/report-result.schema.json +3 -3
  108. package/schemas/report-spec.schema.json +98 -100
  109. package/schemas/skeptic.schema.json +86 -0
  110. package/schemas/source.schema.json +21 -2
  111. package/schemas/v2/finding.schema.json +5 -1
  112. package/schemas/v2/methodology-audit.schema.json +5 -1
  113. package/schemas/v2/outcome.schema.json +28 -5
  114. package/schemas/v2/study.schema.json +5 -1
  115. package/schemas/vNext/autoevolve-session.schema.json +34 -1
  116. package/schemas/vNext/eval-snapshot.schema.json +77 -1
  117. package/schemas/vNext/execution-plan.schema.json +50 -1
  118. package/schemas/vNext/gap-priority.schema.json +54 -1
  119. package/schemas/vNext/negative-search-record.schema.json +68 -1
  120. package/schemas/vNext/research-iteration.schema.json +87 -1
  121. package/schemas/vNext/research-strategy.schema.json +62 -1
  122. package/schemas/vNext/skill-experiment.schema.json +90 -1
  123. package/schemas/vNext/task-spec.schema.json +156 -1
  124. package/schemas/vNext/worker-result.schema.json +60 -1
  125. package/schemas/verdict.schema.json +164 -28
  126. package/scripts/build_esl_artifacts.py +2 -2
  127. package/scripts/build_report_variants.py +18 -2
  128. package/scripts/build_result.py +74 -9
  129. package/scripts/check_package_parity.py +85 -0
  130. package/scripts/check_protocol_alignment.py +375 -0
  131. package/scripts/check_versioned_schemas.py +254 -0
  132. package/scripts/claim_audit.py +13 -8
  133. package/scripts/compute_confidence.py +10 -0
  134. package/scripts/did_regression.py +12 -2
  135. package/scripts/evidence_score.py +5 -2
  136. package/scripts/generate_new_projects.py +4 -4
  137. package/scripts/orchestrator.py +120 -24
  138. package/scripts/pre_verdict_gate.py +224 -26
  139. package/scripts/quickstart.py +18 -2
  140. package/scripts/run_workspace.py +7 -1
  141. package/scripts/skill_payload.py +4 -1
  142. package/scripts/test_adversarial_empirical.py +26 -19
  143. package/scripts/validate_schema.py +31 -1
  144. package/skill/agents/evaluation-designer.md +20 -4
  145. package/skill/agents/evidence-analyst.md +19 -3
  146. package/skill/agents/evidence-judge.md +50 -2
  147. package/skill/agents/evidence-retriever.md +20 -3
  148. package/skill/agents/intervention-designer.md +20 -4
  149. package/skill/agents/method-reviewer.md +18 -2
  150. package/skill/agents/{education-planner.md → research-planner.md} +19 -3
  151. package/skill/agents/skeptic.md +18 -2
  152. package/skill/roles/registry.yaml +11 -11
  153. package/skill/sub-skills/aihot-trend-analysis/SKILL.md +28 -9
  154. package/skill/sub-skills/contradiction-analysis/SKILL.md +31 -11
  155. package/skill/sub-skills/data-analysis/SKILL.md +34 -15
  156. package/skill/sub-skills/ethics-review/SKILL.md +33 -10
  157. package/skill/sub-skills/evidence-extraction/SKILL.md +29 -11
  158. package/skill/sub-skills/evidence-review/SKILL.md +31 -12
  159. package/skill/sub-skills/gap-analysis/SKILL.md +31 -9
  160. package/skill/sub-skills/literature-review/SKILL.md +35 -14
  161. package/skill/sub-skills/methodology-audit/SKILL.md +29 -12
  162. package/skill/sub-skills/report-generation/SKILL.md +28 -0
  163. package/skill/sub-skills/research-planning/SKILL.md +41 -14
  164. package/skill/sub-skills/study-design/SKILL.md +30 -9
  165. package/skill/task-briefs/adjudicate.md +32 -7
  166. package/skill/task-briefs/applicability.md +37 -2
  167. package/skill/task-briefs/audit.md +32 -7
  168. package/skill/task-briefs/challenge.md +34 -5
  169. package/skill/task-briefs/evaluate.md +30 -5
  170. package/skill/task-briefs/extract.md +31 -8
  171. package/skill/task-briefs/frame.md +39 -10
  172. package/skill/task-briefs/intervene.md +32 -6
  173. package/skill/task-briefs/present.md +32 -8
  174. package/skill/task-briefs/projection.md +36 -2
  175. package/skill/task-briefs/retrieve.md +36 -6
  176. package/skill/workflows/decision-and-pilot.md +76 -1
  177. package/skill/workflows/evaluate-and-update.md +83 -0
  178. package/skill/workflows/evidence-review.md +104 -0
  179. package/visualization/eduevidence-report/scripts/build_figures.py +25 -3
  180. package/visualization/eduevidence-report/scripts/build_infographics.py +5 -1
  181. package/visualization/eduevidence-report/scripts/build_report.py +512 -65
  182. package/visualization/eduevidence-report/scripts/charts_data.py +2 -0
  183. package/visualization/eduevidence-report/scripts/lieflat_engine.py +349 -38
  184. package/visualization/eduevidence-report/scripts/zh_labels.py +80 -1
  185. package/web/architecture.html +14885 -0
  186. package/web/studio/assets/index-B8tkF44Q.css +1 -0
  187. package/web/studio/index.html +2 -2
  188. package/web/studio/assets/index-CzXocaGv.css +0 -1
  189. /package/web/studio/assets/{index-pa7jD7n4.js → index-CQ6Keoyc.js} +0 -0
@@ -1,8 +1,8 @@
1
1
  {
2
- "result_sha256": "b498e00a88080c0645a9f0c2f4f513dc916f1ec4a2c6c69592363c99adb5dc17",
3
- "result_zh_sha256": "0b14bf748bf92008a7813cc1aa7db3e30b5ad726e85b5f408b2cebc2dcb94b1c",
2
+ "result_sha256": "3bce7a28f588643a3064c3d82eb749e3999cb18cb812b85c9ad6ecc2a731c6a6",
3
+ "result_zh_sha256": "4d58876e31d6f095e733c674217bbdaf2204dd66097fe5e5ccf4545a397583f8",
4
4
  "renderer_version": "1.0.0",
5
- "git_commit": "0db3403087cf3708bfb7fb0597fa8a289d4fd382",
5
+ "git_commit": "56f5dc3edc6e805c2614767208a83009b0f0a097",
6
6
  "evidence_count": 12,
7
7
  "source_count": 8,
8
8
  "themes": [
@@ -1,5 +1,5 @@
1
1
  {
2
- "generated_at": "2026-08-24T12:39:14+0800",
2
+ "generated_at": "2026-09-13T11:17:51+0800",
3
3
  "checked": 8,
4
4
  "bad": 0,
5
5
  "results": [
@@ -0,0 +1,107 @@
1
+ {
2
+ "decision_question": "大一 C 语言课程是否应该允许学生使用生成式 AI 编程助手?",
3
+ "target_population": "university first-year computer science students learning C programming for the first time",
4
+ "target_context": "16-week lecture-lab course, 60 students, TA support, offline",
5
+ "supported_claims": [
6
+ "AI coding assistants reliably increase task performance during training (completion speed, correctness) — E-001, E-006.",
7
+ "Unguarded generative AI access can harm independent problem solving when access is removed — E-004.",
8
+ "Guardrail design (hints instead of answers) substantially mitigates the negative learning effect — E-005.",
9
+ "Task performance gains do not automatically imply learning gains — E-004 vs E-006 (within-study contrast).",
10
+ "Tool capability is substantial: Codex solves roughly half to three-quarters of CS1 exam-style questions — E-010.",
11
+ "Professional-developer RCT shows ~55% faster task completion with Copilot; directness limited by professional population — E-008.",
12
+ "LLM code explanations rate comparable to student-authored explanations, viable as scaffold material — E-011."
13
+ ],
14
+ "uncertain_claims": [
15
+ "Whether AI coding assistants improve or preserve actual programming learning in university novices — no direct university-level RCT in reviewed set [无直接证据]",
16
+ "Whether one-week neutral retention (Kazemitabaar 2023) extends to a semester — E-003.",
17
+ "Whether benchmark quality findings (E-009) and explanation-quality ratings (E-011) translate into classroom learning gains.",
18
+ "How comprehension/ownership difficulties documented in usability studies (E-012) behave over a full semester with guardrails."
19
+ ],
20
+ "contradicted_claims": [
21
+ "The claim 'AI tools always improve learning' is contradicted by E-004 (unguarded access, -17% independent exam).",
22
+ "The claim 'speed gains equal learning gains' is contradicted by the task-vs-learning separation across E-001/E-006/E-008 vs E-004."
23
+ ],
24
+ "reason_for_disagreement": "Disagreement comes from outcome separation (task vs learning), tool design (guarded vs unguarded), and population (K-12 / professionals vs university). Task-performance evidence is consistently positive across randomized and benchmark studies; the only study measuring independent performance after AI removal shows harm without guardrails; usability and artifact studies add dependence and quality caveats rather than resolving the learning question.",
25
+ "methodology_summary": "Eight real sources: three randomized experiments (Kazemitabaar 2023 n=69 K-12; Bastani 2025 n≈950 high-school mathematics; Peng 2023 n=95 professional developers, preprint), one ESL writing mixed-methods study (Marzuki 2024), plus benchmark/capability/usability studies (Yetistiren 2023; Finnie-Ansley 2022; explanation-comparison 2023; Vaithilingam 2022). No direct RCT in university programming courses. Internal validity of the core RCTs is strong; directness to first-year university C programming is weak. All sources carry registry-verified DOIs (see benchmarks/doi-audit/report.md).",
26
+ "outcome_specific_findings": {
27
+ "completion_time": "positive during training and professional tasks (E-001, E-008)",
28
+ "independent_problem_solving": "neutral-to-negative without guardrails (E-002, E-004)",
29
+ "retention": "neutral over 1 week (E-003)",
30
+ "assignment_score": "positive during practice, negative on closed-book exam (E-004, E-006); tool itself scores passing-level on CS1 questions (E-010)",
31
+ "code_quality": "mixed on benchmarks; security concerns documented (E-009)",
32
+ "metacognition": "LLM explanations compare well (E-011) while novice ownership/debugging difficulties persist (E-012)",
33
+ "ai_dependency": "documented crutch behavior with unguarded tool (E-004, E-005, E-012)"
34
+ },
35
+ "short_term_effect": "Task performance reliably increases; learning effect null-to-negative without guardrails.",
36
+ "long_term_effect": "No evidence beyond one week; long-term learning effect unknown.",
37
+ "transfer_effect": "No full transfer evidence; manual code-modification not harmed in one small study (E-002).",
38
+ "risk_effect": "AI dependency and over-reliance risk is real and documented for unguarded usage (E-004) and foreshadowed by usability findings (E-012).",
39
+ "applicability": {
40
+ "suitable_for": "pilot in first-year C course with guardrailed usage policy",
41
+ "not_suitable_for": "unrestricted AI adoption without usage policy",
42
+ "required_conditions": [
43
+ "guardrailed AI usage policy (hints not answers, modeled on GPT Tutor arm)",
44
+ "no-AI transfer assessment",
45
+ "TA support"
46
+ ]
47
+ },
48
+ "confidence": "Moderate",
49
+ "confidence_breakdown": {
50
+ "score": 0.586,
51
+ "evidence_quality": 0.758,
52
+ "consistency": 0.667,
53
+ "directness": 0.458,
54
+ "evidence_count": 12,
55
+ "independent_studies": 8,
56
+ "independent_samples": 8,
57
+ "count_term": 1.0,
58
+ "conflict_penalty": 0.15,
59
+ "unsupported_penalty": 0.0
60
+ },
61
+ "what_can_be_claimed": [
62
+ "AI coding assistants raise task performance for novices during training.",
63
+ "Unguarded access carries a real risk of hurting independent problem solving.",
64
+ "Guardrail design can mitigate that risk.",
65
+ "Direct evidence for university C programming learning is missing.",
66
+ "Tool capability headroom is large (CS1 question pass rates; professional speed RCT)."
67
+ ],
68
+ "what_cannot_be_claimed": [
69
+ "AI coding assistants improve (or even preserve) university students' programming learning.",
70
+ "Any long-term or retention benefit.",
71
+ "Any claim about which students benefit, based on university samples.",
72
+ "That benchmark or usability findings substitute for classroom learning outcomes."
73
+ ],
74
+ "missing_evidence": [
75
+ "RCT of AI coding assistants in university programming courses with retention and no-AI transfer tests.",
76
+ "Studies varying AI usage policy within the same course.",
77
+ "Longitudinal data on AI dependency beyond one course.",
78
+ "Peer-reviewed replication of the professional speed RCT (Peng et al. remains a preprint)."
79
+ ],
80
+ "recommended_action": "pilot",
81
+ "decision_rationale": "Positive task-performance evidence plus documented unguarded-access risk, mixed quality/usability signals, and missing university-level learning evidence → bounded, guardrailed pilot with evaluation, not full adoption.",
82
+ "exceeds_evidence_boundary": [
83
+ "Claiming 'AI coding assistants improve learning' exceeds the boundary: direct learning-effect evidence is missing.",
84
+ "Claiming 'AI works for everyone' exceeds the boundary: population and subject mismatch."
85
+ ],
86
+ "confidence_score": 0.586,
87
+ "confidence_policy_version": "2026-08-12.v3",
88
+ "independent_studies": 8,
89
+ "independent_samples": 8,
90
+ "raw_model_confidence": "Moderate",
91
+ "raw_model_confidence_breakdown": {
92
+ "score": 0.5,
93
+ "evidence_quality": 0.7,
94
+ "consistency": 0.6,
95
+ "directness": 0.4,
96
+ "evidence_count": 12,
97
+ "independent_studies": 8,
98
+ "independent_samples": 8,
99
+ "count_term": 1.0,
100
+ "conflict_penalty": 0.0,
101
+ "unsupported_penalty": 0.0
102
+ },
103
+ "strongest_support": "AI coding assistants reliably speed up practice work: completion rate 1.15x and time 0.57x in a randomised trial of 69 novices.",
104
+ "key_uncertainty": "No university-level RCT measures learning directly, and the one large trial that did - unguarded GPT-4 - saw independent exam scores fall 17%.",
105
+ "main_risk": "Unguarded access can raise practice performance while lowering independent exam performance, and learners may not notice the gap.",
106
+ "next_action": "Run a phased CS1 pilot with hints-not-answers guardrails, weekly lab use, and a no-AI transfer exam that can stop the pilot."
107
+ }
@@ -0,0 +1,101 @@
1
+ {
2
+ "gate_version": "2026-08-13.v1",
3
+ "checked_at": "2026-09-13T03:17:13.667942+00:00",
4
+ "workspace": "examples/ai-coding-assistant-evidence",
5
+ "items": {
6
+ "research_frame_valid": {
7
+ "title": "Research Frame valid",
8
+ "status": "pass",
9
+ "detail": "frame.json valid (question=我们准备在大学一年级 C 语言课程中允许学生使用生成式 AI 编程助手。它到底会不会提高学习效果?应该怎样引入?)",
10
+ "critical": true,
11
+ "blocks_high": true
12
+ },
13
+ "sources_valid": {
14
+ "title": "Sources valid",
15
+ "status": "pass",
16
+ "detail": "8 source record(s) schema-valid",
17
+ "critical": true,
18
+ "blocks_high": true
19
+ },
20
+ "evidence_schema_valid": {
21
+ "title": "Evidence Schema valid",
22
+ "status": "pass",
23
+ "detail": "12 evidence record(s) schema-valid",
24
+ "critical": true,
25
+ "blocks_high": true
26
+ },
27
+ "source_dedupe": {
28
+ "title": "Source dedupe",
29
+ "status": "pass",
30
+ "detail": "8 unique source(s), no duplicates",
31
+ "critical": true,
32
+ "blocks_high": true
33
+ },
34
+ "counter_evidence_search": {
35
+ "title": "Counter-evidence search",
36
+ "status": "pass",
37
+ "detail": "search_performed=true; 9/9 checks run; findings=9; contradictory_evidence_found=True",
38
+ "critical": true,
39
+ "blocks_high": true
40
+ },
41
+ "methodology_audit": {
42
+ "title": "Methodology audit",
43
+ "status": "pass",
44
+ "detail": "methodology verdict=CONCERN; task/learning separated",
45
+ "critical": true,
46
+ "blocks_high": true
47
+ },
48
+ "claim_evidence_audit": {
49
+ "title": "Claim-Evidence Audit",
50
+ "status": "warn",
51
+ "detail": "supported claim cites contradicting evidence E-004 (may be an intentional negative finding); supported claim cites contradicting evidence E-004 (may be an intentional negative finding) (+3 more)",
52
+ "critical": true,
53
+ "blocks_high": false
54
+ },
55
+ "outcome_mapping": {
56
+ "title": "Outcome mapping",
57
+ "status": "warn",
58
+ "detail": "outcome keys known; frame-declared outcomes without evidence: ai_dependency, reduced_transfer",
59
+ "critical": false,
60
+ "blocks_high": false
61
+ },
62
+ "scope_calibration": {
63
+ "title": "Scope calibration",
64
+ "status": "pass",
65
+ "detail": "claims bounded: can=5, cannot=4, exceeds_boundary=2",
66
+ "critical": false,
67
+ "blocks_high": true
68
+ },
69
+ "independent_study_count": {
70
+ "title": "Independent study-sample count",
71
+ "status": "pass",
72
+ "detail": "independent studies=8, samples=8",
73
+ "critical": true,
74
+ "blocks_high": true
75
+ },
76
+ "deterministic_confidence": {
77
+ "title": "Deterministic confidence",
78
+ "status": "pass",
79
+ "detail": "deterministic confidence=Moderate (score=0.586, policy=2026-08-12.v3)",
80
+ "critical": true,
81
+ "blocks_high": true
82
+ },
83
+ "decision_action_consistency": {
84
+ "title": "Decision action consistency",
85
+ "status": "pass",
86
+ "detail": "action=pilot is within the conservative bound",
87
+ "critical": true,
88
+ "blocks_high": false
89
+ }
90
+ },
91
+ "passed": true,
92
+ "critical_failures": [],
93
+ "high_confidence_allowed": true,
94
+ "max_confidence": "High",
95
+ "enforcement": {
96
+ "rule": "confidence capped at {max}; High confidence requires a fully passing gate and >= 2 independent studies",
97
+ "max_confidence": "High",
98
+ "requires_action_change": false,
99
+ "action_override": null
100
+ }
101
+ }
@@ -108,30 +108,41 @@
108
108
  "catalog_ref": "L9 Bubble Almanac",
109
109
  "source": "evidence.year_x_dimension"
110
110
  },
111
+ {
112
+ "chart_id": "lieflat-matrix-heat.svg",
113
+ "type": "matrix_heat",
114
+ "catalog_ref": "L16 Matrix Heat",
115
+ "source": "evidence.year_x_outcome_counts"
116
+ },
111
117
  {
112
118
  "chart_id": "lieflat-tick-rows.svg",
113
119
  "type": "tick_rows",
114
120
  "catalog_ref": "F5 Tick Rows",
115
121
  "source": "outcomes.direction_counts"
116
- }
117
- ],
118
- "suppressed": [
122
+ },
123
+ {
124
+ "chart_id": "lieflat-paired-rungs.svg",
125
+ "type": "paired_rungs",
126
+ "catalog_ref": "F6 Paired Rungs",
127
+ "source": "outcomes.paired_counts"
128
+ },
119
129
  {
120
- "chart_id": "lieflat-forest-plot.svg",
121
- "type": "forest_plot",
122
- "catalog_ref": "FOREST-PLOT (publication figure)",
123
- "reason": "fewer than 3 studies with numeric effect size (got 0)"
130
+ "chart_id": "lieflat-brand-spectrum.svg",
131
+ "type": "brand_spectrum",
132
+ "catalog_ref": "L7 Brand Spectrum",
133
+ "source": "outcomes.bipolar_axes"
124
134
  },
125
135
  {
126
- "chart_id": "lieflat-dot-cascade.svg",
127
- "type": "dot_cascade",
128
- "catalog_ref": "L2 Dot Cascade",
129
- "reason": "fewer than 3 studies with numeric effect size (got 0)"
136
+ "chart_id": "lieflat-hundred-field.svg",
137
+ "type": "hundred_field",
138
+ "catalog_ref": "L14 Hundred Field",
139
+ "source": "evidence.study_type_composition"
130
140
  }
131
141
  ],
142
+ "suppressed": [],
132
143
  "rejected": [],
133
144
  "warnings": [
134
- "visual_layout missing or fully invalid — using deterministic safe combination (forest_plot + dot_cascade + bubble_almanac + tick_rows)"
145
+ "visual_layout missing or fully invalid — data-driven fallback selected bubble_almanac, matrix_heat, tick_rows, paired_rungs, brand_spectrum, hundred_field"
135
146
  ]
136
147
  },
137
148
  "charts": [