eduevidence 6.0.0 → 6.3.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (267) hide show
  1. package/CHANGELOG.md +395 -0
  2. package/CONTRIBUTING.md +105 -0
  3. package/README.md +113 -49
  4. package/README.zh-CN.md +39 -12
  5. package/SKILL.md +15 -5
  6. package/assets/readme/landing-tour.gif +0 -0
  7. package/assets/readme/studio-tour.gif +0 -0
  8. package/benchmarks/evidence-library.json +277 -1
  9. package/bin/eduevidence.js +2 -1
  10. package/docs/architecture.md +325 -46
  11. package/docs/demo-workplace-ai.md +1 -1
  12. package/docs/install-guide.md +1 -1
  13. package/docs/j-ev-experimental.md +250 -0
  14. package/docs/orchestration-role-model.md +1 -1
  15. package/docs/release-closeout/README.md +1 -1
  16. package/docs/reproducibility.md +138 -0
  17. package/docs/sciverse-api.md +125 -0
  18. package/domains/_neutral/copy/few_shots.json +21 -0
  19. package/domains/_neutral/copy/framing_lexicon.json +19 -0
  20. package/domains/_neutral/copy/module_labels.json +5 -0
  21. package/domains/_neutral/copy/module_labels_footer.json +102 -0
  22. package/domains/_neutral/copy/module_labels_modules.json +204 -0
  23. package/domains/_neutral/copy/module_labels_nav.json +126 -0
  24. package/domains/_neutral/copy/module_labels_summary.json +98 -0
  25. package/domains/_neutral/copy/module_labels_tables.json +164 -0
  26. package/domains/_neutral/copy/module_labels_v2.json +90 -0
  27. package/domains/_neutral/copy/risk_constructs.json +20 -0
  28. package/domains/_neutral/copy/section_titles.json +66 -0
  29. package/domains/_neutral/copy/terminology.json +11 -0
  30. package/domains/check_copy_packs.py +103 -0
  31. package/domains/education/copy/few_shots.json +22 -0
  32. package/domains/education/copy/framing_enums.json +167 -0
  33. package/domains/education/copy/framing_lexicon.json +166 -0
  34. package/domains/education/copy/module_labels.json +169 -0
  35. package/domains/education/copy/risk_constructs.json +48 -0
  36. package/domains/education/copy/section_titles.json +186 -0
  37. package/domains/education/copy/terminology.json +70 -0
  38. package/domains/education/manifest.json +1 -1
  39. package/domains/education/outcome_taxonomy.json +2 -2
  40. package/domains/manifest.json +1 -1
  41. package/domains/policy/copy/few_shots.json +22 -0
  42. package/domains/policy/copy/framing_enums.json +94 -0
  43. package/domains/policy/copy/framing_lexicon.json +174 -0
  44. package/domains/policy/copy/module_labels.json +168 -0
  45. package/domains/policy/copy/risk_constructs.json +33 -0
  46. package/domains/policy/copy/section_titles.json +186 -0
  47. package/domains/policy/copy/terminology.json +64 -0
  48. package/eduevidence_cli.py +10 -0
  49. package/engine/capabilities.py +57 -5
  50. package/engine/decision_policy.py +167 -0
  51. package/engine/evidence_graph.py +14 -10
  52. package/engine/gaps.py +42 -22
  53. package/engine/ids.py +2 -0
  54. package/engine/library.py +6 -2
  55. package/engine/library_builtin.py +7 -4
  56. package/engine/living.py +34 -4
  57. package/engine/migration.py +88 -3
  58. package/engine/orchestration.py +5 -5
  59. package/engine/paths.py +2 -0
  60. package/engine/pilot.py +34 -32
  61. package/engine/taxonomy.py +211 -0
  62. package/engine/tribunal.py +49 -43
  63. package/engine/versions.py +1 -1
  64. package/examples/ai-coding-assistant-evidence/EduEvidence_Report.html +1361 -147
  65. package/examples/ai-coding-assistant-evidence/artifact_manifest.json +3 -3
  66. package/examples/ai-coding-assistant-evidence/citation_check.json +1 -1
  67. package/examples/ai-coding-assistant-evidence/final_verdict.json +107 -0
  68. package/examples/ai-coding-assistant-evidence/gate_report.json +101 -0
  69. package/examples/ai-coding-assistant-evidence/report_spec.json +23 -12
  70. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_academic.html +448 -128
  71. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_claude.html +448 -128
  72. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_datalab-dark.html +448 -128
  73. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_datalab.html +448 -128
  74. package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_presentation.html +448 -128
  75. package/examples/ai-coding-assistant-evidence/reports-5themes/report_academic.html +1360 -146
  76. package/examples/ai-coding-assistant-evidence/reports-5themes/report_claude.html +1360 -146
  77. package/examples/ai-coding-assistant-evidence/reports-5themes/report_datalab-dark.html +1360 -146
  78. package/examples/ai-coding-assistant-evidence/reports-5themes/report_datalab.html +1360 -146
  79. package/examples/ai-coding-assistant-evidence/reports-5themes/report_presentation.html +1360 -146
  80. package/examples/ai-coding-assistant-evidence/result.json +13 -9
  81. package/examples/ai-coding-assistant-evidence/result.zh.json +45 -41
  82. package/examples/ai-coding-assistant-evidence/skeptic.json +72 -0
  83. package/examples/ai-coding-assistant-evidence/verdict.json +6 -2
  84. package/examples/spaced-retrieval-practice/EduEvidence_Report.html +2728 -0
  85. package/examples/spaced-retrieval-practice/applicability.json +14 -0
  86. package/examples/spaced-retrieval-practice/artifact_manifest.json +15 -0
  87. package/examples/spaced-retrieval-practice/claims.jsonl +3 -0
  88. package/examples/spaced-retrieval-practice/evidence.jsonl +6 -0
  89. package/examples/spaced-retrieval-practice/final_verdict.json +93 -0
  90. package/examples/spaced-retrieval-practice/frame.json +58 -0
  91. package/examples/spaced-retrieval-practice/gate_report.json +101 -0
  92. package/examples/spaced-retrieval-practice/methodology.json +78 -0
  93. package/examples/spaced-retrieval-practice/report.html +2522 -0
  94. package/examples/spaced-retrieval-practice/report_spec.json +212 -0
  95. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_academic.html +2728 -0
  96. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_claude.html +2728 -0
  97. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_datalab-dark.html +2728 -0
  98. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_datalab.html +2728 -0
  99. package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_presentation.html +2728 -0
  100. package/examples/spaced-retrieval-practice/reports-5themes/report_academic.html +2728 -0
  101. package/examples/spaced-retrieval-practice/reports-5themes/report_claude.html +2728 -0
  102. package/examples/spaced-retrieval-practice/reports-5themes/report_datalab-dark.html +2728 -0
  103. package/examples/spaced-retrieval-practice/reports-5themes/report_datalab.html +2728 -0
  104. package/examples/spaced-retrieval-practice/reports-5themes/report_presentation.html +2728 -0
  105. package/examples/spaced-retrieval-practice/result.json +942 -0
  106. package/examples/spaced-retrieval-practice/result.zh.json +942 -0
  107. package/examples/spaced-retrieval-practice/skeptic.json +70 -0
  108. package/examples/spaced-retrieval-practice/sources.jsonl +7 -0
  109. package/examples/spaced-retrieval-practice/verdict.json +93 -0
  110. package/examples/workplace-ai-assistant/EduEvidence_Report.html +2814 -0
  111. package/examples/workplace-ai-assistant/artifact_manifest.json +15 -0
  112. package/examples/workplace-ai-assistant/claims.jsonl +4 -4
  113. package/examples/workplace-ai-assistant/evidence.jsonl +4 -4
  114. package/examples/workplace-ai-assistant/evidence_graph.json +15 -15
  115. package/examples/workplace-ai-assistant/final_verdict.json +78 -0
  116. package/examples/workplace-ai-assistant/gate_report.json +101 -0
  117. package/examples/workplace-ai-assistant/report_spec.json +209 -40
  118. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_academic.html +449 -119
  119. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_claude.html +449 -119
  120. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_datalab-dark.html +449 -119
  121. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_datalab.html +449 -119
  122. package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_presentation.html +449 -119
  123. package/examples/workplace-ai-assistant/reports-5themes/report_academic.html +2814 -0
  124. package/examples/workplace-ai-assistant/reports-5themes/report_claude.html +2814 -0
  125. package/examples/workplace-ai-assistant/reports-5themes/report_datalab-dark.html +2814 -0
  126. package/examples/workplace-ai-assistant/reports-5themes/report_datalab.html +2814 -0
  127. package/examples/workplace-ai-assistant/reports-5themes/report_presentation.html +2814 -0
  128. package/examples/workplace-ai-assistant/result.json +82 -20
  129. package/examples/workplace-ai-assistant/result.zh.json +82 -20
  130. package/examples/workplace-ai-assistant/skeptic.json +72 -0
  131. package/examples/workplace-ai-assistant/verdict.json +36 -10
  132. package/integrations/agent_mcp.py +2 -2
  133. package/integrations/jev/__init__.py +115 -0
  134. package/integrations/jev/approval.py +212 -0
  135. package/integrations/jev/cli.py +84 -0
  136. package/integrations/jev/config.py +112 -0
  137. package/integrations/jev/gateway.py +128 -0
  138. package/integrations/jev/modes.py +38 -0
  139. package/integrations/jev/tools_classify.py +88 -0
  140. package/integrations/jev/tools_extract.py +111 -0
  141. package/integrations/jev/tools_rerank.py +71 -0
  142. package/integrations/jev/tools_screen.py +87 -0
  143. package/integrations/jev/tools_verify.py +95 -0
  144. package/integrations/jev_mcp.py +22 -0
  145. package/integrations/semantic_decide.py +286 -0
  146. package/integrations/semdecide_cli.py +55 -0
  147. package/package.json +19 -2
  148. package/pyproject.toml +4 -3
  149. package/references/report-copy-style.md +107 -0
  150. package/references/retrieval-compliance.md +75 -0
  151. package/references/retrieval-protocol.md +20 -0
  152. package/retrieval/audit.py +27 -3
  153. package/retrieval/fetch.py +96 -0
  154. package/retrieval/sciverse.py +398 -0
  155. package/retrieval/search.py +47 -7
  156. package/schemas/applicability.schema.json +94 -0
  157. package/schemas/chart-spec.schema.json +10 -3
  158. package/schemas/evidence.schema.json +316 -43
  159. package/schemas/fetch-result.schema.json +2 -1
  160. package/schemas/report-result.schema.json +3 -3
  161. package/schemas/report-spec.schema.json +98 -100
  162. package/schemas/skeptic.schema.json +86 -0
  163. package/schemas/source.schema.json +21 -2
  164. package/schemas/v2/decision-snapshot.schema.json +20 -9
  165. package/schemas/v2/finding.schema.json +5 -1
  166. package/schemas/v2/intake.schema.json +191 -0
  167. package/schemas/v2/methodology-audit.schema.json +5 -1
  168. package/schemas/v2/outcome.schema.json +28 -5
  169. package/schemas/v2/study.schema.json +5 -1
  170. package/schemas/vNext/autoevolve-session.schema.json +34 -1
  171. package/schemas/vNext/eval-snapshot.schema.json +77 -1
  172. package/schemas/vNext/execution-plan.schema.json +50 -1
  173. package/schemas/vNext/gap-priority.schema.json +54 -1
  174. package/schemas/vNext/negative-search-record.schema.json +68 -1
  175. package/schemas/vNext/research-iteration.schema.json +87 -1
  176. package/schemas/vNext/research-strategy.schema.json +62 -1
  177. package/schemas/vNext/skill-experiment.schema.json +90 -1
  178. package/schemas/vNext/task-spec.schema.json +156 -1
  179. package/schemas/vNext/worker-result.schema.json +60 -1
  180. package/schemas/verdict.schema.json +164 -28
  181. package/scripts/build_evidence_library.py +15 -5
  182. package/scripts/build_report_variants.py +18 -2
  183. package/scripts/build_result.py +74 -9
  184. package/scripts/check_package_parity.py +85 -0
  185. package/scripts/check_protocol_alignment.py +375 -0
  186. package/scripts/check_versioned_schemas.py +254 -0
  187. package/scripts/claim_audit.py +13 -8
  188. package/scripts/compute_confidence.py +10 -0
  189. package/scripts/dashboard_server.py +13 -2
  190. package/scripts/did_regression.py +12 -2
  191. package/scripts/evidence_score.py +5 -2
  192. package/scripts/intake/__init__.py +31 -0
  193. package/scripts/intake/__main__.py +18 -0
  194. package/scripts/intake/background.py +78 -0
  195. package/scripts/intake/browser.py +79 -0
  196. package/scripts/intake/cli.py +57 -0
  197. package/scripts/intake/constants.py +57 -0
  198. package/scripts/intake/depth.py +53 -0
  199. package/scripts/intake/enhancements.py +106 -0
  200. package/scripts/intake/hooks.py +90 -0
  201. package/scripts/intake/prefs.py +76 -0
  202. package/scripts/intake/prompts.py +85 -0
  203. package/scripts/intake/session.py +152 -0
  204. package/scripts/lint_file_layers.py +126 -0
  205. package/scripts/orchestrator.py +187 -40
  206. package/scripts/pre_verdict_gate.py +241 -29
  207. package/scripts/quickstart.py +18 -2
  208. package/scripts/run_workspace.py +7 -1
  209. package/scripts/skill_lint.py +11 -1
  210. package/scripts/skill_payload.py +6 -3
  211. package/scripts/test_adversarial_empirical.py +96 -25
  212. package/scripts/validate_schema.py +31 -1
  213. package/skill/agents/evaluation-designer.md +20 -4
  214. package/skill/agents/evidence-analyst.md +19 -3
  215. package/skill/agents/evidence-judge.md +98 -8
  216. package/skill/agents/evidence-retriever.md +20 -3
  217. package/skill/agents/intervention-designer.md +20 -4
  218. package/skill/agents/method-reviewer.md +18 -2
  219. package/skill/agents/{education-planner.md → research-planner.md} +19 -3
  220. package/skill/agents/skeptic.md +18 -2
  221. package/skill/roles/registry.yaml +11 -11
  222. package/skill/sub-skills/aihot-trend-analysis/SKILL.md +28 -9
  223. package/skill/sub-skills/contradiction-analysis/SKILL.md +31 -11
  224. package/skill/sub-skills/data-analysis/SKILL.md +34 -15
  225. package/skill/sub-skills/ethics-review/SKILL.md +33 -10
  226. package/skill/sub-skills/evidence-extraction/SKILL.md +29 -11
  227. package/skill/sub-skills/evidence-review/SKILL.md +31 -12
  228. package/skill/sub-skills/gap-analysis/SKILL.md +31 -9
  229. package/skill/sub-skills/literature-review/SKILL.md +35 -14
  230. package/skill/sub-skills/methodology-audit/SKILL.md +29 -12
  231. package/skill/sub-skills/report-generation/SKILL.md +28 -0
  232. package/skill/sub-skills/research-planning/SKILL.md +41 -14
  233. package/skill/sub-skills/study-design/SKILL.md +30 -9
  234. package/skill/task-briefs/adjudicate.md +32 -7
  235. package/skill/task-briefs/applicability.md +37 -2
  236. package/skill/task-briefs/audit.md +32 -7
  237. package/skill/task-briefs/challenge.md +34 -5
  238. package/skill/task-briefs/evaluate.md +30 -5
  239. package/skill/task-briefs/extract.md +31 -8
  240. package/skill/task-briefs/frame.md +39 -10
  241. package/skill/task-briefs/intervene.md +32 -6
  242. package/skill/task-briefs/present.md +32 -8
  243. package/skill/task-briefs/projection.md +36 -2
  244. package/skill/task-briefs/retrieve.md +36 -6
  245. package/skill/workflows/decision-and-pilot.md +76 -1
  246. package/skill/workflows/evaluate-and-update.md +83 -0
  247. package/skill/workflows/evidence-review.md +104 -0
  248. package/skill/workflows/experimental-jev.md +170 -0
  249. package/skill/workflows/intake.md +120 -0
  250. package/visualization/eduevidence-report/scripts/build_figures.py +25 -3
  251. package/visualization/eduevidence-report/scripts/build_infographics.py +37 -15
  252. package/visualization/eduevidence-report/scripts/build_report.py +435 -575
  253. package/visualization/eduevidence-report/scripts/charts_data.py +2 -0
  254. package/visualization/eduevidence-report/scripts/lieflat_engine.py +349 -38
  255. package/visualization/eduevidence-report/scripts/report_copy_pack.py +296 -0
  256. package/visualization/eduevidence-report/scripts/report_copy_policy_guard.py +47 -0
  257. package/visualization/eduevidence-report/scripts/zh_labels.py +141 -1
  258. package/web/architecture.html +14885 -0
  259. package/web/studio/assets/index-B8tkF44Q.css +1 -0
  260. package/web/studio/index.html +2 -2
  261. package/scripts/build_esl_artifacts.py +0 -1921
  262. package/scripts/build_killer_demo.py +0 -295
  263. package/scripts/enrich_projects_human_and_lieflat.py +0 -315
  264. package/scripts/generate_new_projects.py +0 -686
  265. package/scripts/sync_killer_demo_report.py +0 -270
  266. package/web/studio/assets/index-CzXocaGv.css +0 -1
  267. /package/web/studio/assets/{index-pa7jD7n4.js → index-CQ6Keoyc.js} +0 -0
@@ -4941,6 +4941,282 @@
4941
4941
  "assessment_edtech"
4942
4942
  ]
4943
4943
  },
4944
+ {
4945
+ "entry_id": "lib-ai-coding-assistant-E-001",
4946
+ "source_id": "S-2023-kazemitabaar",
4947
+ "title": "Studying the effect of AI Code Generators on Supporting Novice Learners in Introductory Programming",
4948
+ "year": 2023,
4949
+ "outcome_token": "completion_time",
4950
+ "outcome_tokens": [
4951
+ "completion_time"
4952
+ ],
4953
+ "direction": "support",
4954
+ "study_type": "rct",
4955
+ "claim_text": "AI coding assistants significantly increase task completion speed and completion rate during training.",
4956
+ "effect_summary": "1.15x completion rate, 0.57x time, 1.8x correctness",
4957
+ "confidence_markers": [
4958
+ "evidence_level:strong",
4959
+ "quality_score:9.0",
4960
+ "confidence:0.7",
4961
+ "decision_relation:support_adoption"
4962
+ ],
4963
+ "domains": [
4964
+ "ai-coding-assistant"
4965
+ ]
4966
+ },
4967
+ {
4968
+ "entry_id": "lib-ai-coding-assistant-E-002",
4969
+ "source_id": "S-2023-kazemitabaar",
4970
+ "title": "Studying the effect of AI Code Generators on Supporting Novice Learners in Introductory Programming",
4971
+ "year": 2023,
4972
+ "outcome_token": "independent_problem_solving",
4973
+ "outcome_tokens": [
4974
+ "independent_problem_solving"
4975
+ ],
4976
+ "direction": "neutral",
4977
+ "study_type": "rct",
4978
+ "claim_text": "Access to AI code generation did not decrease performance on manual code-modification tasks.",
4979
+ "effect_summary": "no significant difference between groups",
4980
+ "confidence_markers": [
4981
+ "evidence_level:moderate",
4982
+ "quality_score:7.0",
4983
+ "confidence:0.5",
4984
+ "decision_relation:neutral"
4985
+ ],
4986
+ "domains": [
4987
+ "ai-coding-assistant"
4988
+ ]
4989
+ },
4990
+ {
4991
+ "entry_id": "lib-ai-coding-assistant-E-003",
4992
+ "source_id": "S-2023-kazemitabaar",
4993
+ "title": "Studying the effect of AI Code Generators on Supporting Novice Learners in Introductory Programming",
4994
+ "year": 2023,
4995
+ "outcome_token": "retention",
4996
+ "outcome_tokens": [
4997
+ "retention"
4998
+ ],
4999
+ "direction": "neutral",
5000
+ "study_type": "rct",
5001
+ "claim_text": "One week after training, retention differences between Codex and baseline groups did not reach statistical significance.",
5002
+ "effect_summary": "slightly better for Codex group but not significant",
5003
+ "confidence_markers": [
5004
+ "evidence_level:strong",
5005
+ "quality_score:9.0",
5006
+ "confidence:0.5",
5007
+ "decision_relation:neutral"
5008
+ ],
5009
+ "domains": [
5010
+ "ai-coding-assistant"
5011
+ ]
5012
+ },
5013
+ {
5014
+ "entry_id": "lib-ai-coding-assistant-E-004",
5015
+ "source_id": "S-2025-bastani",
5016
+ "title": "Generative AI without guardrails can harm learning: Evidence from high school mathematics",
5017
+ "year": 2025,
5018
+ "outcome_token": "independent_problem_solving",
5019
+ "outcome_tokens": [
5020
+ "independent_problem_solving"
5021
+ ],
5022
+ "direction": "contradict",
5023
+ "study_type": "rct",
5024
+ "claim_text": "Students with unguarded GPT-4 access performed 17% worse on the independent exam than the control group, despite higher practice performance.",
5025
+ "effect_summary": "negative_17_percent_on_independent_exam",
5026
+ "confidence_markers": [
5027
+ "evidence_level:strong",
5028
+ "quality_score:8.0",
5029
+ "confidence:0.75",
5030
+ "decision_relation:oppose_adoption"
5031
+ ],
5032
+ "domains": [
5033
+ "ai-coding-assistant"
5034
+ ]
5035
+ },
5036
+ {
5037
+ "entry_id": "lib-ai-coding-assistant-E-005",
5038
+ "source_id": "S-2025-bastani",
5039
+ "title": "Generative AI without guardrails can harm learning: Evidence from high school mathematics",
5040
+ "year": 2025,
5041
+ "outcome_token": "independent_problem_solving",
5042
+ "outcome_tokens": [
5043
+ "independent_problem_solving"
5044
+ ],
5045
+ "direction": "support",
5046
+ "study_type": "rct",
5047
+ "claim_text": "Guardrail design of the AI tutor (hints instead of answers, teacher-informed prompts) largely eliminated the negative learning effect.",
5048
+ "effect_summary": "negative effect essentially eradicated, no positive effect observed",
5049
+ "confidence_markers": [
5050
+ "evidence_level:strong",
5051
+ "quality_score:8.0",
5052
+ "confidence:0.75",
5053
+ "decision_relation:conditional"
5054
+ ],
5055
+ "domains": [
5056
+ "ai-coding-assistant"
5057
+ ]
5058
+ },
5059
+ {
5060
+ "entry_id": "lib-ai-coding-assistant-E-006",
5061
+ "source_id": "S-2025-bastani",
5062
+ "title": "Generative AI without guardrails can harm learning: Evidence from high school mathematics",
5063
+ "year": 2025,
5064
+ "outcome_token": "assignment_score",
5065
+ "outcome_tokens": [
5066
+ "assignment_score"
5067
+ ],
5068
+ "direction": "support",
5069
+ "study_type": "rct",
5070
+ "claim_text": "Access to GPT-4 during practice improves task performance (48% for GPT Base, 127% for GPT Tutor) — but this task performance does not transfer to independent exam performance.",
5071
+ "effect_summary": "48-127 percent improvement on practice problems",
5072
+ "confidence_markers": [
5073
+ "evidence_level:strong",
5074
+ "quality_score:8.0",
5075
+ "confidence:0.75",
5076
+ "decision_relation:conditional"
5077
+ ],
5078
+ "domains": [
5079
+ "ai-coding-assistant"
5080
+ ]
5081
+ },
5082
+ {
5083
+ "entry_id": "lib-ai-coding-assistant-E-007",
5084
+ "source_id": "S-2024-marzuki",
5085
+ "title": "Impact of ChatGPT on ESL students' academic writing skills",
5086
+ "year": 2024,
5087
+ "outcome_token": "knowledge_gain",
5088
+ "outcome_tokens": [
5089
+ "knowledge_gain"
5090
+ ],
5091
+ "direction": "support",
5092
+ "study_type": "mixed_methods",
5093
+ "claim_text": "ChatGPT as a formative feedback tool produced a significant positive impact on students' academic writing skills with positive student perceptions.",
5094
+ "effect_summary": "significant positive impact on writing skills",
5095
+ "confidence_markers": [
5096
+ "evidence_level:moderate",
5097
+ "quality_score:6.0",
5098
+ "confidence:0.55",
5099
+ "decision_relation:support_adoption"
5100
+ ],
5101
+ "domains": [
5102
+ "ai-coding-assistant"
5103
+ ]
5104
+ },
5105
+ {
5106
+ "entry_id": "lib-ai-coding-assistant-E-008",
5107
+ "source_id": "S-2023-peng",
5108
+ "title": "The Impact of AI on Developer Productivity: Evidence from GitHub Copilot",
5109
+ "year": 2023,
5110
+ "outcome_token": "completion_time",
5111
+ "outcome_tokens": [
5112
+ "completion_time"
5113
+ ],
5114
+ "direction": "support",
5115
+ "study_type": "rct",
5116
+ "claim_text": "Professional developers with Copilot access completed a standardized coding task about 55% faster than the control group (RCT, n=95).",
5117
+ "effect_summary": "~55.8% faster task completion in Copilot group",
5118
+ "confidence_markers": [
5119
+ "evidence_level:moderate",
5120
+ "quality_score:8.0",
5121
+ "confidence:0.6",
5122
+ "decision_relation:conditional"
5123
+ ],
5124
+ "domains": [
5125
+ "ai-coding-assistant"
5126
+ ]
5127
+ },
5128
+ {
5129
+ "entry_id": "lib-ai-coding-assistant-E-009",
5130
+ "source_id": "S-2023-yetistiren",
5131
+ "title": "GitHub Copilot AI Pair Programmer: Asset or Liability?",
5132
+ "year": 2023,
5133
+ "outcome_token": "code_quality",
5134
+ "outcome_tokens": [
5135
+ "code_quality"
5136
+ ],
5137
+ "direction": "support",
5138
+ "study_type": "observational",
5139
+ "claim_text": "Systematic benchmark evaluation reports mixed quality results for Copilot-generated code relative to human code: correctness competitive on parts of the benchmark while security-relevant defects are documented.",
5140
+ "effect_summary": "mixed quality profile; no single-direction summary",
5141
+ "confidence_markers": [
5142
+ "evidence_level:moderate",
5143
+ "quality_score:7.0",
5144
+ "confidence:0.55",
5145
+ "decision_relation:conditional"
5146
+ ],
5147
+ "domains": [
5148
+ "ai-coding-assistant"
5149
+ ]
5150
+ },
5151
+ {
5152
+ "entry_id": "lib-ai-coding-assistant-E-010",
5153
+ "source_id": "S-2022-finnie-ansley",
5154
+ "title": "Using GitHub Copilot to Solve Introductory Programming Problems",
5155
+ "year": 2022,
5156
+ "outcome_token": "assignment_score",
5157
+ "outcome_tokens": [
5158
+ "assignment_score"
5159
+ ],
5160
+ "direction": "support",
5161
+ "study_type": "observational",
5162
+ "claim_text": "Codex produced passing-level solutions for roughly half to three-quarters of CS1 exam-style questions depending on the dataset, indicating substantial task-capability headroom available to novices.",
5163
+ "effect_summary": "passing solutions on ~50-75% of questions across datasets",
5164
+ "confidence_markers": [
5165
+ "evidence_level:moderate",
5166
+ "quality_score:7.0",
5167
+ "confidence:0.55",
5168
+ "decision_relation:conditional"
5169
+ ],
5170
+ "domains": [
5171
+ "ai-coding-assistant"
5172
+ ]
5173
+ },
5174
+ {
5175
+ "entry_id": "lib-ai-coding-assistant-E-011",
5176
+ "source_id": "S-2023-explanations-compare",
5177
+ "title": "Comparing Code Explanations Created by Students and Large Language Models",
5178
+ "year": 2023,
5179
+ "outcome_token": "metacognition",
5180
+ "outcome_tokens": [
5181
+ "metacognition"
5182
+ ],
5183
+ "direction": "support",
5184
+ "study_type": "observational",
5185
+ "claim_text": "Controlled comparisons find LLM-generated code explanations comparable to (in places better than) student-authored explanations, suggesting viability as explanatory scaffold material rather than as a replacement for student explanation practice.",
5186
+ "effect_summary": "comparable-or-better rated quality vs student explanations",
5187
+ "confidence_markers": [
5188
+ "evidence_level:moderate",
5189
+ "quality_score:7.0",
5190
+ "confidence:0.55",
5191
+ "decision_relation:conditional"
5192
+ ],
5193
+ "domains": [
5194
+ "ai-coding-assistant"
5195
+ ]
5196
+ },
5197
+ {
5198
+ "entry_id": "lib-ai-coding-assistant-E-012",
5199
+ "source_id": "S-2022-vaithilingam",
5200
+ "title": "Expectation vs. Experience: Evaluating the Usability of Code Generation Tools",
5201
+ "year": 2022,
5202
+ "outcome_token": "over_reliance",
5203
+ "outcome_tokens": [
5204
+ "over_reliance"
5205
+ ],
5206
+ "direction": "support",
5207
+ "study_type": "qualitative",
5208
+ "claim_text": "Despite faster first-task completion, participants struggled to understand and debug AI-generated solutions and reported low ownership of the final program - documenting metacognitive and dependence risks that pure speed metrics miss.",
5209
+ "effect_summary": "documented comprehension/ownership difficulties despite speed gain",
5210
+ "confidence_markers": [
5211
+ "evidence_level:moderate",
5212
+ "quality_score:7.0",
5213
+ "confidence:0.55",
5214
+ "decision_relation:conditional"
5215
+ ],
5216
+ "domains": [
5217
+ "ai-coding-assistant"
5218
+ ]
5219
+ },
4944
5220
  {
4945
5221
  "entry_id": "lib-ai-tutor-E-001",
4946
5222
  "source_id": "S-2025-bastani",
@@ -5264,5 +5540,5 @@
5264
5540
  ]
5265
5541
  }
5266
5542
  ],
5267
- "coverage_note": "内置证据库:由 30 份金标准标注(benchmarks/annotations/gold-Q01..Q30 的 key_claims/key_supporting_sources/known_contradictions/correct_outcome_types)+ 3 个示例工作流 evidence.jsonl(ai-coding-assistant / ai-tutor / ai-writing-assistant)抽取生成;按 (source_id, outcome_token, claim_text) 去重合并。direction 语义为采纳方向:support=支持采纳(初步裁决=>pilot),contradict=反对采纳(=>reject),neutral=中性;金标准条目按 expected_decision_range 粗粒度映射方向(纯 reject 问题反向映射),conflict 与混合方向问题的单条断言方向可能不精确。仅用于离线初步裁决(preliminary,保守),从不直接给出 adopt。"
5543
+ "coverage_note": "内置证据库:由 30 份金标准标注(benchmarks/annotations/gold-Q01..Q30 的 key_claims/key_supporting_sources/known_contradictions/correct_outcome_types)+ 3 个示例工作流 evidence.jsonl(ai-coding-assistant / ai-tutor / ai-writing-assistant)抽取生成;按 (source_id, outcome_token, claim_text) 去重合并。direction 语义为采纳方向:support=支持采纳(离线初筛上限 => pilot,永不输出 adopt),contradict=反对采纳(oppose-only 时 => reject;与 support 并存属未决冲突,应交 engine.decision_policy.decision_outcome 判为 INSUFFICIENT_EVIDENCE),neutral=中性;金标准条目按 expected_decision_range 粗粒度映射方向(纯 reject 问题反向映射),conflict 与混合方向问题的单条断言方向可能不精确。仅用于离线初步裁决(preliminary,保守),从不直接给出 adopt;完整 ADOPT 只能由 engine.decision_policy.decision_outcome 给出(High + support + 主要结果 directness 2)。"
5268
5544
  }
@@ -17,7 +17,8 @@ const PKG_ROOT = path.resolve(__dirname, '..');
17
17
  const INSTALL_SH = path.join(PKG_ROOT, 'install.sh');
18
18
 
19
19
  function usage() {
20
- console.log(`eduevidence ${readPkgVersion()} — Evidence-Based Education Decision Skill
20
+ console.log(`eduevidence ${readPkgVersion()} — Evidence-Based Decision Skill
21
+ for education and organizational policy (multi-domain by contract)
21
22
 
22
23
  Install (Python runtime + self-check):
23
24
  eduevidence install [--dev] [--dry-run]