eduevidence 6.2.0 → 6.3.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/CHANGELOG.md +395 -0
- package/README.md +22 -13
- package/README.zh-CN.md +15 -8
- package/SKILL.md +10 -9
- package/benchmarks/evidence-library.json +277 -1
- package/docs/architecture.md +6 -3
- package/docs/j-ev-experimental.md +250 -0
- package/docs/reproducibility.md +138 -0
- package/domains/_neutral/copy/few_shots.json +21 -0
- package/domains/_neutral/copy/framing_lexicon.json +19 -0
- package/domains/_neutral/copy/module_labels.json +5 -0
- package/domains/_neutral/copy/module_labels_footer.json +102 -0
- package/domains/_neutral/copy/module_labels_modules.json +204 -0
- package/domains/_neutral/copy/module_labels_nav.json +126 -0
- package/domains/_neutral/copy/module_labels_summary.json +98 -0
- package/domains/_neutral/copy/module_labels_tables.json +164 -0
- package/domains/_neutral/copy/module_labels_v2.json +90 -0
- package/domains/_neutral/copy/risk_constructs.json +20 -0
- package/domains/_neutral/copy/section_titles.json +66 -0
- package/domains/_neutral/copy/terminology.json +11 -0
- package/domains/check_copy_packs.py +103 -0
- package/domains/education/copy/few_shots.json +22 -0
- package/domains/education/copy/framing_enums.json +167 -0
- package/domains/education/copy/framing_lexicon.json +166 -0
- package/domains/education/copy/module_labels.json +169 -0
- package/domains/education/copy/risk_constructs.json +48 -0
- package/domains/education/copy/section_titles.json +186 -0
- package/domains/education/copy/terminology.json +70 -0
- package/domains/education/manifest.json +1 -1
- package/domains/education/outcome_taxonomy.json +2 -2
- package/domains/manifest.json +1 -1
- package/domains/policy/copy/few_shots.json +22 -0
- package/domains/policy/copy/framing_enums.json +94 -0
- package/domains/policy/copy/framing_lexicon.json +174 -0
- package/domains/policy/copy/module_labels.json +168 -0
- package/domains/policy/copy/risk_constructs.json +33 -0
- package/domains/policy/copy/section_titles.json +186 -0
- package/domains/policy/copy/terminology.json +64 -0
- package/engine/capabilities.py +57 -5
- package/engine/decision_policy.py +88 -17
- package/engine/library_builtin.py +7 -4
- package/engine/tribunal.py +17 -23
- package/engine/versions.py +1 -1
- package/examples/ai-coding-assistant-evidence/EduEvidence_Report.html +4 -4
- package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_academic.html +4 -4
- package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_claude.html +4 -4
- package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_datalab-dark.html +4 -4
- package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_datalab.html +4 -4
- package/examples/ai-coding-assistant-evidence/reports-5themes/EduEvidence_Report_presentation.html +4 -4
- package/examples/spaced-retrieval-practice/EduEvidence_Report.html +2728 -0
- package/examples/spaced-retrieval-practice/report.html +2522 -0
- package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_academic.html +4 -4
- package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_claude.html +4 -4
- package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_datalab-dark.html +4 -4
- package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_datalab.html +4 -4
- package/examples/spaced-retrieval-practice/reports-5themes/EduEvidence_Report_presentation.html +4 -4
- package/examples/workplace-ai-assistant/EduEvidence_Report.html +2814 -0
- package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_academic.html +36 -36
- package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_claude.html +36 -36
- package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_datalab-dark.html +36 -36
- package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_datalab.html +36 -36
- package/examples/workplace-ai-assistant/reports-5themes/EduEvidence_Report_presentation.html +36 -36
- package/integrations/jev/__init__.py +115 -0
- package/integrations/jev/approval.py +212 -0
- package/integrations/jev/cli.py +84 -0
- package/integrations/jev/config.py +112 -0
- package/integrations/jev/gateway.py +128 -0
- package/integrations/jev/modes.py +38 -0
- package/integrations/jev/tools_classify.py +88 -0
- package/integrations/jev/tools_extract.py +111 -0
- package/integrations/jev/tools_rerank.py +71 -0
- package/integrations/jev/tools_screen.py +87 -0
- package/integrations/jev/tools_verify.py +95 -0
- package/integrations/jev_mcp.py +22 -0
- package/integrations/semantic_decide.py +286 -0
- package/integrations/semdecide_cli.py +55 -0
- package/package.json +9 -1
- package/pyproject.toml +1 -1
- package/references/report-copy-style.md +43 -3
- package/schemas/v2/decision-snapshot.schema.json +20 -9
- package/schemas/v2/intake.schema.json +191 -0
- package/scripts/build_evidence_library.py +15 -5
- package/scripts/dashboard_server.py +13 -2
- package/scripts/intake/__init__.py +31 -0
- package/scripts/intake/__main__.py +18 -0
- package/scripts/intake/background.py +78 -0
- package/scripts/intake/browser.py +79 -0
- package/scripts/intake/cli.py +57 -0
- package/scripts/intake/constants.py +57 -0
- package/scripts/intake/depth.py +53 -0
- package/scripts/intake/enhancements.py +106 -0
- package/scripts/intake/hooks.py +90 -0
- package/scripts/intake/prefs.py +76 -0
- package/scripts/intake/prompts.py +85 -0
- package/scripts/intake/session.py +152 -0
- package/scripts/lint_file_layers.py +126 -0
- package/scripts/orchestrator.py +68 -17
- package/scripts/pre_verdict_gate.py +21 -7
- package/scripts/skill_lint.py +11 -1
- package/scripts/skill_payload.py +3 -3
- package/scripts/test_adversarial_empirical.py +70 -6
- package/skill/agents/evidence-judge.md +49 -7
- package/skill/workflows/experimental-jev.md +170 -0
- package/skill/workflows/intake.md +120 -0
- package/visualization/eduevidence-report/scripts/build_infographics.py +32 -14
- package/visualization/eduevidence-report/scripts/build_report.py +75 -662
- package/visualization/eduevidence-report/scripts/report_copy_pack.py +296 -0
- package/visualization/eduevidence-report/scripts/report_copy_policy_guard.py +47 -0
- package/visualization/eduevidence-report/scripts/zh_labels.py +61 -0
- package/scripts/build_esl_artifacts.py +0 -1921
- package/scripts/build_killer_demo.py +0 -295
- package/scripts/enrich_projects_human_and_lieflat.py +0 -315
- package/scripts/generate_new_projects.py +0 -686
- package/scripts/sync_killer_demo_report.py +0 -270
|
@@ -0,0 +1,186 @@
|
|
|
1
|
+
{
|
|
2
|
+
"domain": "education",
|
|
3
|
+
"purpose": "Section titles/leads + default full-report chapter plan (education voice).",
|
|
4
|
+
"section_titles": {
|
|
5
|
+
"zh": {
|
|
6
|
+
"01": "01 执行决策",
|
|
7
|
+
"02": "02 结果证据概览",
|
|
8
|
+
"03": "03 证据矩阵",
|
|
9
|
+
"04": "04 证据裁决",
|
|
10
|
+
"05": "05 方法学审计",
|
|
11
|
+
"06": "06 冲突分析",
|
|
12
|
+
"07": "07 主张-证据追溯",
|
|
13
|
+
"08": "08 适用性",
|
|
14
|
+
"09": "09 教学干预",
|
|
15
|
+
"10": "10 评价方案",
|
|
16
|
+
"11": "11 基准测试",
|
|
17
|
+
"12": "12 来源与溯源"
|
|
18
|
+
},
|
|
19
|
+
"en": {
|
|
20
|
+
"01": "01 Executive Decision",
|
|
21
|
+
"02": "02 Outcome Evidence Overview",
|
|
22
|
+
"03": "03 Evidence Matrix",
|
|
23
|
+
"04": "04 Evidence Tribunal",
|
|
24
|
+
"05": "05 Methodology Audit",
|
|
25
|
+
"06": "06 Conflict Analysis",
|
|
26
|
+
"07": "07 Claim-Evidence Trace",
|
|
27
|
+
"08": "08 Applicability",
|
|
28
|
+
"09": "09 Teaching Intervention",
|
|
29
|
+
"10": "10 Evaluation Plan",
|
|
30
|
+
"11": "11 Benchmark",
|
|
31
|
+
"12": "12 Sources & Provenance"
|
|
32
|
+
}
|
|
33
|
+
},
|
|
34
|
+
"section_leads": {
|
|
35
|
+
"zh": {
|
|
36
|
+
"01": "本节先给结论:最终怎么裁决、置信度多高、靠哪几条证据。",
|
|
37
|
+
"02": "一图看清:哪些学习结果有支持证据、哪些被反驳。",
|
|
38
|
+
"03": "每条证据来自哪项研究、测了什么、方向与质量如何;可筛选、可搜索。",
|
|
39
|
+
"04": "证据允许主张什么、不允许主张什么;缺失的关键证据是什么。",
|
|
40
|
+
"05": "研究质量可靠吗?哪些方法学问题让结论打折。",
|
|
41
|
+
"06": "不同研究为何结论不同;分歧出在哪一环。",
|
|
42
|
+
"07": "从结论到证据到原始来源,每一步都可追查。",
|
|
43
|
+
"08": "结论适用于谁、什么课程与结果、需要什么条件。",
|
|
44
|
+
"09": "试点怎么分阶段放开 AI 规则;什么情况必须叫停。",
|
|
45
|
+
"10": "如何验证效果:指标、对照、成功阈值。",
|
|
46
|
+
"11": "EduEvidence 自身基准表现:引用精度与成本。",
|
|
47
|
+
"12": "每篇文献是谁、出自哪里、如何获取。"
|
|
48
|
+
},
|
|
49
|
+
"en": {
|
|
50
|
+
"01": "The verdict first: what we decide, at what confidence, on which evidence.",
|
|
51
|
+
"02": "At a glance: which learning outcomes have supporting evidence, which are contradicted.",
|
|
52
|
+
"03": "Where each piece of evidence comes from, what it measures, its direction and quality — filterable and searchable.",
|
|
53
|
+
"04": "What the evidence lets us claim, what it does not, and what is still missing.",
|
|
54
|
+
"05": "How reliable are these studies, and which methodological concerns discount the conclusions.",
|
|
55
|
+
"06": "Why studies disagree — and where exactly they diverge.",
|
|
56
|
+
"07": "Every step from conclusion to evidence to source stays traceable.",
|
|
57
|
+
"08": "Who the conclusion applies to, for which course and outcomes, under what conditions.",
|
|
58
|
+
"09": "How AI usage rules phase in during a pilot, and when we must stop.",
|
|
59
|
+
"10": "How we verify real effects: metrics, comparison, success threshold.",
|
|
60
|
+
"11": "How EduEvidence itself performs: citation precision and cost.",
|
|
61
|
+
"12": "Who wrote each cited study, where it came from, how it was fetched."
|
|
62
|
+
}
|
|
63
|
+
},
|
|
64
|
+
"default_chapters": [
|
|
65
|
+
{
|
|
66
|
+
"key": "decision",
|
|
67
|
+
"title_zh": "结论、裁决与研究边界",
|
|
68
|
+
"title_en": "Decision, Adjudication & Research Boundary",
|
|
69
|
+
"lead_zh": "先明确最终裁决与研究边界,再解释为什么。",
|
|
70
|
+
"lead_en": "State the final adjudication and research boundary before explaining why.",
|
|
71
|
+
"modules": [
|
|
72
|
+
"decision",
|
|
73
|
+
"scope"
|
|
74
|
+
]
|
|
75
|
+
},
|
|
76
|
+
{
|
|
77
|
+
"key": "evidence",
|
|
78
|
+
"title_zh": "关键证据与结果分离",
|
|
79
|
+
"title_en": "Key Evidence & Outcome Separation",
|
|
80
|
+
"lead_zh": "把任务表现、真实学习、保持与风险放在同一证据地图中,但不混为一谈。",
|
|
81
|
+
"lead_en": "Place task performance, actual learning, retention and risk on one evidence map without conflating them.",
|
|
82
|
+
"modules": [
|
|
83
|
+
"retrieval",
|
|
84
|
+
"outcomes",
|
|
85
|
+
"evidence"
|
|
86
|
+
]
|
|
87
|
+
},
|
|
88
|
+
{
|
|
89
|
+
"key": "quality",
|
|
90
|
+
"title_zh": "证据可信度、反证与方法审计",
|
|
91
|
+
"title_en": "Evidence Quality, Counterevidence & Method Audit",
|
|
92
|
+
"lead_zh": "检查证据为什么可信、哪里冲突,以及哪些结论必须降级。",
|
|
93
|
+
"lead_en": "Examine why evidence is credible, where it conflicts, and which conclusions require downgrading.",
|
|
94
|
+
"modules": [
|
|
95
|
+
"quality",
|
|
96
|
+
"conflicts",
|
|
97
|
+
"trace"
|
|
98
|
+
]
|
|
99
|
+
},
|
|
100
|
+
{
|
|
101
|
+
"key": "action",
|
|
102
|
+
"title_zh": "适用范围与教学行动",
|
|
103
|
+
"title_en": "Applicability & Teaching Action",
|
|
104
|
+
"lead_zh": "把可外推范围、护栏和教学动作连接到具体证据。",
|
|
105
|
+
"lead_en": "Connect applicability, guardrails and teaching actions to specific evidence.",
|
|
106
|
+
"modules": [
|
|
107
|
+
"applicability",
|
|
108
|
+
"intervention"
|
|
109
|
+
]
|
|
110
|
+
},
|
|
111
|
+
{
|
|
112
|
+
"key": "evaluation",
|
|
113
|
+
"title_zh": "试点设计、评估与停止条件",
|
|
114
|
+
"title_en": "Pilot, Evaluation & Stop Conditions",
|
|
115
|
+
"lead_zh": "用独立学习结果验证试点,并预先写清停止条件。",
|
|
116
|
+
"lead_en": "Validate the pilot with independent learning outcomes and pre-specified stop conditions.",
|
|
117
|
+
"modules": [
|
|
118
|
+
"evaluation"
|
|
119
|
+
]
|
|
120
|
+
},
|
|
121
|
+
{
|
|
122
|
+
"key": "sources",
|
|
123
|
+
"title_zh": "来源、溯源与附录",
|
|
124
|
+
"title_en": "Sources, Traceability & Appendix",
|
|
125
|
+
"lead_zh": "保留原始来源、URL、证据 ID 和获取信息,确保可回查。",
|
|
126
|
+
"lead_en": "Preserve original sources, URLs, evidence IDs and retrieval metadata for auditability.",
|
|
127
|
+
"modules": [
|
|
128
|
+
"sources"
|
|
129
|
+
]
|
|
130
|
+
}
|
|
131
|
+
],
|
|
132
|
+
"brief_blocks": {
|
|
133
|
+
"zh": {
|
|
134
|
+
"decision": [
|
|
135
|
+
"先看结论",
|
|
136
|
+
"该不该做、置信度多高、最关键的证据边界在哪。"
|
|
137
|
+
],
|
|
138
|
+
"lieflat": [
|
|
139
|
+
"Lieflat 实证手作画廊",
|
|
140
|
+
"AI 按数据形状从 Lieflat 目录选型编排;每张图的数字都可溯源到 result.json。"
|
|
141
|
+
],
|
|
142
|
+
"outcomes": [
|
|
143
|
+
"任务表现 ≠ 学习效果",
|
|
144
|
+
"只展示真正有解释力的结果分离;正向、负向与零效应按 effect_direction 编码。"
|
|
145
|
+
],
|
|
146
|
+
"tribunal": [
|
|
147
|
+
"证据裁决",
|
|
148
|
+
"支持、不确定、被反驳与缺失证据分开放置,不把长段落平铺在同一层。"
|
|
149
|
+
],
|
|
150
|
+
"action": [
|
|
151
|
+
"从证据到行动",
|
|
152
|
+
"适用性、护栏、停止条件与评价连成一条可执行路径。"
|
|
153
|
+
],
|
|
154
|
+
"sources": [
|
|
155
|
+
"关键来源",
|
|
156
|
+
"摘要页只列最关键的来源;完整溯源在完整报告中展开。"
|
|
157
|
+
]
|
|
158
|
+
},
|
|
159
|
+
"en": {
|
|
160
|
+
"decision": [
|
|
161
|
+
"Decision first",
|
|
162
|
+
"What to do, how confident we are, and the most important evidence boundary."
|
|
163
|
+
],
|
|
164
|
+
"lieflat": [
|
|
165
|
+
"Lieflat Editorial Gallery",
|
|
166
|
+
"Charts selected and composed by AI from the Lieflat catalog; every number traces back to result.json."
|
|
167
|
+
],
|
|
168
|
+
"outcomes": [
|
|
169
|
+
"Task performance ≠ learning",
|
|
170
|
+
"Only informative outcome separation; positive, negative and null effects use effect_direction."
|
|
171
|
+
],
|
|
172
|
+
"tribunal": [
|
|
173
|
+
"Evidence tribunal",
|
|
174
|
+
"Supported, uncertain, contradicted and missing evidence stay separated instead of flattened into long prose."
|
|
175
|
+
],
|
|
176
|
+
"action": [
|
|
177
|
+
"Evidence to action",
|
|
178
|
+
"Applicability, guardrails, stop conditions and evaluation form one executable path."
|
|
179
|
+
],
|
|
180
|
+
"sources": [
|
|
181
|
+
"Key sources",
|
|
182
|
+
"Only the key sources in the brief; full traceability expands in the full report."
|
|
183
|
+
]
|
|
184
|
+
}
|
|
185
|
+
}
|
|
186
|
+
}
|
|
@@ -0,0 +1,70 @@
|
|
|
1
|
+
{
|
|
2
|
+
"domain": "education",
|
|
3
|
+
"purpose": "Domain terminology: methodology item labels + display names for internal fields.",
|
|
4
|
+
"methodology_labels": {
|
|
5
|
+
"zh": {
|
|
6
|
+
"control_group": "对照组",
|
|
7
|
+
"randomization": "随机分配",
|
|
8
|
+
"pre_test": "前测",
|
|
9
|
+
"post_test": "后测",
|
|
10
|
+
"retention_test": "保持测试",
|
|
11
|
+
"transfer_test": "迁移测试",
|
|
12
|
+
"sample_bias": "样本偏差",
|
|
13
|
+
"self_selection": "自我选择偏差",
|
|
14
|
+
"measurement_validity": "测量效度",
|
|
15
|
+
"confounders": "混杂因素",
|
|
16
|
+
"instructor_effect": "教师效应",
|
|
17
|
+
"novelty_effect": "新奇效应",
|
|
18
|
+
"tool_version_effect": "工具版本效应",
|
|
19
|
+
"ai_usage_policy": "AI 使用规则",
|
|
20
|
+
"dropout": "样本流失"
|
|
21
|
+
},
|
|
22
|
+
"en": {
|
|
23
|
+
"control_group": "Control group",
|
|
24
|
+
"randomization": "Randomization",
|
|
25
|
+
"pre_test": "Pre-test",
|
|
26
|
+
"post_test": "Post-test",
|
|
27
|
+
"retention_test": "Retention test",
|
|
28
|
+
"transfer_test": "Transfer test",
|
|
29
|
+
"sample_bias": "Sample bias",
|
|
30
|
+
"self_selection": "Self-selection bias",
|
|
31
|
+
"measurement_validity": "Measurement validity",
|
|
32
|
+
"confounders": "Confounders",
|
|
33
|
+
"instructor_effect": "Instructor effect",
|
|
34
|
+
"novelty_effect": "Novelty effect",
|
|
35
|
+
"tool_version_effect": "Tool-version effect",
|
|
36
|
+
"ai_usage_policy": "AI usage policy",
|
|
37
|
+
"dropout": "Attrition"
|
|
38
|
+
}
|
|
39
|
+
},
|
|
40
|
+
"field_labels": {
|
|
41
|
+
"zh": {
|
|
42
|
+
"target_learners": "目标学习者",
|
|
43
|
+
"target_population": "目标人群",
|
|
44
|
+
"ai_usage_policy": "AI 使用规则",
|
|
45
|
+
"ai_usage_rule": "AI 规则",
|
|
46
|
+
"task_vs_learning_guard": "任务 vs 学习护栏",
|
|
47
|
+
"learning_metrics": "学习指标",
|
|
48
|
+
"process_metrics": "过程指标",
|
|
49
|
+
"risk_metrics": "风险指标",
|
|
50
|
+
"task_performance": "任务表现",
|
|
51
|
+
"learning_gain": "学习收获",
|
|
52
|
+
"retention": "保持",
|
|
53
|
+
"transfer": "迁移"
|
|
54
|
+
},
|
|
55
|
+
"en": {
|
|
56
|
+
"target_learners": "Target learners",
|
|
57
|
+
"target_population": "Target population",
|
|
58
|
+
"ai_usage_policy": "AI usage policy",
|
|
59
|
+
"ai_usage_rule": "AI rule",
|
|
60
|
+
"task_vs_learning_guard": "Task vs learning guard",
|
|
61
|
+
"learning_metrics": "Learning metrics",
|
|
62
|
+
"process_metrics": "Process metrics",
|
|
63
|
+
"risk_metrics": "Risk metrics",
|
|
64
|
+
"task_performance": "Task performance",
|
|
65
|
+
"learning_gain": "Learning gain",
|
|
66
|
+
"retention": "Retention",
|
|
67
|
+
"transfer": "Transfer"
|
|
68
|
+
}
|
|
69
|
+
}
|
|
70
|
+
}
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
"description": "education 域只是“指向现有契约”的注册:不新增任何逻辑路径,不引入新 schema / 新校验器 / 新方法学。引擎层 load_domain('education') 只做契约存在性校验;领域选择由主 agent 接 CLI 完成。",
|
|
5
5
|
"contracts_source": {
|
|
6
6
|
"frame_schema": "schemas/education-frame.schema.json(EducationResearchFrame,顶层 additionalProperties:false,required: question/decision_target)",
|
|
7
|
-
"outcome_taxonomy": "20 token
|
|
7
|
+
"outcome_taxonomy": "20 token 与四类映射(learning/task_performance/process/risk)都在 domains/education/outcome_taxonomy.json;唯一读取者为 engine/taxonomy.py(未知 token fail-closed)",
|
|
8
8
|
"methodology_checklist": "15 项取自 skill/agents/method-reviewer.md;每项 status 枚举 met/partial/missing/not_applicable",
|
|
9
9
|
"golds_dir": "benchmarks/annotations(30 条人工金标 gold-Q*.json)",
|
|
10
10
|
"references_dir": "references(11 篇教育证据方法学文档)"
|
|
@@ -1,8 +1,8 @@
|
|
|
1
1
|
{
|
|
2
2
|
"domain": "education",
|
|
3
3
|
"source": {
|
|
4
|
-
"
|
|
5
|
-
"
|
|
4
|
+
"authority": "engine/taxonomy.py",
|
|
5
|
+
"note": "Sole reader of this contract. Token set and token-to-category mapping live in this file; callers must use engine/taxonomy.py (fail-closed) and must not keep private copies of either table."
|
|
6
6
|
},
|
|
7
7
|
"categories": {
|
|
8
8
|
"learning": {
|
package/domains/manifest.json
CHANGED
|
@@ -5,7 +5,7 @@
|
|
|
5
5
|
{
|
|
6
6
|
"id": "education",
|
|
7
7
|
"name": "education",
|
|
8
|
-
"description": "教育研究领域(注册域,零逻辑路径):框架复用 schemas/education-frame.schema.json(EducationResearchFrame
|
|
8
|
+
"description": "教育研究领域(注册域,零逻辑路径):框架复用 schemas/education-frame.schema.json(EducationResearchFrame),结果分类(20 token 与四类映射)由 domains/education/outcome_taxonomy.json 声明、engine/taxonomy.py 单一读取(未知 token fail-closed),方法学清单复用 skill/agents/method-reviewer.md 的 15 项审查清单。",
|
|
9
9
|
"frame_schema": "schemas/education-frame.schema.json",
|
|
10
10
|
"outcome_taxonomy": "domains/education/outcome_taxonomy.json",
|
|
11
11
|
"methodology_checklist": "domains/education/manifest.json#/methodology_checklist",
|
|
@@ -0,0 +1,22 @@
|
|
|
1
|
+
{
|
|
2
|
+
"domain": "policy",
|
|
3
|
+
"purpose": "One written example per decision state, policy voice. No teaching vocabulary; intervention is 干预方案.",
|
|
4
|
+
"states": {
|
|
5
|
+
"adopt": {
|
|
6
|
+
"zh": "独立效果、成本与公平证据一致为正,实施风险可控 → 在目标辖区采用该干预方案,并保留持续监测与退出条款。",
|
|
7
|
+
"en": "Independent effect, cost and equity evidence is consistently positive with controllable implementation risk → adopt the intervention plan in the target jurisdiction, with ongoing monitoring and exit clauses."
|
|
8
|
+
},
|
|
9
|
+
"pilot": {
|
|
10
|
+
"zh": "试点期产出改善明确,但独立效果与成本证据仍薄 → 有界、分阶段、带对照与成本核算的试点,而不是全面铺开。",
|
|
11
|
+
"en": "Pilot-period output gains are clear, but independent effect and cost evidence is still thin → a bounded, phased pilot with comparison and cost accounting rather than full rollout."
|
|
12
|
+
},
|
|
13
|
+
"reject": {
|
|
14
|
+
"zh": "已观察到实施风险或分配恶化,且无补偿性目标效果证据 → 不建议采用,并写明受影响人群与补救边界。",
|
|
15
|
+
"en": "Implementation risk or distributional harm is already observed with no compensatory target-effect evidence → do not adopt, and state affected populations and the remediation boundary."
|
|
16
|
+
},
|
|
17
|
+
"insufficient_evidence": {
|
|
18
|
+
"zh": "现有研究未测量独立效果、成本或分配,或辖区与目标人群错位 → 证据不足,先补关键测量设计,再谈决策。",
|
|
19
|
+
"en": "Existing studies do not measure independent effects, cost or distribution, or the jurisdiction mismatches the target population → insufficient evidence; fill the key measurement design before deciding."
|
|
20
|
+
}
|
|
21
|
+
}
|
|
22
|
+
}
|
|
@@ -0,0 +1,94 @@
|
|
|
1
|
+
{
|
|
2
|
+
"domain": "policy",
|
|
3
|
+
"purpose": "policy frame enum display labels (framing_lexicon part 2).",
|
|
4
|
+
"frame_enums": {
|
|
5
|
+
"zh": {
|
|
6
|
+
"policy_decision": "政策决策",
|
|
7
|
+
"evidence_review": "证据评审",
|
|
8
|
+
"pilot_design": "试点设计",
|
|
9
|
+
"evaluation_design": "评价设计",
|
|
10
|
+
"primary": "主要结果",
|
|
11
|
+
"secondary": "次要结果",
|
|
12
|
+
"risk": "风险结果",
|
|
13
|
+
"adopt": "采纳",
|
|
14
|
+
"modify": "修改",
|
|
15
|
+
"terminate": "终止",
|
|
16
|
+
"maintain": "维持",
|
|
17
|
+
"evaluate_impact": "评估影响",
|
|
18
|
+
"rct": "随机对照试验",
|
|
19
|
+
"quasi_experimental": "准实验",
|
|
20
|
+
"observational": "观察性研究",
|
|
21
|
+
"meta_analysis": "元分析",
|
|
22
|
+
"systematic_review": "系统综述",
|
|
23
|
+
"qualitative": "质性研究",
|
|
24
|
+
"mixed_methods": "混合方法",
|
|
25
|
+
"survey": "调查研究",
|
|
26
|
+
"case_study": "案例研究",
|
|
27
|
+
"literature_review": "文献综述",
|
|
28
|
+
"online": "线上",
|
|
29
|
+
"offline": "线下",
|
|
30
|
+
"hybrid": "混合",
|
|
31
|
+
"policy_effectiveness": "政策有效性",
|
|
32
|
+
"cost_effectiveness": "成本效益",
|
|
33
|
+
"equity": "公平性",
|
|
34
|
+
"feasibility": "可行性",
|
|
35
|
+
"implementation_risk": "实施风险",
|
|
36
|
+
"institutional_reform": "机构改革",
|
|
37
|
+
"human_supervised_customer_support_assistant": "人工监督的客服助手",
|
|
38
|
+
"approved_knowledge_base_and_agent_review": "已批准知识库 + 坐席复核",
|
|
39
|
+
"enterprise_customer_support_staff": "企业客服人员",
|
|
40
|
+
"enterprise_customer_support": "企业客户支持",
|
|
41
|
+
"autonomous_agents_and_high_stakes_specialist_advice": "自主代理与高风险专业建议",
|
|
42
|
+
"under_design_pending_evidence_review": "设计中(待证据评审)",
|
|
43
|
+
"knowledge_work": "知识型工作",
|
|
44
|
+
"software_development": "软件开发",
|
|
45
|
+
"daily": "每日",
|
|
46
|
+
"weekly": "每周",
|
|
47
|
+
"monthly": "每月"
|
|
48
|
+
},
|
|
49
|
+
"en": {
|
|
50
|
+
"policy_decision": "Policy decision",
|
|
51
|
+
"evidence_review": "Evidence review",
|
|
52
|
+
"pilot_design": "Pilot design",
|
|
53
|
+
"evaluation_design": "Evaluation design",
|
|
54
|
+
"primary": "Primary outcomes",
|
|
55
|
+
"secondary": "Secondary outcomes",
|
|
56
|
+
"risk": "Risk outcomes",
|
|
57
|
+
"adopt": "Adopt",
|
|
58
|
+
"modify": "Modify",
|
|
59
|
+
"terminate": "Terminate",
|
|
60
|
+
"maintain": "Maintain",
|
|
61
|
+
"evaluate_impact": "Evaluate impact",
|
|
62
|
+
"rct": "Randomised controlled trial",
|
|
63
|
+
"quasi_experimental": "Quasi-experimental",
|
|
64
|
+
"observational": "Observational",
|
|
65
|
+
"meta_analysis": "Meta-analysis",
|
|
66
|
+
"systematic_review": "Systematic review",
|
|
67
|
+
"qualitative": "Qualitative",
|
|
68
|
+
"mixed_methods": "Mixed methods",
|
|
69
|
+
"survey": "Survey",
|
|
70
|
+
"case_study": "Case study",
|
|
71
|
+
"literature_review": "Literature review",
|
|
72
|
+
"online": "Online",
|
|
73
|
+
"offline": "Offline",
|
|
74
|
+
"hybrid": "Hybrid",
|
|
75
|
+
"policy_effectiveness": "Policy effectiveness",
|
|
76
|
+
"cost_effectiveness": "Cost-effectiveness",
|
|
77
|
+
"equity": "Equity",
|
|
78
|
+
"feasibility": "Feasibility",
|
|
79
|
+
"implementation_risk": "Implementation risk",
|
|
80
|
+
"institutional_reform": "Institutional reform",
|
|
81
|
+
"human_supervised_customer_support_assistant": "Human-supervised customer-support assistant",
|
|
82
|
+
"approved_knowledge_base_and_agent_review": "Approved knowledge base + agent review",
|
|
83
|
+
"enterprise_customer_support_staff": "Customer-support staff",
|
|
84
|
+
"enterprise_customer_support": "Enterprise customer support",
|
|
85
|
+
"autonomous_agents_and_high_stakes_specialist_advice": "Autonomous agents and high-stakes specialist advice",
|
|
86
|
+
"under_design_pending_evidence_review": "Under design (pending evidence review)",
|
|
87
|
+
"knowledge_work": "Knowledge work",
|
|
88
|
+
"software_development": "Software development",
|
|
89
|
+
"daily": "Daily",
|
|
90
|
+
"weekly": "Weekly",
|
|
91
|
+
"monthly": "Monthly"
|
|
92
|
+
}
|
|
93
|
+
}
|
|
94
|
+
}
|
|
@@ -0,0 +1,174 @@
|
|
|
1
|
+
{
|
|
2
|
+
"domain": "policy",
|
|
3
|
+
"purpose": "policy frame field/subfield labels (framing_lexicon part 1).",
|
|
4
|
+
"frame_fields": {
|
|
5
|
+
"zh": {
|
|
6
|
+
"question": "研究问题",
|
|
7
|
+
"decision_object": "决策对象",
|
|
8
|
+
"intervention": "干预方案",
|
|
9
|
+
"population": "目标人群",
|
|
10
|
+
"stakeholders": "利益相关者",
|
|
11
|
+
"outcomes": "结果构念",
|
|
12
|
+
"context": "政策情境",
|
|
13
|
+
"scope": "研究范围",
|
|
14
|
+
"comparison": "比较条件",
|
|
15
|
+
"success": "决策成功条件",
|
|
16
|
+
"success_condition": "决策成功条件",
|
|
17
|
+
"learner": "目标人群",
|
|
18
|
+
"course": "适用情境"
|
|
19
|
+
},
|
|
20
|
+
"en": {
|
|
21
|
+
"question": "Research question",
|
|
22
|
+
"decision_object": "Decision object",
|
|
23
|
+
"intervention": "Intervention plan",
|
|
24
|
+
"population": "Target population",
|
|
25
|
+
"stakeholders": "Stakeholders",
|
|
26
|
+
"outcomes": "Outcome constructs",
|
|
27
|
+
"context": "Policy context",
|
|
28
|
+
"scope": "Research scope",
|
|
29
|
+
"comparison": "Comparison",
|
|
30
|
+
"success": "Decision success condition",
|
|
31
|
+
"success_condition": "Decision success condition",
|
|
32
|
+
"learner": "Target population",
|
|
33
|
+
"course": "Operating context"
|
|
34
|
+
}
|
|
35
|
+
},
|
|
36
|
+
"frame_subfields": {
|
|
37
|
+
"zh": {
|
|
38
|
+
"learner": {
|
|
39
|
+
"education_level": "覆盖层级",
|
|
40
|
+
"major": "岗位序列",
|
|
41
|
+
"prior_knowledge": "先验经验",
|
|
42
|
+
"special_characteristics": "人群特征"
|
|
43
|
+
},
|
|
44
|
+
"course": {
|
|
45
|
+
"subject": "业务域",
|
|
46
|
+
"course_type": "组织形态",
|
|
47
|
+
"duration": "实施周期"
|
|
48
|
+
},
|
|
49
|
+
"intervention": {
|
|
50
|
+
"teaching_method": "干预方式",
|
|
51
|
+
"ai_tool": "技术工具",
|
|
52
|
+
"allowed_usage": "允许使用",
|
|
53
|
+
"frequency": "使用频率",
|
|
54
|
+
"duration": "干预周期",
|
|
55
|
+
"policy_name": "政策名称",
|
|
56
|
+
"policy_type": "政策类型",
|
|
57
|
+
"jurisdiction": "适用辖区",
|
|
58
|
+
"mechanism": "作用机制"
|
|
59
|
+
},
|
|
60
|
+
"population": {
|
|
61
|
+
"definition": "人群定义",
|
|
62
|
+
"size": "覆盖规模",
|
|
63
|
+
"vulnerability": "脆弱性"
|
|
64
|
+
},
|
|
65
|
+
"scope": {
|
|
66
|
+
"time_range": "时间范围",
|
|
67
|
+
"geography": "地域",
|
|
68
|
+
"study_types": "研究设计"
|
|
69
|
+
}
|
|
70
|
+
},
|
|
71
|
+
"en": {
|
|
72
|
+
"learner": {
|
|
73
|
+
"education_level": "Coverage level",
|
|
74
|
+
"major": "Role track",
|
|
75
|
+
"prior_knowledge": "Prior experience",
|
|
76
|
+
"special_characteristics": "Population characteristics"
|
|
77
|
+
},
|
|
78
|
+
"course": {
|
|
79
|
+
"subject": "Domain",
|
|
80
|
+
"course_type": "Organisation type",
|
|
81
|
+
"duration": "Rollout window"
|
|
82
|
+
},
|
|
83
|
+
"intervention": {
|
|
84
|
+
"teaching_method": "Method",
|
|
85
|
+
"ai_tool": "Technology",
|
|
86
|
+
"allowed_usage": "Allowed usage",
|
|
87
|
+
"frequency": "Frequency",
|
|
88
|
+
"duration": "Duration",
|
|
89
|
+
"policy_name": "Policy",
|
|
90
|
+
"policy_type": "Policy type",
|
|
91
|
+
"jurisdiction": "Jurisdiction",
|
|
92
|
+
"mechanism": "Mechanism"
|
|
93
|
+
},
|
|
94
|
+
"population": {
|
|
95
|
+
"definition": "Definition",
|
|
96
|
+
"size": "Scale",
|
|
97
|
+
"vulnerability": "Vulnerability"
|
|
98
|
+
},
|
|
99
|
+
"scope": {
|
|
100
|
+
"time_range": "Time range",
|
|
101
|
+
"geography": "Geography",
|
|
102
|
+
"study_types": "Study designs"
|
|
103
|
+
}
|
|
104
|
+
}
|
|
105
|
+
},
|
|
106
|
+
"evidence_detail_labels": {
|
|
107
|
+
"zh": {
|
|
108
|
+
"study_id": "研究 ID",
|
|
109
|
+
"sample_id": "样本 ID",
|
|
110
|
+
"title": "研究标题",
|
|
111
|
+
"year": "年份",
|
|
112
|
+
"study_type": "研究设计",
|
|
113
|
+
"education_level": "覆盖层级",
|
|
114
|
+
"population": "研究人群",
|
|
115
|
+
"sample_size": "样本量",
|
|
116
|
+
"intervention": "干预方案",
|
|
117
|
+
"comparison": "对照 / 比较条件",
|
|
118
|
+
"outcome_measure": "结果测量",
|
|
119
|
+
"claim": "完整主张",
|
|
120
|
+
"effect": "效应 / 结果",
|
|
121
|
+
"effect_direction": "效应方向",
|
|
122
|
+
"relation_to_claim": "与主张关系",
|
|
123
|
+
"duration": "干预时长",
|
|
124
|
+
"method": "方法",
|
|
125
|
+
"strengths": "优势",
|
|
126
|
+
"limitations": "局限",
|
|
127
|
+
"confounders": "混杂因素",
|
|
128
|
+
"quality_dimensions": "质量维度",
|
|
129
|
+
"quality_score": "质量分",
|
|
130
|
+
"evidence_level": "证据等级",
|
|
131
|
+
"directness": "直接性",
|
|
132
|
+
"applicability": "适用性",
|
|
133
|
+
"confidence": "置信度",
|
|
134
|
+
"status": "证据状态",
|
|
135
|
+
"source_location": "来源位置",
|
|
136
|
+
"source_title": "来源标题",
|
|
137
|
+
"source_url": "可验证链接",
|
|
138
|
+
"claim_id": "Claim ID"
|
|
139
|
+
},
|
|
140
|
+
"en": {
|
|
141
|
+
"study_id": "Study ID",
|
|
142
|
+
"sample_id": "Sample ID",
|
|
143
|
+
"title": "Study title",
|
|
144
|
+
"year": "Year",
|
|
145
|
+
"study_type": "Study design",
|
|
146
|
+
"education_level": "Coverage level",
|
|
147
|
+
"population": "Population",
|
|
148
|
+
"sample_size": "Sample size",
|
|
149
|
+
"intervention": "Intervention plan",
|
|
150
|
+
"comparison": "Comparison",
|
|
151
|
+
"outcome_measure": "Outcome measure",
|
|
152
|
+
"claim": "Full claim",
|
|
153
|
+
"effect": "Effect / result",
|
|
154
|
+
"effect_direction": "Effect direction",
|
|
155
|
+
"relation_to_claim": "Relation to claim",
|
|
156
|
+
"duration": "Duration",
|
|
157
|
+
"method": "Method",
|
|
158
|
+
"strengths": "Strengths",
|
|
159
|
+
"limitations": "Limitations",
|
|
160
|
+
"confounders": "Confounders",
|
|
161
|
+
"quality_dimensions": "Quality dimensions",
|
|
162
|
+
"quality_score": "Quality score",
|
|
163
|
+
"evidence_level": "Evidence level",
|
|
164
|
+
"directness": "Directness",
|
|
165
|
+
"applicability": "Applicability",
|
|
166
|
+
"confidence": "Confidence",
|
|
167
|
+
"status": "Evidence status",
|
|
168
|
+
"source_location": "Source location",
|
|
169
|
+
"source_title": "Source title",
|
|
170
|
+
"source_url": "Verifiable URL",
|
|
171
|
+
"claim_id": "Claim ID"
|
|
172
|
+
}
|
|
173
|
+
}
|
|
174
|
+
}
|