feat: 初始化 CS 知识应试强化系统
- 纯 HTML/CSS/JS 单页应用,支持填空题和选择题交互 - 5 个主题共 100 道题(GC/JVM、AI 结构化输出、记忆管理、幻觉、工具链) - 按主题组织、按题型拆分的 JSON 数据结构 - JSON Schema 校验 + 提示词模板 + 题型模板 - 自定义导入功能(粘贴/上传 JSON) - 项目文档(需求 + 架构) Co-Authored-By: Claude Code <noreply@anthropic.com>
This commit is contained in:
@@ -0,0 +1,188 @@
|
||||
{
|
||||
"topic": "ai-hallucination",
|
||||
"type": "fill_blank",
|
||||
"schema_version": "1.0.0",
|
||||
"generated": "2026-09-02T00:00:00Z",
|
||||
"questions": [
|
||||
{
|
||||
"id": "fb-001",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "在 Critic 架构流程中,执行 Agent 产出结果后,Critic Agent 提供评分和反馈,若未达标则反馈回 Agent 进行______。",
|
||||
"explanation": "Critic 流程中,未达标时反馈返回给 Agent 修改后重新评估。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"修改"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
},
|
||||
{
|
||||
"id": "fb-002",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "Critic 评估维度中,准确性(Accuracy)的权重是______。",
|
||||
"explanation": "准确性(事实正确性)在所有评估维度中权重最高,为 0.3。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"0.3"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
},
|
||||
{
|
||||
"id": "fb-003",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "三层评估体系由自动评估、AI 评估和______评估组成。",
|
||||
"explanation": "三层评估:自动(Schema/测试/Lint)、AI(Critic 评分)、人工(Code Review/验收)。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"人工"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
},
|
||||
{
|
||||
"id": "fb-004",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "评分实现中,加权总分的通过阈值是______。",
|
||||
"explanation": "评分系统要求加权总分 ≥ 0.8 才算通过。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"0.8"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
},
|
||||
{
|
||||
"id": "fb-005",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "当 AI 生成与客观事实不符的内容(如编造不存在的 API),这叫做______幻觉。",
|
||||
"explanation": "事实性幻觉指生成与客观事实不符的内容,如编造不存在的函数或 API。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"事实性"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
},
|
||||
{
|
||||
"id": "fb-006",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "Context Engineering 的核心思想是用精准、充分的上下文来减少模型的______空间。",
|
||||
"explanation": "Context Engineering 旨在通过提供精准上下文来最小化模型的猜测空间。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"猜测"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
},
|
||||
{
|
||||
"id": "fb-007",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "上下文组成包括:System Prompt、Knowledge Base、History 和______。",
|
||||
"explanation": "四大组件:系统提示、知识库、对话历史和检索到的上下文(RAG)。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"Retrieved Context"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
},
|
||||
{
|
||||
"id": "fb-008",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "Sub-agent Adversarial Verify 中使用______投票来决定是否接受输出。",
|
||||
"explanation": "系统使用多数投票:多数验证 Agent 通过则接受,否则重新生成。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"多数"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
},
|
||||
{
|
||||
"id": "fb-009",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "验证提示要求 Agent 对每个声明给出的裁决是:CONFIRMED、REFUTED 或______。",
|
||||
"explanation": "每个声明的裁决为三选一:CONFIRMED(确认)、REFUTED(反驳)、UNCERTAIN(不确定)。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"UNCERTAIN"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
},
|
||||
{
|
||||
"id": "fb-010",
|
||||
"type": "fill_blank",
|
||||
"difficulty": 2,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "RAG 的全称是 Retrieval ______ Generation。",
|
||||
"explanation": "RAG 全称 Retrieval Augmented Generation(检索增强生成)。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"answer": [
|
||||
"Augmented"
|
||||
],
|
||||
"answer_rule": "any"
|
||||
}
|
||||
]
|
||||
}
|
||||
@@ -0,0 +1,28 @@
|
||||
{
|
||||
"slug": "ai-hallucination",
|
||||
"name": "AI 幻觉 / 目标评价 / 规范",
|
||||
"description": "幻觉分类、Critic 架构、Context Engineering、评估体系",
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic",
|
||||
"context-engineering",
|
||||
"rag"
|
||||
],
|
||||
"difficulty_range": [
|
||||
1,
|
||||
5
|
||||
],
|
||||
"schema_version": "1.0.0",
|
||||
"question_files": [
|
||||
"fill_blank",
|
||||
"single_choice"
|
||||
],
|
||||
"stats": {
|
||||
"total": 20,
|
||||
"by_type": {
|
||||
"fill_blank": 10,
|
||||
"single_choice": 10
|
||||
}
|
||||
}
|
||||
}
|
||||
@@ -0,0 +1,218 @@
|
||||
{
|
||||
"topic": "ai-hallucination",
|
||||
"type": "single_choice",
|
||||
"schema_version": "1.0.0",
|
||||
"generated": "2026-09-02T00:00:00Z",
|
||||
"questions": [
|
||||
{
|
||||
"id": "sc-001",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "Critic 评估维度中,完整性(Completeness)的权重是多少?",
|
||||
"explanation": "完整性的权重是 0.25。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "0.15",
|
||||
"B": "0.2",
|
||||
"C": "0.25",
|
||||
"D": "0.3"
|
||||
},
|
||||
"answer": "C"
|
||||
},
|
||||
{
|
||||
"id": "sc-002",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "以下哪项不属于自动评估层?",
|
||||
"explanation": "Critic Agent 评分属于 AI 评估层,不属于自动评估层。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "Schema 校验",
|
||||
"B": "单测通过率",
|
||||
"C": "Critic Agent 评分",
|
||||
"D": "Lint 零告警"
|
||||
},
|
||||
"answer": "C"
|
||||
},
|
||||
{
|
||||
"id": "sc-003",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "当 AI 生成内容与输入上下文不一致时(如文档说 A,回答说 B),属于哪种幻觉?",
|
||||
"explanation": "忠实性幻觉指生成内容与输入上下文不一致。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "事实性幻觉",
|
||||
"B": "忠实性幻觉",
|
||||
"C": "推理性幻觉",
|
||||
"D": "指令性幻觉"
|
||||
},
|
||||
"answer": "B"
|
||||
},
|
||||
{
|
||||
"id": "sc-004",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "以下哪项不是上下文组成的组成部分?",
|
||||
"explanation": "模型参数不是上下文组成部分。上下文由 System Prompt、Knowledge Base、History 和 Retrieved Context 组成。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "System Prompt",
|
||||
"B": "Knowledge Base",
|
||||
"C": "Model Parameters(模型参数)",
|
||||
"D": "Retrieved Context"
|
||||
},
|
||||
"answer": "C"
|
||||
},
|
||||
{
|
||||
"id": "sc-005",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "Hook 机制中,自动化检查的正确顺序是?",
|
||||
"explanation": "正确顺序:Pre-commit Hook(类型检查)→ Lint Hook → 单测 Hook → 合并。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "Lint → 类型检查 → 单测 → 合并",
|
||||
"B": "类型检查 → Lint → 单测 → 合并",
|
||||
"C": "单测 → 类型检查 → Lint → 合并",
|
||||
"D": "类型检查 → 单测 → Lint → 合并"
|
||||
},
|
||||
"answer": "B"
|
||||
},
|
||||
{
|
||||
"id": "sc-006",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "注入领域知识作为 Skill 的目的是什么?",
|
||||
"explanation": "Skill 在 AI 不了解项目特定技术栈或业务领域时注入领域知识。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "加速模型推理",
|
||||
"B": "减少 token 使用",
|
||||
"C": "帮助不了解项目特定技术栈或业务领域的 AI",
|
||||
"D": "替代人类开发者"
|
||||
},
|
||||
"answer": "C"
|
||||
},
|
||||
{
|
||||
"id": "sc-007",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "RAG 流程中,嵌入和向量检索之后的下一步是什么?",
|
||||
"explanation": "RAG 流程:问题 → 嵌入 → 向量检索 → Top-K 文档 → 上下文组装 → LLM → 回答。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "LLM 生成回答",
|
||||
"B": "Top-K 文档检索",
|
||||
"C": "上下文组装",
|
||||
"D": "基于文档的回答"
|
||||
},
|
||||
"answer": "B"
|
||||
},
|
||||
{
|
||||
"id": "sc-008",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "CLAUDE.md 示例中,接口使用什么前缀命名?",
|
||||
"explanation": "CLAUDE.md 示例规定接口使用 I 前缀命名。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "_interface",
|
||||
"B": "I",
|
||||
"C": "Int",
|
||||
"D": "Interface"
|
||||
},
|
||||
"answer": "B"
|
||||
},
|
||||
{
|
||||
"id": "sc-009",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "统一返回类使用什么格式?",
|
||||
"explanation": "项目规范统一使用 Result<T> 包装:{code, message, data}。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "{status, error, result}",
|
||||
"B": "{success, message, payload}",
|
||||
"C": "{code, message, data}",
|
||||
"D": "{result, description, body}"
|
||||
},
|
||||
"answer": "C"
|
||||
},
|
||||
{
|
||||
"id": "sc-010",
|
||||
"type": "single_choice",
|
||||
"difficulty": 3,
|
||||
"tags": [
|
||||
"ai",
|
||||
"hallucination",
|
||||
"critic"
|
||||
],
|
||||
"question": "在 AI 辅助的瀑布模型中,进入下一阶段前必须完成什么?",
|
||||
"explanation": "每个阶段需要人工验收后才能进入下一阶段。",
|
||||
"source": null,
|
||||
"related": [],
|
||||
"options": {
|
||||
"A": "AI 自验证",
|
||||
"B": "自动化测试",
|
||||
"C": "人工验收",
|
||||
"D": "同行评审"
|
||||
},
|
||||
"answer": "C"
|
||||
}
|
||||
]
|
||||
}
|
||||
Reference in New Issue
Block a user