mirror of
https://github.com/rohitg00/ai-engineering-from-scratch.git
synced 2026-10-02 01:54:39 +08:00
feat(phase-18/29): add quiz.json
This commit is contained in:
+78
@@ -0,0 +1,78 @@
|
||||
{
|
||||
"lesson": "29-moderation-systems-openai-perspective-llamaguard",
|
||||
"title": "Moderation Systems - OpenAI, Perspective, Llama Guard",
|
||||
"questions": [
|
||||
{
|
||||
"stage": "pre",
|
||||
"question": "What is OpenAI's omni-moderation-latest (2024)?",
|
||||
"options": [
|
||||
"A reasoning-focused frontier model",
|
||||
"A GPT-4o-based moderation classifier that handles text and images in one call, returning a 13-category response schema (harassment, hate, self-harm, sexual, violence, illicit, with sub-categories), free for most developers",
|
||||
"An RLHF reward model",
|
||||
"A C2PA watermarking endpoint"
|
||||
],
|
||||
"correct": 1,
|
||||
"explanation": ""
|
||||
},
|
||||
{
|
||||
"stage": "check",
|
||||
"question": "Which of these is NOT one of OpenAI Moderation's primary top-level categories?",
|
||||
"options": [
|
||||
"harassment",
|
||||
"self-harm",
|
||||
"election-interference",
|
||||
"violence"
|
||||
],
|
||||
"correct": 2,
|
||||
"explanation": ""
|
||||
},
|
||||
{
|
||||
"stage": "check",
|
||||
"question": "What are the three layers of the standard 2026 moderation pattern?",
|
||||
"options": [
|
||||
"Tokenizer / model / detokenizer",
|
||||
"Input moderation (pre-generation), Output moderation (post-generation), Custom moderation (domain rules)",
|
||||
"Probe / detector / harness",
|
||||
"Telescopic / periscopic / microscopic"
|
||||
],
|
||||
"correct": 1,
|
||||
"explanation": ""
|
||||
},
|
||||
{
|
||||
"stage": "check",
|
||||
"question": "What is Perspective API (Google Jigsaw) primarily designed to score?",
|
||||
"options": [
|
||||
"Reasoning quality",
|
||||
"Toxicity (with severe-toxicity, insult, profanity, threat, identity-attack sub-dimensions) as a pre-LLM-era content-moderation baseline",
|
||||
"Watermark strength",
|
||||
"Reward over user satisfaction"
|
||||
],
|
||||
"correct": 1,
|
||||
"explanation": ""
|
||||
},
|
||||
{
|
||||
"stage": "post",
|
||||
"question": "What is one failure mode of input-only moderation?",
|
||||
"options": [
|
||||
"It always blocks legitimate content",
|
||||
"It does not catch model output failures or hallucinations; encoding attacks (Lessons 12-14) can bypass input classifiers entirely",
|
||||
"It cannot run on GPUs",
|
||||
"It removes the need for output classifiers"
|
||||
],
|
||||
"correct": 1,
|
||||
"explanation": ""
|
||||
},
|
||||
{
|
||||
"stage": "post",
|
||||
"question": "What is the Azure Content Moderator deprecation and migration target?",
|
||||
"options": [
|
||||
"Deprecated February 2024, retired February 2027; migration target is Azure AI Content Safety (LLM-based, integrated with Azure OpenAI)",
|
||||
"Deprecated immediately with no replacement",
|
||||
"Replaced by Perspective API",
|
||||
"Merged into Llama Guard 4"
|
||||
],
|
||||
"correct": 0,
|
||||
"explanation": ""
|
||||
}
|
||||
]
|
||||
}
|
||||
Reference in New Issue
Block a user