{
  "video_id": "2607.01153",
  "channel": "cs.AI",
  "title": "Adversarial Pragmatics for AI Safety Evaluation: A Benchmark for Instruction Conflict, Embedded Commands, and Policy Ambiguity",
  "scores": {
    "depth": 3,
    "novelty": 3,
    "aria_relevance": 2,
    "production_ready": 2
  },
  "junk_penalty": 0,
  "avg_score": 2.5,
  "effective_score": 2.5,
  "verdict": "promote",
  "one_line_reason": "Rigoroses Benchmark-Design für Safety-Evaluation mit linguistischer Kontrolle und Metodologie für Ambiguität-Diagnose; starke Novelty und Depth, aber Aria-Relevanz eher indirekt (Safety-Eval als Enabler für Agent-Robustheit).",
  "model": "claude-haiku-4-5-20251001",
  "cost_usd": 0.001862,
  "triaged_at": "2026-07-02T15:00:11.709394+00:00"
}