{
 "generated_at_utc": "2026-09-15T16:32:08Z",
 "escalations": [
  {
   "id": "esc-1",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-2",
   "reason": "passages_drifted",
   "question": "2 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-3",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5385 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-4",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4118 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-5",
   "reason": "cross_paper_overlap",
   "question": "1 claim(s) in 2606.19887 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-6",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-7",
   "reason": "checker_citation_not_found",
   "question": "1 of 23 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-8",
   "reason": "grounded_recovery_low",
   "question": "Only 0.6667 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-9",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.6875 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-10",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Safety alignment suppresses engagement with sensitive knowledge, making it difficult for models to identify and reason about multimodal risk elements.\" already answered by prior work? Retrieval returned 9 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-11",
   "reason": "checker_citation_not_found",
   "question": "2 of 26 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-12",
   "reason": "grounded_recovery_low",
   "question": "Only 0.6111 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-13",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.6111 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-14",
   "reason": "cross_paper_overlap",
   "question": "3 claim(s) in 2606.09711 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-15",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c1 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-16",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c2 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-17",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c3 names \"Source B\", \"Source C\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-18",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c4 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-19",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c5 names \"PRIME\", \"Source A\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-20",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c6 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-21",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c7 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-22",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c9 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-23",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c11 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-24",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c12 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-25",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c13 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-26",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c16 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-27",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c17 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-28",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c18 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-29",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c19 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-30",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c20 names \"100\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-31",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.09711/c22 names \"PRIME\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-32",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-33",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-34",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5217 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-35",
   "reason": "cross_paper_overlap",
   "question": "8 claim(s) in 2606.02630 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-36",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.02630/c3 names \"four\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-37",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.02630/c14 names \"six\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-38",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The paper characterizes four degradation trajectory signatures (Compliance Creep, Diminishing Returns, Pattern Recognition, Spike-and-Abandonment) that describe\" already answered by prior work? Retrieval returned 18 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-39",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Turn 2 is the critical vulnerability window for safety intervention in multi-turn medical conversations.\" already answered by prior work? Retrieval returned 23 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-40",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-41",
   "reason": "grounded_recovery_low",
   "question": "Only 0.6111 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-42",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5714 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-43",
   "reason": "cross_paper_overlap",
   "question": "15 claim(s) in 2512.00349 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-44",
   "reason": "cross_paper_overlap",
   "question": "8 claim(s) in 2606.17478 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-45",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.17478/c4 names \"GPT-OSS-20B\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-46",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.17478/c11 names \"LatentQA\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-47",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Threshold OR ensembles of monitor families reduce false negatives but raise realized Alpaca-control FPR.\" already answered by prior work? Retrieval returned 13 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-48",
   "reason": "passages_drifted",
   "question": "2 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-49",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4167 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-50",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.3684 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-51",
   "reason": "cross_paper_overlap",
   "question": "7 claim(s) in 2606.10747 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-52",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Instruction-induced misalignment produces salient behavioral cues: when the fine-tuned model organism is paired with a risky system prompt, pure observation alr\" already answered by prior work? Retrieval returned 12 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-53",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-54",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4444 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-55",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4167 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-56",
   "reason": "cross_paper_overlap",
   "question": "14 claim(s) in 2509.02655 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-57",
   "reason": "grounded_recovery_low",
   "question": "Only 0.625 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-58",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5263 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-59",
   "reason": "cross_paper_overlap",
   "question": "12 claim(s) in 2606.08682 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-60",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.08682/c12 names \"AS-induced EM\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-61",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.08682/c13 names \"AS-induced EM\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-62",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Among the tested injection layer groups on Qwen3.5-27B, layers 22-25 give the highest EM rate, while injecting into layers 24-25 yields near-zero emergent misal\" already answered by prior work? Retrieval returned 12 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-63",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-64",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-65",
   "reason": "checker_citation_not_found",
   "question": "1 of 23 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-66",
   "reason": "grounded_recovery_low",
   "question": "Only 0.6154 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-67",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4706 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-68",
   "reason": "cross_paper_overlap",
   "question": "17 claim(s) in 2606.28863 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-69",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.28863/c12 names \"RLHF\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-70",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.28863/c19 names \"three\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-71",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.28863/c22 names \"four\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-72",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The paper claims the triadic test partitions cases into in-class and out-of-class: an honest safety filter and incidental distribution shift fall outside, while\" already answered by prior work? Retrieval returned 12 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-73",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The paper reports that among its thirty documented cases no two share the same (trigger, swap, origin) triple and that cases distribute across twenty-two of the\" already answered by prior work? Retrieval returned 12 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-74",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The paper reports that of its thirty documented cases, twelve are upward swaps, twelve are downward, and six are lateral (persona switch), and that the default \" already answered by prior work? Retrieval returned 12 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-75",
   "reason": "grounded_recovery_low",
   "question": "Only 0.48 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-76",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.48 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-77",
   "reason": "cross_paper_overlap",
   "question": "27 claim(s) in 2605.24197 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-78",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2605.24197/c1 names \"Multi-agent\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-79",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2605.24197/c12 names \"Multi-agent\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-80",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2605.24197/c20 names \"RewardBench\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-81",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2605.24197/c27 names \"six\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-82",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Under epsilon-close priors and likelihoods and a sufficiently informative evidence lower bound, role posteriors remain delta-close; consequently, without distin\" already answered by prior work? Retrieval returned 18 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-83",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-84",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-85",
   "reason": "checker_citation_not_found",
   "question": "20 of 137 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-86",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5926 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-87",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5714 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-88",
   "reason": "cross_paper_overlap",
   "question": "10 claim(s) in 2603.00829 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-89",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2603.00829/c7 names \"Kimi K2\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-90",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2603.00829/c12 names \"LLMs\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-91",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Attempts to improve on grid-search-selected prompts via additional iterative refinement (human or automated) generally do not yield further gains and instead in\" already answered by prior work? Retrieval returned 17 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-92",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"A grid search over 3 candidate models and 15 candidate prompts yields monitors with test-set partial AUROC of 0.853 (Gloom) and 0.866 (STRIDE).\" already answered by prior work? Retrieval returned 11 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-93",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Human-guided prompt refinement on STRIDE yields a statistically significant improvement over the best prompt-sweep prompt, an isolated exception to the general \" already answered by prior work? Retrieval returned 9 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-94",
   "reason": "checker_citation_not_found",
   "question": "3 of 45 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-95",
   "reason": "grounded_recovery_low",
   "question": "Only 0.3571 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-96",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.2632 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-97",
   "reason": "cross_paper_overlap",
   "question": "11 claim(s) in 2606.11409 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-98",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.11409/c5 names \"HarmBench\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-99",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.11409/c7 names \"PAIR\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-100",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.11409/c8 names \"Qwen3-4B-SafeRL\", \"50\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-101",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Gradient-based GCG suffixes optimized on an open-weight surrogate (Qwen2.5-0.5B-Instruct) can transfer to a separate target model (Qwen3-8B), eliciting non-triv\" already answered by prior work? Retrieval returned 16 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-102",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-103",
   "reason": "checker_contradicted",
   "question": "Claim \"Safety-aligned RL on Qwen3-4B raises aggregate adversarial compute cost for JailBroken and PAIR while leaving some harm categories disproportionately exploitabl\" was judged contradicted by a blinded checker reading only the passages cited for it. Decide whether it should remain in the graph.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-104",
   "reason": "checker_citation_not_found",
   "question": "1 of 22 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-105",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4167 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-106",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.3571 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-107",
   "reason": "cross_paper_overlap",
   "question": "13 claim(s) in 2606.20626 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-108",
   "reason": "cross_paper_overlap",
   "question": "15 claim(s) in 2501.14940 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-109",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2501.14940/c6 names \"Contextual Integrity\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-110",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2501.14940/c15 names \"CASE-Bench\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-111",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Each task (query-context pair) was annotated by 21 annotators, a number determined by statistical power analysis, and the dataset contains 47,000+ human annotat\" already answered by prior work? Retrieval returned 16 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-112",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The paper applies Contextual Integrity (CI) theory parameters to formalize context, describing this as the first instance of using CI theory to build a foundati\" already answered by prior work? Retrieval returned 17 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-113",
   "reason": "passage_fabricated",
   "question": "1 passage(s) the reader cited do not appear in the source paper at all. Is the extractor inventing text, or is the PDF extraction mangling honest quotes?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-114",
   "reason": "passages_drifted",
   "question": "2 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-115",
   "reason": "checker_citation_not_found",
   "question": "1 of 25 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-116",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5714 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-117",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.6471 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-118",
   "reason": "cross_paper_overlap",
   "question": "16 claim(s) in 2606.00027 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-119",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The highest-scoring domains were Safety & Reliability and Medical Errors (each averaging around 0.96), while Bias, Fairness & Equity (0.95 ± 0.04 SD) and Clinic\" already answered by prior work? Retrieval returned 12 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-120",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Operationally complex categories including Liability, Accountability, and Medical Coding & Billing were the most challenging (domain means between 0.79 and 0.83\" already answered by prior work? Retrieval returned 19 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-121",
   "reason": "grounded_recovery_low",
   "question": "Only 0.6667 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-122",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5455 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-123",
   "reason": "cross_paper_overlap",
   "question": "18 claim(s) in 2606.04435 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-124",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.04435/c3 names \"Confidence Inflation Cascade\", \"Context Poisoning\", \"Inference Cascade\", \"Poisoning Cascade\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-125",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.04435/c19 names \"HITL-AP\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-126",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.04435/c21 names \"Confidence Inflation Cascade\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-127",
   "reason": "passage_fabricated",
   "question": "4 passage(s) the reader cited do not appear in the source paper at all. Is the extractor inventing text, or is the PDF extraction mangling honest quotes?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-128",
   "reason": "passages_drifted",
   "question": "2 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-129",
   "reason": "checker_citation_not_found",
   "question": "6 of 48 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-130",
   "reason": "grounded_recovery_low",
   "question": "Only 0.375 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-131",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.3913 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-132",
   "reason": "cross_paper_overlap",
   "question": "12 claim(s) in 2606.05391 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-133",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.05391/c8 names \"Co-planning\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-134",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.05391/c11 names \"two\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-135",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-136",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4118 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-137",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.3889 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-138",
   "reason": "cross_paper_overlap",
   "question": "13 claim(s) in 2606.12918 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-139",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.12918/c2 names \"Shapley-guided\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-140",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.12918/c4 names \"2\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-141",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.12918/c6 names \"Agent-level Shapley\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-142",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The paper designs a closed-loop, Shapley-guided autonomous red-teaming agent that selects a coalition of agents and jointly generates coordinated, role-aware ad\" already answered by prior work? Retrieval returned 22 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-143",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Prior red-teaming methods (TAMAS, GCA, AutoTransform, AiTM) yield near-zero attack success rates in most settings on hierarchical MAS, especially under limited \" already answered by prior work? Retrieval returned 14 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-144",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Agent-level Shapley value distributions are highly skewed and task-dependent: only a small subset of agents contributes significantly to attack success, and the\" already answered by prior work? Retrieval returned 10 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-145",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4615 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-146",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4286 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-147",
   "reason": "cross_paper_overlap",
   "question": "14 claim(s) in 2606.05566 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-148",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.05566/c8 names \"JBB-Behaviors\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-149",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The system operates with an average latency of approximately 50 ms on CPU, making it suitable for production deployment under cost and infrastructure constraint\" already answered by prior work? Retrieval returned 13 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-150",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The protectai-v2 model (184M parameters) achieves a perfect F1 of 1.000 on the awall-test benchmark but collapses to F1 = 0.000 on the unseen JBB-Behaviors pool\" already answered by prior work? Retrieval returned 16 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-151",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-152",
   "reason": "grounded_recovery_low",
   "question": "Only 0.375 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-153",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.2609 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-154",
   "reason": "cross_paper_overlap",
   "question": "14 claim(s) in 2606.05233 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-155",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.05233/c1 names \"158\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-156",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.05233/c4 names \"Anthropic\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-157",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.05233/c10 names \"17\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-158",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-159",
   "reason": "checker_citation_not_found",
   "question": "3 of 32 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-160",
   "reason": "grounded_recovery_low",
   "question": "Only 0.3125 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-161",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.3125 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-162",
   "reason": "cross_paper_overlap",
   "question": "16 claim(s) in 2606.03810 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-163",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.03810/c17 names \"three\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-164",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The paper reports that the regularization methods ACT and BCT produce larger effects than label-generation methods, strongly suppressing reward hacking and emer\" already answered by prior work? Retrieval returned 13 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-165",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-166",
   "reason": "passages_drifted",
   "question": "4 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-167",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5385 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-168",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.45 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-169",
   "reason": "cross_paper_overlap",
   "question": "17 claim(s) in 2606.01322 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-170",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.01322/c1 names \"TukaBench\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-171",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.01322/c5 names \"Afri-JBB-Culture\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-172",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.01322/c6 names \"Code-switched\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-173",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.01322/c7 names \"Boundary Point Jailbreaking\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-174",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.01322/c11 names \"EFUSED\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-175",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.01322/c13 names \"JAILBROKEN\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-176",
   "reason": "passages_drifted",
   "question": "4 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-177",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5333 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-178",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4737 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-179",
   "reason": "cross_paper_overlap",
   "question": "27 claim(s) in 2506.04018 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-180",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2506.04018/c3 names \"two\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-181",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Two conditions are required for behavior to qualify as misaligned: acting contrary to the deployer's intended goals (rather than following malicious instruction\" already answered by prior work? Retrieval returned 12 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-182",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Most evaluations use the pre-built InspectAI basic agent, a simple ReAct loop with task-specific tools and a reflective prompt.\" already answered by prior work? Retrieval returned 20 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-183",
   "reason": "grounded_recovery_low",
   "question": "Only 0.6333 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-184",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.6333 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-185",
   "reason": "cross_paper_overlap",
   "question": "17 claim(s) in 2606.24081 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-186",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.24081/c4 names \"eleven\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-187",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Methods that must be reconstructed primarily from paper text (PGJ, R2A, Low-Effort) show larger deviations, with PGJ at 7.2% error and R2A at 16.1% error, attri\" already answered by prior work? Retrieval returned 16 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-188",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-189",
   "reason": "passages_drifted",
   "question": "2 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-190",
   "reason": "grounded_recovery_low",
   "question": "Only 0.3529 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-191",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.3158 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-192",
   "reason": "cross_paper_overlap",
   "question": "13 claim(s) in 2604.23130 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-193",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2604.23130/c2 names \"Single-token-driven\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-194",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2604.23130/c7 names \"17\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-195",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Hierarchical-linkage steering is the most selective and least effective of the three strategies, because its cluster-size constraint (merged cluster at most 50 \" already answered by prior work? Retrieval returned 9 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-196",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The harm-responsible features are largely prompt-specific: 17.4% of steered responses on original adversarial prompts received a higher harmfulness score than t\" already answered by prior work? Retrieval returned 16 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-197",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Among responses that began as non-harmful content (default score 1), 3.70% were driven to maximal harm (score 5) and a further 1.10% to score 4, so 4.8% of non-\" already answered by prior work? Retrieval returned 16 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-198",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4615 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-199",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.3333 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-200",
   "reason": "cross_paper_overlap",
   "question": "18 claim(s) in 2606.24014 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-201",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.24014/c1 names \"80\", \"50\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-202",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.24014/c14 names \"twelve\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-203",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Across the evaluated OpenAI models, alignment evaluation scores show weak positive cross-model correlation (mean Spearman's rho = 0.107) and the first principal\" already answered by prior work? Retrieval returned 20 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-204",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-205",
   "reason": "checker_citation_not_found",
   "question": "1 of 33 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-206",
   "reason": "grounded_recovery_low",
   "question": "Only 0.2941 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-207",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.25 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-208",
   "reason": "cross_paper_overlap",
   "question": "15 claim(s) in 2601.19072 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-209",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2601.19072/c10 names \"97\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-210",
   "reason": "checker_contradicted",
   "question": "Claim \"Tree of thought is consistently the top-performing strategy, direct assessment is second best, multi-step reasoning and few-shot achieve lower performance, and \" was judged contradicted by a blinded checker reading only the passages cited for it. Decide whether it should remain in the graph.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-211",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4667 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-212",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4667 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-213",
   "reason": "cross_paper_overlap",
   "question": "15 claim(s) in 2602.02557 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-214",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2602.02557/c15 names \"Text-transferred\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-215",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Text attacks achieve the highest average StrongReject (SR) score across the evaluated omni-models, revealing a text-centric vulnerability.\" already answered by prior work? Retrieval returned 10 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-216",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-217",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-218",
   "reason": "cross_paper_overlap",
   "question": "18 claim(s) in 2606.28332 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-219",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.28332/c2 names \"General-purpose\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-220",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.28332/c15 names \"Safe Helpfulness\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-221",
   "reason": "passages_drifted",
   "question": "4 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-222",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5714 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-223",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4545 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-224",
   "reason": "cross_paper_overlap",
   "question": "22 claim(s) in 2606.07631 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-225",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.07631/c2 names \"EM-relevant\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-226",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Alignment-relevant trait directions are necessary for low-FNR detection: the alignment feature set reaches 2.2% FNR under RF while semantic and random 7D contro\" already answered by prior work? Retrieval returned 17 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-227",
   "reason": "passage_fabricated",
   "question": "1 passage(s) the reader cited do not appear in the source paper at all. Is the extractor inventing text, or is the PDF extraction mangling honest quotes?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-228",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-229",
   "reason": "checker_citation_not_found",
   "question": "1 of 44 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-230",
   "reason": "grounded_recovery_low",
   "question": "Only 0.2857 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-231",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.3043 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-232",
   "reason": "cross_paper_overlap",
   "question": "14 claim(s) in 2603.26846 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-233",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2603.26846/c2 names \"two\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-234",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2603.26846/c8 names \"Instruct\", \"Llama-3\", \"3.3\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-235",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2603.26846/c9 names \"Gaussian\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-236",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2603.26846/c12 names \"Lagrange\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-237",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Among the four stability metrics, semantic entropy (SE) maintains the most consistent separability across CoT and Response, whereas PPL, Pmax, and Cosine Sim ar\" already answered by prior work? Retrieval returned 15 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-238",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"CoT Monitor induces obfuscated reward hacking, paradoxically worsening Actual Deception while collapsing CoT Faithfulness.\" already answered by prior work? Retrieval returned 15 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-239",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"SAR retains general model capability, performing within normal fluctuation ranges and avoiding alignment tax or capability collapse.\" already answered by prior work? Retrieval returned 11 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-240",
   "reason": "critic_high_severity",
   "question": "2 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-241",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-242",
   "reason": "checker_citation_not_found",
   "question": "2 of 21 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-243",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5385 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-244",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-245",
   "reason": "cross_paper_overlap",
   "question": "13 claim(s) in 2505.14289 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-246",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Effective adversarial semantics form a dense, continuous 'semantic attack space' in the model's latent representation rather than isolated sparse points, which \" already answered by prior work? Retrieval returned 15 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-247",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"An 'alignment paradox' exists: models with more extensive alignment training sometimes show increased vulnerability to EVA's attacks, because alignment training\" already answered by prior work? Retrieval returned 13 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-248",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Successful adversarial payloads are not uniformly distributed over persuasion dimensions but concentrate on two attractors—trust-aligned and urgency-aligned sem\" already answered by prior work? Retrieval returned 13 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-249",
   "reason": "passages_drifted",
   "question": "2 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-250",
   "reason": "checker_citation_not_found",
   "question": "1 of 24 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-251",
   "reason": "grounded_recovery_low",
   "question": "Only 0.7273 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-252",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.6923 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-253",
   "reason": "cross_paper_overlap",
   "question": "10 claim(s) in 2606.15396 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-254",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.15396/c4 names \"three\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-255",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"MDPO dynamically adjusts the KL penalty coefficient based on the policy model's real-time responsiveness to sample difficulty, using normalized reward gaps, out\" already answered by prior work? Retrieval returned 15 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-256",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-257",
   "reason": "passages_drifted",
   "question": "2 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-258",
   "reason": "grounded_recovery_low",
   "question": "Only 0.375 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-259",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.2667 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-260",
   "reason": "cross_paper_overlap",
   "question": "15 claim(s) in 2606.00033 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-261",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5882 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-262",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4167 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-263",
   "reason": "cross_paper_overlap",
   "question": "14 claim(s) in 1805.03090 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-264",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 1805.03090/c5 names \"POMDPs\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-265",
   "reason": "grounded_recovery_low",
   "question": "Only 0.2667 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-266",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.2353 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-267",
   "reason": "cross_paper_overlap",
   "question": "19 claim(s) in 2604.26360 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-268",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2604.26360/c13 names \"Hopper-v4\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-269",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The reciprocal reliability filter is derived from risk-sensitive mean-variance utility and replaces an unbounded linear penalty that can become negative under h\" already answered by prior work? Retrieval returned 15 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-270",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The reciprocal reliability filter satisfies positivity, monotonicity, boundedness, identity at zero uncertainty, and Lipschitz continuity.\" already answered by prior work? Retrieval returned 15 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-271",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The magnitude of reward misspecification is assumed to be bounded by a monotonically increasing function of epistemic and aleatoric uncertainty.\" already answered by prior work? Retrieval returned 13 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-272",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-273",
   "reason": "checker_citation_not_found",
   "question": "1 of 37 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-274",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4211 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-275",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-276",
   "reason": "cross_paper_overlap",
   "question": "14 claim(s) in 2606.07706 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-277",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-278",
   "reason": "grounded_recovery_low",
   "question": "Only 0.7857 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-279",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.7333 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-280",
   "reason": "cross_paper_overlap",
   "question": "22 claim(s) in 2606.07612 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-281",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.07612/c16 names \"Fine-tuning\", \"Instruct\", \"Llama-3\", \"3.1\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-282",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.07612/c19 names \"3.15\", \"5.18\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-283",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-284",
   "reason": "checker_citation_not_found",
   "question": "3 of 56 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-285",
   "reason": "grounded_recovery_low",
   "question": "Only 0.2381 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-286",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.2273 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-287",
   "reason": "cross_paper_overlap",
   "question": "17 claim(s) in 2502.20914 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-288",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2502.20914/c13 names \"Gaussian\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-289",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"An exhaustive circuit-first pass on the illustrative example network (k = 3, n = 1, loss cutoff 10^-3) yielded 59 circuits and 114,230 interpretations.\" already answered by prior work? Retrieval returned 17 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-290",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4737 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-291",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4737 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-292",
   "reason": "cross_paper_overlap",
   "question": "12 claim(s) in 2604.24668 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-293",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2604.24668/c5 names \"50\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-294",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2604.24668/c9 names \"two\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-295",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-296",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4545 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-297",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.375 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-298",
   "reason": "cross_paper_overlap",
   "question": "17 claim(s) in 2606.18988 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-299",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.18988/c4 names \"VAC-GRPO\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-300",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.18988/c15 names \"four\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-301",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The authors construct Deception-10K, described as the first fine-grained audio-visual Chain-of-Thought dataset, comprising 10,000 video-reasoning pairs (~50 hou\" already answered by prior work? Retrieval returned 18 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-302",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The paper proposes Visual-Audio Consistency Group Relative Policy Optimization (VAC-GRPO) with a progressive training strategy that stratifies data into four di\" already answered by prior work? Retrieval returned 17 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-303",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Most existing baseline multimodal large language models perform around the random-guess baseline of 50% on deception detection despite identical prompts.\" already answered by prior work? Retrieval returned 13 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-304",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-305",
   "reason": "passage_fabricated",
   "question": "1 passage(s) the reader cited do not appear in the source paper at all. Is the extractor inventing text, or is the PDF extraction mangling honest quotes?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-306",
   "reason": "passages_drifted",
   "question": "3 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-307",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5385 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-308",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4706 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-309",
   "reason": "cross_paper_overlap",
   "question": "14 claim(s) in 2606.08451 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-310",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.08451/c8 names \"six\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-311",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.08451/c13 names \"38\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-312",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"In the most severe cases, models agree with harmful prompts over 70% of the time in zero-shot languages, defaulting to explicit agreement with safety-critical p\" already answered by prior work? Retrieval returned 20 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-313",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-314",
   "reason": "passages_drifted",
   "question": "2 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-315",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5833 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-316",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-317",
   "reason": "cross_paper_overlap",
   "question": "16 claim(s) in 2606.10106 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-318",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.10106/c8 names \"four\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-319",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.10106/c9 names \"four\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-320",
   "reason": "checker_citation_not_found",
   "question": "4 of 34 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-321",
   "reason": "grounded_recovery_low",
   "question": "Only 0.625 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-322",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-323",
   "reason": "cross_paper_overlap",
   "question": "19 claim(s) in 2606.07532 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-324",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.07532/c2 names \"Justice\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-325",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.07532/c6 names \"Experiment\", \"200\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-326",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.07532/c10 names \"ChatEval\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-327",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"A single-model ablation isolating identity stripping (Experiment 3 / pilot14) produces a directional accuracy gain of 6.0 percentage points over unstripped cont\" already answered by prior work? Retrieval returned 13 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-328",
   "reason": "checker_citation_not_found",
   "question": "2 of 59 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-329",
   "reason": "grounded_recovery_low",
   "question": "Only 0.6 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-330",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.48 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-331",
   "reason": "cross_paper_overlap",
   "question": "18 claim(s) in 2606.20814 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-332",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.20814/c2 names \"Instruct\", \"Qwen2\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-333",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The relationship between evaluation sample size and the Max Score Difference roughly follows a power law.\" already answered by prior work? Retrieval returned 14 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-334",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"The raw (non-JSON, non-template) format of the Initial EM questions is almost always the most misaligned format.\" already answered by prior work? Retrieval returned 14 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-335",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Even with variations in learning schedules, in-domain loss has a dominant effect on the level of misalignment in the paper's experiment setting.\" already answered by prior work? Retrieval returned 19 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-336",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Across the training-dynamics experiments summarized in Table 3, training loss still guides the level of misalignment.\" already answered by prior work? Retrieval returned 22 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-337",
   "reason": "passages_drifted",
   "question": "2 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-338",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5625 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-339",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5556 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-340",
   "reason": "cross_paper_overlap",
   "question": "14 claim(s) in 2606.26793 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-341",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.26793/c2 names \"Prior Sampling\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-342",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-343",
   "reason": "checker_citation_not_found",
   "question": "1 of 20 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-344",
   "reason": "grounded_recovery_low",
   "question": "Only 0.625 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-345",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.625 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-346",
   "reason": "cross_paper_overlap",
   "question": "8 claim(s) in 2602.18008 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-347",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2602.18008/c2 names \"Existing LLM-based\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-348",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Models generated by NIMMGen can be used for counterfactual intervention simulation: increasing simulated social distancing strength produces systematic reductio\" already answered by prior work? Retrieval returned 21 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-349",
   "reason": "checker_citation_not_found",
   "question": "3 of 22 quote(s) produced by the checker could not be located in the material it was shown. Its verdicts rest on text it may not have read correctly.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-350",
   "reason": "grounded_recovery_low",
   "question": "Only 0.75 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-351",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-352",
   "reason": "cross_paper_overlap",
   "question": "21 claim(s) in 2606.27188 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-353",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.27188/c1 names \"Agentic Business Process\", \"Business Process Management\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-354",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.27188/c5 names \"Model Context Protocol\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-355",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.27188/c16 names \"two\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-356",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.27188/c18 names \"CUGA FLO\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-357",
   "reason": "grounded_recovery_low",
   "question": "Only 0.3333 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-358",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.3333 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-359",
   "reason": "cross_paper_overlap",
   "question": "27 claim(s) in 2605.11047 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-360",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Qwen3.5-Plus, DeepSeek-v4-Flash, and DeepSeek-v4-Pro show consistently high AGS across the six risk categories, indicating the generated traps transfer beyond t\" already answered by prior work? Retrieval returned 11 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-361",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-362",
   "reason": "grounded_recovery_low",
   "question": "Only 0.4615 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-363",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.4444 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-364",
   "reason": "cross_paper_overlap",
   "question": "16 claim(s) in 2510.00845 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-365",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Bootstrap resampling of the input dataset yields the lowest structural consistency and highest variability of discovered circuits (Jaccard µ = 0.561, CV = 0.335\" already answered by prior work? Retrieval returned 16 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-366",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Circuits discovered under bootstrap resampling also have the highest average circuit error (0.440), meaning they are structurally different and less faithful to\" already answered by prior work? Retrieval returned 12 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-367",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Shifting the meta-distribution (meta-dataset or prompt paraphrasing) yields more stable circuits than bootstrap resampling, with higher Jaccard indices (0.790 a\" already answered by prior work? Retrieval returned 15 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-368",
   "reason": "prior_art_uncertain",
   "question": "Is claim \"Circuit discovery methods do not scale trivially: stability degrades for larger models, with gpt2-small yielding relatively clustered results while Llama-3.2 (1\" already answered by prior work? Retrieval returned 9 candidates and the adjudicator could not decide.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-369",
   "reason": "passages_drifted",
   "question": "1 passage(s) are neither verbatim nor invented: they open with real source text and then diverge into paraphrase. Should the extractor be required to quote contiguously, or should near-quotes be accepted with a marker?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-370",
   "reason": "grounded_recovery_low",
   "question": "Only 0.5333 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-371",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.5 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-372",
   "reason": "cross_paper_overlap",
   "question": "26 claim(s) in 2606.21399 closely restate a claim the society already holds from another paper. Is this an independent confirmation, a duplicate that should be merged, or a contradiction that should be recorded?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-373",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.21399/c2 names \"two\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-374",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.21399/c6 names \"two\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-375",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.21399/c10 names \"0.423\", \"0.394\", \"0.436\", \"0.417\", \"four\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-376",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.21399/c15 names \"ScienceWorld\", \"23\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-377",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.21399/c17 names \"Prompt-only\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-378",
   "reason": "detail_not_in_cited_evidence",
   "question": "Claim 2606.21399/c20 names \"On ALFWorld\", which is real paper vocabulary but appears in none of the passages cited for this claim, nor in the window shown to the checker. Is this misattribution or partial citation?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-379",
   "reason": "critic_high_severity",
   "question": "1 high-severity objection(s) stand against this extraction. Do any invalidate a claim that should not remain in the graph?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-380",
   "reason": "grounded_recovery_low",
   "question": "Only 0.2692 of source-grounded claims were recovered by the second reader pass. Is this coverage diversity, or extractor unreliability? The distinction determines whether the 0.50 raw overlap is bad news.",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  },
  {
   "id": "esc-381",
   "reason": "extraction_unstable",
   "question": "The two reader passes agreed on only 0.2692 of claims. Should extractions below this threshold be treated as unusable?",
   "status": "open",
   "raised_at": null,
   "resolved_at": null,
   "claim": null
  }
 ]
}