teleo-codex/inbox/queue/.extraction-debug/2026-01-29-metr-time-horizon-1-1.json
Teleo Agents 98d283e794
Some checks are pending
Sync Graph Data to teleo-app / sync (push) Waiting to run
extract: 2026-01-29-metr-time-horizon-1-1
Pentagon-Agent: Epimetheus <3D35839A-7722-4740-B93D-51157F7D5E70>
2026-03-24 00:17:26 +00:00

33 lines
No EOL
1.3 KiB
JSON

{
"rejected_claims": [
{
"filename": "metr-time-horizon-benchmark-saturating-at-governance-relevant-capability-levels.md",
"issues": [
"missing_attribution_extractor"
]
},
{
"filename": "ai-capability-evaluation-scaffold-sensitivity-introduces-cross-model-comparison-uncertainty.md",
"issues": [
"missing_attribution_extractor"
]
}
],
"validation_stats": {
"total": 2,
"kept": 0,
"fixed": 3,
"rejected": 2,
"fixes_applied": [
"metr-time-horizon-benchmark-saturating-at-governance-relevant-capability-levels.md:set_created:2026-03-24",
"metr-time-horizon-benchmark-saturating-at-governance-relevant-capability-levels.md:stripped_wiki_link:verification degrades faster than capability grows",
"ai-capability-evaluation-scaffold-sensitivity-introduces-cross-model-comparison-uncertainty.md:set_created:2026-03-24"
],
"rejections": [
"metr-time-horizon-benchmark-saturating-at-governance-relevant-capability-levels.md:missing_attribution_extractor",
"ai-capability-evaluation-scaffold-sensitivity-introduces-cross-model-comparison-uncertainty.md:missing_attribution_extractor"
]
},
"model": "anthropic/claude-sonnet-4.5",
"date": "2026-03-24"
}