{ "rejected_claims": [ { "filename": "benchmark-based-capability-metrics-overstate-real-world-autonomous-performance-because-automated-scoring-ignores-production-readiness-requirements.md", "issues": [ "missing_attribution_extractor" ] }, { "filename": "ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md", "issues": [ "missing_attribution_extractor" ] } ], "validation_stats": { "total": 2, "kept": 0, "fixed": 5, "rejected": 2, "fixes_applied": [ "benchmark-based-capability-metrics-overstate-real-world-autonomous-performance-because-automated-scoring-ignores-production-readiness-requirements.md:set_created:2026-03-27", "benchmark-based-capability-metrics-overstate-real-world-autonomous-performance-because-automated-scoring-ignores-production-readiness-requirements.md:stripped_wiki_link:verification degrades faster than capability grows", "ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md:set_created:2026-03-27", "ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md:stripped_wiki_link:adoption lag exceeds capability limits as primary bottleneck", "ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md:stripped_wiki_link:the gap between theoretical AI capability and observed deplo" ], "rejections": [ "benchmark-based-capability-metrics-overstate-real-world-autonomous-performance-because-automated-scoring-ignores-production-readiness-requirements.md:missing_attribution_extractor", "ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md:missing_attribution_extractor" ] }, "model": "anthropic/claude-sonnet-4.5", "date": "2026-03-27" }