35 lines
No EOL
2 KiB
JSON
35 lines
No EOL
2 KiB
JSON
{
|
|
"rejected_claims": [
|
|
{
|
|
"filename": "benchmark-based-capability-metrics-overstate-real-world-autonomous-performance-because-automated-scoring-ignores-production-readiness-requirements.md",
|
|
"issues": [
|
|
"missing_attribution_extractor"
|
|
]
|
|
},
|
|
{
|
|
"filename": "ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md",
|
|
"issues": [
|
|
"missing_attribution_extractor"
|
|
]
|
|
}
|
|
],
|
|
"validation_stats": {
|
|
"total": 2,
|
|
"kept": 0,
|
|
"fixed": 5,
|
|
"rejected": 2,
|
|
"fixes_applied": [
|
|
"benchmark-based-capability-metrics-overstate-real-world-autonomous-performance-because-automated-scoring-ignores-production-readiness-requirements.md:set_created:2026-03-27",
|
|
"benchmark-based-capability-metrics-overstate-real-world-autonomous-performance-because-automated-scoring-ignores-production-readiness-requirements.md:stripped_wiki_link:verification degrades faster than capability grows",
|
|
"ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md:set_created:2026-03-27",
|
|
"ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md:stripped_wiki_link:adoption lag exceeds capability limits as primary bottleneck",
|
|
"ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md:stripped_wiki_link:the gap between theoretical AI capability and observed deplo"
|
|
],
|
|
"rejections": [
|
|
"benchmark-based-capability-metrics-overstate-real-world-autonomous-performance-because-automated-scoring-ignores-production-readiness-requirements.md:missing_attribution_extractor",
|
|
"ai-tools-reduced-experienced-developer-productivity-19-percent-in-rct-despite-developer-expectations-of-speedup.md:missing_attribution_extractor"
|
|
]
|
|
},
|
|
"model": "anthropic/claude-sonnet-4.5",
|
|
"date": "2026-03-27"
|
|
} |