{
  "id": 3901,
  "url": "https://arxiv.org/abs/2605.23058v1",
  "title": "A measurement substrate for agentic Kubernetes operations: Methodology and a case study in retrieval-compounding falsification",
  "summary": "Empirical claims about autonomous Kubernetes operations agents are largely unfalsifiable. Published work reports observational results without controlled comparisons against an agent-disabled baseline, selection bias is endemic, pre-registered decision matrices are absent, and samples are typically too small for the noise level of the underlying scoring system. The cause is the same gap that limits the agents themselves: code agents have a verification substrate that turns \"did it work\" into a f",
  "authors": "Joshua Odmark, Gideon Rubin, Deon van der Vyver",
  "category": "research",
  "topics": "bias-fairness,agents-autonomy",
  "orgs": null,
  "regions": null,
  "published_at": "2026-05-21T21:47:52.000Z",
  "fetched_at": "2026-07-14T16:30:36.741Z",
  "source_slug": "arxiv-ethics",
  "source_name": "arXiv",
  "source_homepage": "https://arxiv.org",
  "ethics_ai_record_url": "https://ethics.ai/record/3901",
  "original_url": "https://arxiv.org/abs/2605.23058v1",
  "evidence_status": "source-only",
  "attribution": "via ethics.ai"
}