{
  "id": 250147,
  "title": "The Low Frequency Trap: Video Language Models Fail at Simple Event Bookkeeping",
  "url": "https://urgent.news/2026/08/06/the-low-frequency-trap-video-language-models-fail-at-simple-event",
  "topic": "culture",
  "section": "Culture",
  "published": "2026-08-06T17:57:06.000Z",
  "source": {
    "name": "arXiv cs.AI",
    "slug": "arxiv-cs-ai",
    "url": "https://arxiv.org/abs/2608.06361v1"
  },
  "original_language": "en",
  "account": null,
  "summary": "Real-world video benchmarks provide broad coverage, but their fixed clips entangle event count, rate, duration, and visual complexity, making failure modes hard to isolate. While existing programmatic benchmarks offer better control, they score only the final answer rather than auditing reported events against executable ground truth. To bridge this gap, we introduce trace-grounded parametric…",
  "key_points": [],
  "editors_take": null,
  "illustration": null,
  "coverage": {
    "outlets": 1,
    "also_reported_by": []
  },
  "ai_generated": true,
  "disclaimer": "Summaries, key points and the editor’s take are written by software from other outlets’ reporting and may contain errors — always check the linked original."
}