{
  "id": 9439046,
  "title": "Mercury 2.5 LLM hits 770 tokens per second",
  "url": "https://urgent.news/2026/09/23/mercury-2-5-llm-hits-770-tokens-per-second",
  "topic": "ai",
  "section": "AI",
  "published": "2026-09-23T22:16:19.000Z",
  "source": {
    "name": "Hacker News",
    "slug": "hacker-news",
    "url": "https://artificialanalysis.ai/models/mercury-2-5"
  },
  "original_language": "en",
  "account": "Mercury 2.5, a model priced moderately among its peers, has demonstrated impressive performance with a token output rate of 770 tokens per second. This model operates within a 260k token context window and has earned a score of 12 on the Artificial Analysis Intelligence Index, placing it slightly below the median of 13. Despite its average intelligence, Mercury 2.5 stands out for its speed and conciseness. It is capable of generating 35 million tokens, which is comparatively concise to the median of 85 million. The model's pricing, both for input and output tokens, aligns with the median values, making it cost-effective for users. With an Elo score of 500 and a weighted average cost per task, Mercury 2.5 proves to be a valuable tool for various applications, particularly in Agentic real-world work tasks.",
  "summary": null,
  "key_points": [],
  "editors_take": null,
  "illustration": null,
  "coverage": {
    "outlets": 1,
    "also_reported_by": []
  },
  "ai_generated": true,
  "disclaimer": "Summaries, key points and the editor’s take are written by software from other outlets’ reporting and may contain errors — always check the linked original."
}