{
  "id": 6651758,
  "title": "Reduce inference cold starts on Amazon SageMaker HyperPod with model caching",
  "url": "https://urgent.news/2026/09/10/reduce-inference-cold-starts-on-amazon-sagemaker-hyperpod-with-model",
  "topic": "ai",
  "section": "AI",
  "published": "2026-09-10T21:37:49.000Z",
  "source": {
    "name": "AWS Machine Learning",
    "slug": "aws-machine-learning",
    "url": "https://aws.amazon.com/blogs/machine-learning/reduce-inference-cold-starts-on-amazon-sagemaker-hyperpod-with-model-caching/"
  },
  "original_language": "en",
  "account": null,
  "summary": "Amazon SageMaker HyperPod now supports model caching for inference, which pre-loads model weights and container images onto cluster nodes so pods read from local NVMe storage instead of downloading over the network. Learn how model caching cuts cold starts from tens of minutes to seconds, how it works, and how to enable it.",
  "key_points": [],
  "editors_take": null,
  "illustration": null,
  "coverage": {
    "outlets": 2,
    "also_reported_by": [
      {
        "outlet": "AWS Machine Learning",
        "title": "Amazon SageMaker Feature Store introduces UpdateRecord for feature-level writes",
        "url": "https://urgent.news/2026/09/08/amazon-sagemaker-feature-store-introduces-updaterecord-for-feature",
        "published": "2026-09-08T18:29:15.000Z"
      }
    ]
  },
  "ai_generated": true,
  "disclaimer": "Summaries, key points and the editor’s take are written by software from other outlets’ reporting and may contain errors — always check the linked original."
}