{"componentChunkName":"component---src-templates-tag-page-js","path":"/tags/prompt-caching/","result":{"data":{"site":{"siteMetadata":{"title":"M.Hassan Ahmed","author":"Hassan11196"}},"allMarkdownRemark":{"totalCount":1,"edges":[{"node":{"excerpt":"Prompt Caching: Cut LLM Cost and Latency Every turn of an agent looks roughly the same from the model’s side. A long system prompt, a list…","fields":{"slug":"/2026-07-27-prompt-caching-cut-llm-cost/"},"frontmatter":{"date":"2026-07-27T00:00:00.000Z","title":"Prompt Caching: Cut LLM Cost and Latency","description":"Prompt caching reuses a request's prefix to cut LLM cost and latency. How the prefix match works, where to put the breakpoint, and the silent cache misses.","tags":["AI","LLM","Prompt Caching","FastAPI","Agents"],"thumbnail":null}}}]}},"pageContext":{"tag":"Prompt Caching"}},"staticQueryHashes":["32046230"]}