{"as_of":"2026-08-20T18:01:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:9cbbaefe8b10718da2aeb7bd366011fa086052d145adac18de3b0d676289814c","coverage":[{"denominator":21,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":21,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-01T13:13:53.265220Z","state":"measured"},{"denominator":21,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":21,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.19214/citation-record","integrity":"/paper/2607.19214/integrity","json":"/paper/2607.19214/citation-record.json","paper":"/paper/2607.19214"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:51.870958Z","title":"Aider change history (v0.53.0: cache keepalive)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:51.870958Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:6f1d64b4346eed4b0b4419769ecf1868bd59b0c6dd48e376e557f83ad7402fad","observation_id":"69dcede4-4127-44f2-9027-e73ea494f222","resolution":{"observed_at":"2026-08-01T13:13:51.870958Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:51.914286Z","title":"Prompt caching","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:51.914286Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:7d725a4c75bb35139c0e3ffdf77fdc64c1f0f327f0b555d3545bef28f938d540","observation_id":"6902f075-26b6-4a3c-a8c0-d850bd0759e8","resolution":{"observed_at":"2026-08-01T13:13:51.914286Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:51.963972Z","title":"Prompt caching: keeping the cache warm","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:51.963972Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:68b1da423c85699c343979e61c78af11950f26da34f86cc6751f057dcf680498","observation_id":"f075fe2f-b888-4055-bcaf-731fd35f7632","resolution":{"observed_at":"2026-08-01T13:13:51.963972Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.044672Z","title":"claude-code-cache-keepalive (plugin, with break-even analysis)","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.044672Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:fa4c69625859e2e1b521b0dbab3b17f3e4010dd4ec17a2f9bfc594ed8d17bce0","observation_id":"0bf08aaa-3c63-4da4-9f34-041aeefd6555","resolution":{"observed_at":"2026-08-01T13:13:52.044672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.30613","last_updated":"2026-05-28T22:06:54Z","snapshot_observed_at":"2026-08-17T01:07:47.305172Z","submitted_at":"2026-05-28T22:06:54Z","title":"CacheProbe: Auditing Prompt Cache Isolation in Gateway APIs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.30613","snapshot_observed_at":"2026-08-01T13:13:52.125516Z","title":"Cacheprobe: Auditing prompt cache isolation in gateway apis.arXiv preprint arXiv:2605.30613, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.125516Z"},"links":{"cited_paper":"/paper/2605.30613","citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:b31aead38d0f952e46aa4995fc1876296ad4121721839ec97ed1ca245ee8fc6b","observation_id":"b05ade23-debc-4ac9-b3f8-83b1daeab2d4","resolution":{"observed_at":"2026-08-01T13:13:52.125516Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.172600Z","title":"Prompt cache: Modular attention reuse for low-latency inference","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.172600Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:fd8d49c2082a93c31b4b4da211279aefd1d0ea3359a28fd3750f931aff64fa99","observation_id":"040e370c-1b2e-4751-bcf0-f7029efc9c74","resolution":{"observed_at":"2026-08-01T13:13:52.172600Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.259359Z","title":"Context caching","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.259359Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:4d1f3959287efd2c5ad7ec95a0eed7a2b494d0a95cb88178347fff61fd2a879b","observation_id":"4d640a2b-c1f0-4e7b-a254-5d9743cb061d","resolution":{"observed_at":"2026-08-01T13:13:52.259359Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.367659Z","title":"Efficient memory management for large language model serving with pagedattention","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.367659Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:4addb19e91c422e7cf820c3ffbfd87d3c341416068fcab93f5bbeb2a6fe55ec6","observation_id":"6c1d2d55-d3a7-46e7-a897-feda188c98bd","resolution":{"observed_at":"2026-08-01T13:13:52.367659Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.437226Z","title":"Don’t break the cache: An evaluation of prompt caching for long-horizon agentic tasks.arXiv preprint arXiv:2601.06007, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.437226Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:55ede48dfdc61f1ded9164d50ec9d5dde7bc0aedaa002658e10df6c317cf4604","observation_id":"8e28510f-e257-424d-84ab-438ec0f07e43","resolution":{"observed_at":"2026-08-01T13:13:52.437226Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.441577Z","title":"Prompt caching","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.441577Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:b33914c856de515800400ba78892c2c832d8044e4fb82a752a9348c1b52d6850","observation_id":"c98029dd-653e-49c7-a602-616d381fea11","resolution":{"observed_at":"2026-08-01T13:13:52.441577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.444982Z","title":"Feature: prompt cache keep-warm pings (is- sue #62475)","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.444982Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:6d2941b7d8e56374ccb2e085e117951e126ad2aa21b795ffade7e324d1619e61","observation_id":"8c35f91b-7f7a-4a2d-9055-bb650a2103ba","resolution":{"observed_at":"2026-08-01T13:13:52.444982Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.495561Z","title":"Openrouter api","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.495561Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:a829d7444be3c9c1389094d79c6fc773e258e98886baf806186377bb9448d2f4","observation_id":"51792600-275a-4ad8-ba5b-ec450899556b","resolution":{"observed_at":"2026-08-01T13:13:52.495561Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.572944Z","title":"Prompt caching best practices","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.572944Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:bbdcf1ede5aa619bb255e3f04d2101c62df59a08fba8cd51da5d33bb56c9f569","observation_id":"b3104ab4-3d6d-4ead-b098-3d334cf6f8c6","resolution":{"observed_at":"2026-08-01T13:13:52.572944Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.644822Z","title":"Prompt cache thrashing: a multi-tenant noisy- neighbor billing scenario","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.644822Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:a171d60f37a0ac56bcd0a3294b19ce0d280b568b0bd46345d765cd3b1ab7e602","observation_id":"6ec96f9c-dcd2-4baa-ae92-4307fb1734e1","resolution":{"observed_at":"2026-08-01T13:13:52.644822Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.712166Z","title":"Efficiently scaling transformer inference.Proceedings of Machine Learning and Systems (MLSys), 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.712166Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:573990740889c8cf70a308861501c61b8787bf8b1d0ab1896cb506976507f5f9","observation_id":"df6d3cce-7520-41e6-b21c-3b1056cc395e","resolution":{"observed_at":"2026-08-01T13:13:52.712166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.786509Z","title":"Cachewise: Understanding workloads and optimizing kvcache man- agement for efficiently serving llm coding agents.arXiv preprint arXiv:2606.16824, 2026","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.786509Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:6bf9cc4335ff352ba13335edd02899186a976e8411f5b0c6c111687a9670b420","observation_id":"80869846-4387-4b03-889d-e2921b921d4d","resolution":{"observed_at":"2026-08-01T13:13:52.786509Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:52.893775Z","title":"We taught our ai agents to take coffee breaks","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:52.893775Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:5a6eadf3b75d222c3be3dd93897dc6808eaf324ca752f602fef53e367037375d","observation_id":"2f67e015-8a19-453a-b3aa-600fa62145c4","resolution":{"observed_at":"2026-08-01T13:13:52.893775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:53.012643Z","title":"Automatic prefix caching (design document)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:53.012643Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:6336dc0be2b9e859de025b7c59f4ff2312946128df5e05ae180b66cf6ffb516c","observation_id":"078c4f29-3100-4b1d-933b-5eeeef610f56","resolution":{"observed_at":"2026-08-01T13:13:53.012643Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:53.091493Z","title":"Kvcache cache in the wild: Charac- terizing and optimizing kvcache reuse at a large cloud provider","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:53.091493Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:8b3023eb2620a233d67afd7d02fdec8d0342b881bb0b4c6aeb3b36a000220b5e","observation_id":"53931263-82d4-4986-8eb5-f437caef57de","resolution":{"observed_at":"2026-08-01T13:13:53.091493Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-01T13:13:53.184672Z","title":"Anthropic prompt cache ttl and cost mechanics","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:53.184672Z"},"links":{"citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:02aa417bddd116e7428caeb201cb6b8eee815222689dc88aa43e97ee3635cca7","observation_id":"e87347c6-f521-431d-bb60-1faeb8ef7fd1","resolution":{"observed_at":"2026-08-01T13:13:53.184672Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.03594","last_updated":"2026-04-22T15:33:51Z","snapshot_observed_at":"2026-08-17T14:15:25.018905Z","submitted_at":"2024-11-29T05:57:37Z","title":"BatchLLM: Optimizing Large Batched LLM Inference with Global Prefix Sharing and Throughput-oriented Token Batching","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.03594","snapshot_observed_at":"2026-08-01T13:13:53.265220Z","title":"Batchllm: Optimizing large batched llm inference with global prefix sharing and throughput- oriented token batching.arXiv preprint arXiv:2412.03594, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-01T13:13:53.265220Z"},"links":{"cited_paper":"/paper/2412.03594","citing_paper":"/paper/2607.19214"},"observation_digest":"sha256:1b43e77ed0ff664578f801e771779955185b6a2fbd51f06b8732d7b78ccb60fb","observation_id":"53923108-590e-41de-ba0c-7d7ae7cd8be7","resolution":{"observed_at":"2026-08-01T13:13:53.265220Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.19214","last_updated":"2026-07-24T00:14:55Z","latest_version":2,"primary_category":"cs.DC","snapshot_observed_at":"2026-08-17T04:17:28.350998Z","submitted_at":"2026-07-21T15:45:47Z","title":"Keeping the Cache Warm Pays: Keepalive Economics for Agentic Workloads"},"reference_resolution":{"displayed":21,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":21,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":21},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 21 of 21 outbound references and 0 inbound Pith citation observations for arXiv:2607.19214."}