{"as_of":"2026-08-17T20:23:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f295bb6d93d9efc665f1296ba1206054a0964d8f621cf402151f3f816ce911bf","coverage":[{"denominator":51,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":51,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T15:58:48.247129Z","state":"measured"},{"denominator":52,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":52,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T05:49:10.326595Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-16T05:49:10.393794Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"cited_work":{"arxiv_id":"2411.13820","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.13820","snapshot_observed_at":"2026-08-16T05:49:10.393794Z","title":"InstCache: A Predictive Cache for LLM Serving","venue":"cs.CL","work_id":"8b231dca-e68a-4bc0-966a-e103c807ecd4","year":2024},"citing_paper":{"arxiv_id":"2504.19720","last_updated":"2025-04-28T12:14:02Z","snapshot_observed_at":"2026-08-17T14:38:27.066872Z","submitted_at":"2025-04-28T12:14:02Z","title":"Taming the Titans: A Survey of Efficient LLM Inference Serving","version":1},"reference_index":146,"source":"arxiv_source","source_observed_at":"2026-08-16T05:49:10.326595Z"},"links":{"cited_paper":"/paper/2411.13820","citing_paper":"/paper/2504.19720"},"observation_digest":"sha256:3f400fbbf6035845ba7e726f7e32c430d0144a179eb191bcaebfe172f3ece1b3","observation_id":"b4ff8910-02e5-416e-92cc-d44a248b28cf","resolution":{"observed_at":"2026-08-16T05:49:10.400777Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2411.13820/citation-record","integrity":"/paper/2411.13820/integrity","json":"/paper/2411.13820/citation-record.json","paper":"/paper/2411.13820"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2112.00861","last_updated":"2021-12-09T21:40:22Z","snapshot_observed_at":"2026-08-10T12:03:10.373653Z","submitted_at":"2021-12-01T22:24:34Z","title":"A General Language Assistant as a Laboratory for Alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.00861","snapshot_observed_at":"2026-08-12T15:58:48.050699Z","title":"Brown, Jack Clark, Sam McCandlish, Chris Olah, and Jared Kaplan","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.050699Z"},"links":{"cited_paper":"/paper/2112.00861","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:fa040c11a64306c81db9b934383a302619a37d2d97636e10f082c0fbfb0390ee","observation_id":"2abab399-c656-4168-8eb6-e8a656bf5c3b","resolution":{"observed_at":"2026-08-12T15:58:48.050699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.871744Z","title":"Baeza-Yates and Felipe Saint-Jean","venue":null,"work_id":"04f26d6f-eb19-4676-8bfc-54e6fe9cdeac","year":2003},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.055211Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:803da72766fe19d77b806520a212aa0de01544be12f90d52af8080262a68c1c7","observation_id":"344203b4-e725-4fb2-90e0-810cc47893bc","resolution":{"observed_at":"2026-08-12T15:58:48.876060Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.855128Z","title":"GPTCache: An open-source semantic cache for LLM applications enabling faster answers and cost savings","venue":null,"work_id":"57d79c06-5c3a-4c24-9f1c-3fa4d7a00a2e","year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.059033Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:187e6c7888e2a977515fa95d2897f2dbbdd5db93ceeb1cef55de94e2880f5096","observation_id":"22190558-f55d-49da-9dde-c0889ff741ab","resolution":{"observed_at":"2026-08-12T15:58:48.861489Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.833310Z","title":"Jordan, Joseph E","venue":null,"work_id":"50a783dd-ab77-44cc-85f3-c734a9a84e10","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.067059Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:5a7f7e661ab77ff135d841844f0242d76c523d33f1558bd6bd91a37e9894f5cc","observation_id":"27d6ce61-9612-489f-a787-f7c86fa94047","resolution":{"observed_at":"2026-08-12T15:58:48.838176Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04434","last_updated":"2024-06-19T06:04:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-05-07T15:56:43Z","title":"DeepSeek-V2: A Strong, Economical, and Efficient Mixture-of-Experts Language Model","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04434","snapshot_observed_at":"2026-08-12T15:58:48.071378Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.071378Z"},"links":{"cited_paper":"/paper/2405.04434","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:300d2451e146cec9b288b7465efe386a7ccec1fe1ac01680d81a8c09e6e5dde9","observation_id":"c1e4f6b5-10d6-4f12-8071-c0bb3b1d3877","resolution":{"observed_at":"2026-08-12T15:58:48.071378Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.19437","last_updated":"2025-02-18T17:26:38Z","snapshot_observed_at":"2026-08-15T17:27:11.980940Z","submitted_at":"2024-12-27T04:03:16Z","title":"DeepSeek-V3 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.19437","snapshot_observed_at":"2026-08-12T15:58:48.075969Z","title":"Zhang, Han Bao, Hanwei Xu, Haocheng Wang, Haowei Zhang, Honghui Ding, Huajian Xin, Huazuo Gao, Hui Li, Hui Qu, J","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.075969Z"},"links":{"cited_paper":"/paper/2412.19437","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:bc7a84013cc457c9f91b110fd408fc0262fe020d1b7a0406de615d2191a742a1","observation_id":"5328728b-8302-47b7-8681-5b26f7cb8a89","resolution":{"observed_at":"2026-08-12T15:58:48.075969Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-12T15:58:48.080782Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.080782Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:4f9c9055a850a408adc97d13655b473e53f8c9ff22b1f121014d17c790c189a8","observation_id":"6d5c70ba-d067-41ca-92d7-61f3448ccb4b","resolution":{"observed_at":"2026-08-12T15:58:48.080782Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.822038Z","title":"Boosting the perfor- mance of web search engines: Caching and prefetching query results by exploiting historical usage data","venue":null,"work_id":"d36a7c50-f895-4003-88a0-967fb87df439","year":2006},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.084979Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:165f59e0c6a9b1be25f41948c434a163519f482a220fa50c1559c7667bbd4f64","observation_id":"b1ff2bd5-0485-4dbd-9abb-efdf031817d6","resolution":{"observed_at":"2026-08-12T15:58:48.825813Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.10997","last_updated":"2024-03-27T09:16:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-18T07:47:33Z","title":"Retrieval-Augmented Generation for Large Language Models: A Survey","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.10997","snapshot_observed_at":"2026-08-12T15:58:48.092990Z","title":"Retrieval-augmented generation for large language models: A survey","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.092990Z"},"links":{"cited_paper":"/paper/2312.10997","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:b0b16e4e37d06a86bf0183ecf20ab8c22109af0bd892ccdfb359593ec8661f06","observation_id":"5ea15936-7068-4813-a226-d02a1d076ed5","resolution":{"observed_at":"2026-08-12T15:58:48.092990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.19708","last_updated":"2024-06-30T23:50:38Z","snapshot_observed_at":"2026-08-16T14:07:07.762606Z","submitted_at":"2024-03-23T10:42:49Z","title":"Cost-Efficient Large Language Model Serving for Multi-turn Conversations with CachedAttention","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.19708","snapshot_observed_at":"2026-08-12T15:58:48.097190Z","title":"Privacy-aware semantic cache for large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.097190Z"},"links":{"cited_paper":"/paper/2403.19708","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:ddc25d0d3008e54489c58c7eba0292b3f79f30a95af3f5ffdf3ce68d616c2906","observation_id":"4cbadb51-9cac-4a66-a3be-15320991c338","resolution":{"observed_at":"2026-08-12T15:58:48.097190Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.101011Z","title":"Efficient memory management for large language model serving with pagedattention","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.101011Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:a8949009b883b0cd9aa03b54fecfb1b479287443e5decacf5da6a14ff627e336","observation_id":"a0f343f8-9ec2-40bc-9b43-b15ea0c5aeb5","resolution":{"observed_at":"2026-08-12T15:58:48.101011Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.802857Z","title":"Predictive caching and prefetching of query results in search engines","venue":null,"work_id":"01ed570a-397d-4989-849e-f4dce2b4ce55","year":2003},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.104821Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:9bf2e976671f7a563260d968ae4f3996925556c0da9da6d69eb9e10ec583bc45","observation_id":"545658c7-f2e4-4a1f-9c57-469b91ab1e67","resolution":{"observed_at":"2026-08-12T15:58:48.807271Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.790891Z","title":"von Riedemann, Cong Zhang, and Jiangchuan Liu","venue":null,"work_id":"f1777ceb-71d6-43c1-aca9-441611b4dead","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.108549Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:e1b37b5f8f9f2a0df79b1190b504264e07fe0bd07d8f37e95926026ab286d67c","observation_id":"e45bdd50-287f-41dd-97dd-ee74c5ef7a9e","resolution":{"observed_at":"2026-08-12T15:58:48.794791Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.778957Z","title":"Three-level caching for efficient query processing in large web search engines","venue":null,"work_id":"ac19456f-3840-4415-bf16-69b6ee514673","year":2005},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.112557Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:5fd14b3f4663cc1d6992f117312221fab1af0dccdc5c3d86d43fda9f18c62b1b","observation_id":"a52ced6d-fa0a-4d03-a700-4390a44dc713","resolution":{"observed_at":"2026-08-12T15:58:48.782959Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.767230Z","title":"New caching techniques for web search engines","venue":null,"work_id":"9ced39de-3789-4aeb-91c6-316aee3a934f","year":2010},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.116208Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:46b3d9278a33fbdf05458f0b40267aed1a331706494aee6b518944a2772f4906","observation_id":"5bfe0ea8-678b-4b98-bbea-1935b14ce1a2","resolution":{"observed_at":"2026-08-12T15:58:48.770863Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.756960Z","title":"Markatos","venue":null,"work_id":"b0ce5fb4-f531-4c17-a900-e24f85be6e28","year":2001},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.119909Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:197188ee5eea1e15c32b5dc9a352c51b6c40403d14d3d3c919c25fddef87a7c2","observation_id":"928c0530-b343-4312-887c-68dbea1d6cdf","resolution":{"observed_at":"2026-08-12T15:58:48.760554Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.746338Z","title":"Context-based semantic caching for llm applications","venue":null,"work_id":"f14f00bc-4f28-4c81-be99-a967e4afe93b","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.123660Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:50766a71c8195f8d2272f9a7d17869afa36dfe91ba80d7a1716c177f3f4224f2","observation_id":"93294676-7c33-42d1-a25a-d542b78a2ac4","resolution":{"observed_at":"2026-08-12T15:58:48.750157Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.735318Z","title":"Openai chatgpt, 2022","venue":null,"work_id":"b0cf169f-2616-47b8-a148-fdd22fd9d1bd","year":2022},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.127561Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:29312529d5a5900e3fdf65d6231808102a5a0478dcd57f6a6268558a7ea397a0","observation_id":"67f5381b-d07a-4041-903e-aac89dae261a","resolution":{"observed_at":"2026-08-12T15:58:48.739502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08138","last_updated":"2024-01-16T06:16:33Z","snapshot_observed_at":"2026-08-16T14:27:14.324570Z","submitted_at":"2024-01-16T06:16:33Z","title":"LLMs for Test Input Generation for Semantic Caches","version":1},"cited_work":{"arxiv_id":"2401.08138","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.08138","snapshot_observed_at":"2026-08-12T15:58:48.354508Z","title":"LLMs for Test Input Generation for Semantic Caches","venue":"cs.SE","work_id":"076de712-c5f2-47e3-bffd-b3df0d95efb2","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.131050Z"},"links":{"cited_paper":"/paper/2401.08138","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:1aefbbecc13dc3cf7b694db6b3fc2222e75c881f740f801c9cf6a6dcd15fccb7","observation_id":"1d4b706a-8d53-46d4-9537-d42729467a35","resolution":{"observed_at":"2026-08-12T15:58:48.361102Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.725732Z","title":"Sentence-bert: Sentence embeddings using siamese bert- networks","venue":null,"work_id":"311864a1-8542-447f-a60e-8710b04696a3","year":2019},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.135036Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:4f40296451e19c4e231f0479752d7b3a318ab12ea359719930bd8b5bcdf9bf31","observation_id":"55c39471-6c74-4266-b09f-e877b2d90aba","resolution":{"observed_at":"2026-08-12T15:58:48.728852Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.716062Z","title":"Fonseca, Wagner Meira Jr., Berthier A","venue":null,"work_id":"3dae246b-8ffa-4a29-989d-68c60f1dfec0","year":2001},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.138752Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:0649ddcc318cfb6a38d167fe65f597dbf9a3cfb12534804a6376b11ab11baed5","observation_id":"fd9dafe8-7aa6-4e33-bbdf-d6f2d4f9579c","resolution":{"observed_at":"2026-08-12T15:58:48.719812Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.142480Z","title":"The early bird catches the leak: Unveiling timing side channels in LLM serving systems","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.142480Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:24e222cb2cb8d14b2a639e59517594eda51a973ac9d8d09149dc513accdbd3b7","observation_id":"70151361-b4af-41ad-8e63-99047fe0a28c","resolution":{"observed_at":"2026-08-12T15:58:48.142480Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.704727Z","title":"MOSS: an open conversational large language model","venue":null,"work_id":"42c72140-db31-43e2-868a-be5ab28564ba","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.145517Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:e03cd262ac04b986770d8817d1f8f8b1ded5a5299febd75918475248f0d3410a","observation_id":"5daaef7f-06e2-4b93-be1b-cdc8956929af","resolution":{"observed_at":"2026-08-12T15:58:48.708827Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.693135Z","title":"Sharegpt, 2023","venue":null,"work_id":"72f6a429-0f2d-4bf4-8eac-b711a50c4118","year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.148683Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:d6ccd293ca673df275be0d150808a01a81eafef676affcd975eb564cea0fa653","observation_id":"46740945-70fb-4e0a-b719-3a7ca5d55525","resolution":{"observed_at":"2026-08-12T15:58:48.696983Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.681317Z","title":"Efficient streaming language models with attention sinks","venue":null,"work_id":"f5439a9e-b553-45eb-a57c-2adf2e3babff","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.151726Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:fdf7e66a8fbee1d62f6fcac71b71ef9357ba2335f77c5b70813e140529edb89c","observation_id":"acafe0b1-14a7-4724-bcf4-0f7c94be0a1f","resolution":{"observed_at":"2026-08-12T15:58:48.684968Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.670205Z","title":"O’Hallaron","venue":null,"work_id":"5bb19433-eafd-4cca-85d4-2c0b098eb596","year":2002},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.155258Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:f0d99e5fefd713e9a589bd7d02292124b22d0d5bac019b8e2b4cb63601756e41","observation_id":"46c2cb6a-1aaf-4612-a74c-4819f842e9f7","resolution":{"observed_at":"2026-08-12T15:58:48.673960Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-08-17T18:50:07.059564Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-12T15:58:48.158767Z","title":"Qwen2.5 technical report","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.158767Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:3f4face30c7ec3f13f0ebf8601d875187372f98df136d7c2f72c37f38f278d9d","observation_id":"db3fb3c9-0a18-4ddc-82b5-c3847940bc78","resolution":{"observed_at":"2026-08-12T15:58:48.158767Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.659452Z","title":"Chunkattention: Efficient self-attention with prefix-aware KV cache and two-phase partition","venue":null,"work_id":"7379c5a2-6d9c-42e9-8f0c-1f10ba33cecf","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.162353Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:71bb0f7ad0e8e8d1ebf1a85980efba7992da131e05cd10877bcd0002fe1f80e0","observation_id":"9134765f-e8b9-4578-b3ff-df789ddf3010","resolution":{"observed_at":"2026-08-12T15:58:48.663350Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.648325Z","title":"Performance of compressed inverted list caching in search engines","venue":null,"work_id":"b0399d78-3918-4d08-bbd8-0971090253b1","year":2008},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.165516Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:5d2176c084cac27fa00f07b9de8f7aef96c540557371fb62eb780a25c2f17745","observation_id":"8c58630c-38c1-4d0f-a4c1-00c5b9f5bc6b","resolution":{"observed_at":"2026-08-12T15:58:48.651757Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.168601Z","title":"Barrett, Zhangyang Wang, and Beidi Chen","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.168601Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:cb38ec5cee1b7feba7e000052c35457bb8bd29bf8ecf8a0d8f86ff15b3f8bc54","observation_id":"eb972039-3c82-483d-9b62-d63b8f6d2df9","resolution":{"observed_at":"2026-08-12T15:58:48.168601Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.631254Z","title":"Wildchat: 1m chatgpt interaction logs in the wild","venue":null,"work_id":"57ca10cc-f6d5-4cd3-a43c-2d1c56c433d5","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.172494Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:0e6038d88f76a2de10354e95e7bc08c68af6b72f35ef776d5f879276e1c3ce21","observation_id":"8d2ec94d-a311-48f4-87d4-cb468bb46d7e","resolution":{"observed_at":"2026-08-12T15:58:48.635067Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.620600Z","title":"Xing, Joseph E","venue":null,"work_id":"905dc580-7b3f-43da-9f82-bef0df03ccfd","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.176135Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:b91cb8e4ab13bbdc53ce3090eaca441e5d2f0ce95e6a13b2d9b932d0e1bad163","observation_id":"9f6f15a9-680b-43cb-ae4d-97bf965e4e33","resolution":{"observed_at":"2026-08-12T15:58:48.624149Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.07104","last_updated":"2024-06-06T00:10:06Z","snapshot_observed_at":"2026-08-16T21:51:10.462948Z","submitted_at":"2023-12-12T09:34:27Z","title":"SGLang: Efficient Execution of Structured Language Model Programs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.07104","snapshot_observed_at":"2026-08-12T15:58:48.183725Z","title":"kinky date","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.183725Z"},"links":{"cited_paper":"/paper/2312.07104","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:08f8d13d3417274ec1b2fe348b4e354f9c22e38110eab28e17350f33471a9114","observation_id":"eb374e5a-b77a-4d85-84b3-6bd15b4e574a","resolution":{"observed_at":"2026-08-12T15:58:48.183725Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.188245Z","title":"Guidelines: • The answer NA means that the abstract and introduction do not include the claims made in the paper","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.188245Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:66efda3ab20531d80442de57d18261db1f08f340614289f055104f18e0809913","observation_id":"5931fcfd-dfff-4c8f-be91-7307a1186dda","resolution":{"observed_at":"2026-08-12T15:58:48.188245Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.598565Z","title":"Limitations","venue":null,"work_id":"a3b4c566-4a34-4475-a661-79478e5138e1","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.192197Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:3365431a7ab677f55aaab285878666806a305dfcec1050bae792e9d3c1abd4e5","observation_id":"ebdb5508-c5ca-4c93-bfd8-f71db2fd09ee","resolution":{"observed_at":"2026-08-12T15:58:48.602113Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.588273Z","title":"Guidelines: • The answer NA means that the paper does not include theoretical results","venue":null,"work_id":"efef348b-f9b4-43cc-bd96-23b68f0bc0d8","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.196316Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:0490696e596deaa884e2b82a92caed5336fc7ad82f02ef3c0b328736f1075193","observation_id":"4cc2ff57-2cde-4661-969f-49d7ea375e7a","resolution":{"observed_at":"2026-08-12T15:58:48.592001Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.575588Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"c5408018-7e2f-4e24-b623-fca79ec0c0c2","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.200359Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:547fd7937173e5836a21a38f733ecb2ad3a212dddaa1c1b51459b492117a20ac","observation_id":"b8396268-01a7-4634-a7ea-ec54dfac6152","resolution":{"observed_at":"2026-08-12T15:58:48.579256Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.564948Z","title":"Guidelines: • The answer NA means that paper does not include experiments requiring code","venue":null,"work_id":"c709714d-fb41-43c7-bba3-63126bc539dc","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.204341Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:9525625733f9260be0b797e6eb8cb36710a6805bcd7aaae38e899675d66653a8","observation_id":"2a19b655-f77a-44ff-a18f-f21f7894b4eb","resolution":{"observed_at":"2026-08-12T15:58:48.568744Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.554452Z","title":"• The experimental setting should be presented in the core of the paper to a level of detail that is necessary to appreciate the results and make sense of them","venue":null,"work_id":"07215d91-0336-4f4e-afe1-91ce149d49ae","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.208712Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:e149d1f2e2351f63fc262ea80fa7ae86026e3e3e114680d7b8248a23ae1ba228","observation_id":"83d019a6-5743-4c91-bbd7-d997dc31b925","resolution":{"observed_at":"2026-08-12T15:58:48.558482Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.544288Z","title":"This experimental setup is also consistent with existing studies[12, 34]","venue":null,"work_id":"97962488-9555-4804-bf07-ee947ccd27e6","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.212336Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:687ed8626c20bb5c76cadc2e1b8d14a37f0ec2c27565c14d7ae6bf4e9046c63b","observation_id":"be2ff848-1b0f-4946-8ae5-4490b02099a2","resolution":{"observed_at":"2026-08-12T15:58:48.547812Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.534267Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"83c5480f-7de0-4658-99d1-da611f255214","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.216428Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:dbdf798b2589c75e5f194879c663e320d07c6df96d310f1a9a9a0ff0dde51f91","observation_id":"800ef9d8-b73f-4f5b-acc0-669ff943bbb2","resolution":{"observed_at":"2026-08-12T15:58:48.537471Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.524317Z","title":"Guidelines: • The answer NA means that the authors have not reviewed the NeurIPS Code of Ethics","venue":null,"work_id":"1fc85633-97ef-431f-a3af-aae69912202b","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.220373Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:e917bf621e9b0f1cc88271ed36f341029cce1944bceab911b38ab1d4973c5b56","observation_id":"1d185f71-5ee1-4cbc-9f7a-2a0b69d7905f","resolution":{"observed_at":"2026-08-12T15:58:48.527890Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.512537Z","title":"Guidelines: • The answer NA means that there is no societal impact of the work performed","venue":null,"work_id":"6c12dcb9-aef9-4882-9edb-3662e0f56125","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.223953Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:525457f630db8bfc3bbba9398de74b1941cc9d0aea52409237c17a137063994c","observation_id":"8451f0d9-f0b5-4737-af24-dca0c05e1a75","resolution":{"observed_at":"2026-08-12T15:58:48.516492Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.500580Z","title":"Guidelines: • The answer NA means that the paper poses no such risks","venue":null,"work_id":"7e6176c6-d58a-4266-bbd5-871793528aac","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.227844Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:a024351a0dd503522375065af36a1411956de2c78a9be34cf4bf5f94e71c653e","observation_id":"9061bc30-2137-4789-b29a-cf01f4c521ac","resolution":{"observed_at":"2026-08-12T15:58:48.504972Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.488042Z","title":"Guidelines: • The answer NA means that the paper does not use existing assets","venue":null,"work_id":"f121fbc3-9650-43d6-8c73-87ae5e7a7e7f","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.231431Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:d5cf6eb4f54a4535ab60ae9bb143cf99c9c32dae7ae74593f64862408425cbae","observation_id":"03375e8b-a654-4846-b014-a03e5f2d0c8f","resolution":{"observed_at":"2026-08-12T15:58:48.492294Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.235405Z","title":"Guidelines: • The answer NA means that the paper does not release new assets","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.235405Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:44727e243c1d553ec54c1832e38f83d20f6e56325814700945718ae0e2d0c7cd","observation_id":"c7a31830-0eae-4c82-986e-d214d17cedc7","resolution":{"observed_at":"2026-08-12T15:58:48.235405Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.239285Z","title":"Guidelines: • The answer NA means that the paper does not involve crowdsourcing nor research with human subjects","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.239285Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:9d366ce3685f847858d02b0e36078fd382a6e4969418367a532cebe3d60e1b5e","observation_id":"d157f936-cc07-43b8-aca2-8388b67ab52e","resolution":{"observed_at":"2026-08-12T15:58:48.239285Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.457648Z","title":"Guidelines: • The answer NA means that the paper does not involve crowdsourcing nor research with human subjects","venue":null,"work_id":"d59d90a2-3efe-4378-b781-507d725370f2","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.243029Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:4f4cc5db4fe052c83be52e09c7f00f7a076cea7d1bf8f6be6a1e6663b93ac360","observation_id":"446d0345-9ff9-4eb6-9f91-209a7b8a567c","resolution":{"observed_at":"2026-08-12T15:58:48.462095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.444465Z","title":"Answer: [NA] Justification: The core method development in this research does not involve LLMs as any important, original, or non-standard components","venue":null,"work_id":"ec04fe49-8440-4de1-8062-9931d9fcb86f","year":2025},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.247129Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:a40e7a0723cc5f362f37c2f171e17c966a1b7989d152efef73b667c7f427dfea","observation_id":"5804080a-9472-46f2-a961-36321cc23137","resolution":{"observed_at":"2026-08-12T15:58:48.448600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.063168Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.063168Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:042b11445cc71de837dac9bdf9f2918ccec0c876451781ec1b0c7eb88030820d","observation_id":"f3280a63-f62c-4be1-a0d4-6a44e68241ed","resolution":{"observed_at":"2026-08-12T15:58:48.063168Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.179849Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.179849Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:97c2334e2a7f63d9e528b89ebef3d4439edb205ecdd5520f1665e1842aee10a1","observation_id":"873708e1-60b9-40f8-83f5-05df7e040935","resolution":{"observed_at":"2026-08-12T15:58:48.179849Z","resolver_source":null,"status":"parse_uncertain"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-16T03:10:21.519704Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving"},"reference_resolution":{"displayed":51,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":2,"unresolved":14,"verified_exact":1,"verified_fuzzy":33},"total_outbound_references":51},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 17 August 2026, this Paper Citation Record lists 51 of 51 outbound references and 1 inbound Pith citation observation for arXiv:2411.13820."}