{"as_of":"2026-08-20T06:01:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:9e80d9c28f23fe6e0f84203975d2a4787eada3f79f337dc2a9f682ac41539d8b","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":8,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":8,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":8,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":8,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T04:40:04.595609Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-18T21:41:51.696957Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2405.06856","last_updated":"2024-05-11T00:00:23Z","snapshot_observed_at":"2026-08-18T13:43:37.657055Z","submitted_at":"2024-05-11T00:00:23Z","title":"Aladdin: Joint Placement and Scaling for SLO-Aware LLM Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.06856","snapshot_observed_at":"2026-08-12T16:49:22.056194Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13157","last_updated":"2024-11-27T03:25:44Z","snapshot_observed_at":"2026-08-16T10:20:38.532309Z","submitted_at":"2024-11-20T09:46:30Z","title":"Closer Look at Efficient Inference Methods: A Survey of Speculative Decoding","version":2},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-12T16:49:22.056194Z"},"links":{"cited_paper":"/paper/2405.06856","citing_paper":"/paper/2411.13157"},"observation_digest":"sha256:1d08ac149eaba37c370a43c7f794106434d6f0a4d8a5a30391c52c64ebcf3c07","observation_id":"7f645844-ee90-45e6-a8e6-24cf1b32dfad","resolution":{"observed_at":"2026-08-12T16:49:22.056194Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.06856","last_updated":"2024-05-11T00:00:23Z","snapshot_observed_at":"2026-08-18T13:43:37.657055Z","submitted_at":"2024-05-11T00:00:23Z","title":"Aladdin: Joint Placement and Scaling for SLO-Aware LLM Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.06856","snapshot_observed_at":"2026-08-11T14:15:31.836279Z","title":"Aladdin: Joint placement and scaling for slo-aware llm serving,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.19823","last_updated":"2024-12-16T20:01:36Z","snapshot_observed_at":"2026-08-18T13:11:09.105489Z","submitted_at":"2024-12-16T20:01:36Z","title":"A Survey on Large Language Models for Communication, Network, and Service Management: Application Insights, Challenges, and Future Directions","version":1},"reference_index":128,"source":"pdf_text","source_observed_at":"2026-08-11T14:15:31.836279Z"},"links":{"cited_paper":"/paper/2405.06856","citing_paper":"/paper/2412.19823"},"observation_digest":"sha256:7ebc12fafdb43960e6a9095b2c82e5eeae02ae74be5775dd426a312466f1cc4b","observation_id":"4dd6cc08-29be-480a-b124-bbfc75db29ab","resolution":{"observed_at":"2026-08-11T14:15:31.836279Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.06856","last_updated":"2024-05-11T00:00:23Z","snapshot_observed_at":"2026-08-18T13:43:37.657055Z","submitted_at":"2024-05-11T00:00:23Z","title":"Aladdin: Joint Placement and Scaling for SLO-Aware LLM Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.06856","snapshot_observed_at":"2026-08-16T04:40:04.595609Z","title":"Aladdin: Joint placement and scaling for slo-aware llm serving","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.07833","last_updated":"2026-06-04T21:46:57Z","snapshot_observed_at":"2026-08-18T21:04:11.844734Z","submitted_at":"2025-05-01T18:58:26Z","title":"Harmonia: End-to-End RAG Serving Optimization","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-16T04:40:04.595609Z"},"links":{"cited_paper":"/paper/2405.06856","citing_paper":"/paper/2505.07833"},"observation_digest":"sha256:4258504a61bf7c7cede64fb261dd9b58cf4088dede470ad8b7cabd16e6ff9d9c","observation_id":"2f32610f-6597-48bc-a4ae-50505ff32eea","resolution":{"observed_at":"2026-08-16T04:40:04.595609Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.06856","last_updated":"2024-05-11T00:00:23Z","snapshot_observed_at":"2026-08-18T13:43:37.657055Z","submitted_at":"2024-05-11T00:00:23Z","title":"Aladdin: Joint Placement and Scaling for SLO-Aware LLM Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.06856","snapshot_observed_at":"2026-08-15T19:59:46.563582Z","title":"Aladdin: Joint placement and scaling for slo-aware llm serving.arXiv preprint arXiv:2405.06856, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14851","last_updated":"2025-06-17T03:49:25Z","snapshot_observed_at":"2026-08-19T12:55:39.450477Z","submitted_at":"2025-06-17T03:49:25Z","title":"Efficient Serving of LLM Applications with Probabilistic Demand Modeling","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-15T19:59:46.563582Z"},"links":{"cited_paper":"/paper/2405.06856","citing_paper":"/paper/2506.14851"},"observation_digest":"sha256:97cd7030012120c15e9cda4b801df4827ec3247b525f7573cc380ffb4e4271d7","observation_id":"f47cb6dd-eabc-4a98-a480-33842cee5a32","resolution":{"observed_at":"2026-08-15T19:59:46.563582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.06856","last_updated":"2024-05-11T00:00:23Z","snapshot_observed_at":"2026-08-18T13:43:37.657055Z","submitted_at":"2024-05-11T00:00:23Z","title":"Aladdin: Joint Placement and Scaling for SLO-Aware LLM Serving","version":1},"cited_work":{"arxiv_id":"2405.06856","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.06856","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Aladdin: Joint placement and scaling for slo-aware llm serving","venue":null,"work_id":"8f07f1e8-04fe-49b7-b64c-ce1fdd3001a3","year":2024},"citing_paper":{"arxiv_id":"2508.15919","last_updated":"2026-04-23T22:48:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-08-21T18:40:20Z","title":"HFX: Joint Design of Algorithms and Systems for Multi-SLO Serving and Fast Scaling","version":3},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-18T21:39:02.560962Z"},"links":{"cited_paper":"/paper/2405.06856","citing_paper":"/paper/2508.15919"},"observation_digest":"sha256:c28ab756a9d1bf251482dbb255cdb27284dce012d41eb56e033bae4f66f52116","observation_id":"eb2ce69d-74b8-4f9e-843e-e38cca40b433","resolution":{"observed_at":"2026-05-18T21:41:51.699897Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.06856","last_updated":"2024-05-11T00:00:23Z","snapshot_observed_at":"2026-08-18T13:43:37.657055Z","submitted_at":"2024-05-11T00:00:23Z","title":"Aladdin: Joint Placement and Scaling for SLO-Aware LLM Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.06856","snapshot_observed_at":"2026-08-05T16:21:09.246158Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.18736","last_updated":"2025-08-26T07:09:09Z","snapshot_observed_at":"2026-08-18T01:05:38.509949Z","submitted_at":"2025-08-26T07:09:09Z","title":"Rethinking Caching for LLM Serving Systems: Beyond Traditional Heuristics","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-05T16:21:09.246158Z"},"links":{"cited_paper":"/paper/2405.06856","citing_paper":"/paper/2508.18736"},"observation_digest":"sha256:60d18717de77935049ff68880152bf12b1335d4fcfd28423a88d9ad53697b6f4","observation_id":"bb85f735-5cd7-433c-8523-843499596bbd","resolution":{"observed_at":"2026-08-05T16:21:09.246158Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.06856","last_updated":"2024-05-11T00:00:23Z","snapshot_observed_at":"2026-08-18T13:43:37.657055Z","submitted_at":"2024-05-11T00:00:23Z","title":"Aladdin: Joint Placement and Scaling for SLO-Aware LLM Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.06856","snapshot_observed_at":"2026-08-15T16:30:35.001673Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.04827","last_updated":"2026-06-22T21:20:39Z","snapshot_observed_at":"2026-08-15T16:24:11.958686Z","submitted_at":"2025-09-05T05:58:16Z","title":"VoltanaLLM: Energy-Efficient and SLO-Aware Disaggregated LLM Serving via Adaptive Frequency Control and State-Space Routing","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T16:30:35.001673Z"},"links":{"cited_paper":"/paper/2405.06856","citing_paper":"/paper/2509.04827"},"observation_digest":"sha256:41ae2691f1dfa65c2b532f5ea9bae1036f897e9767370d167323047b0e0a951f","observation_id":"2dfaf51b-9bfc-4e5f-9075-42aae8d4b50a","resolution":{"observed_at":"2026-08-15T16:30:35.001673Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.06856","last_updated":"2024-05-11T00:00:23Z","snapshot_observed_at":"2026-08-18T13:43:37.657055Z","submitted_at":"2024-05-11T00:00:23Z","title":"Aladdin: Joint Placement and Scaling for SLO-Aware LLM Serving","version":1},"cited_work":{"arxiv_id":"2405.06856","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2405.06856","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Aladdin: Joint placement and scaling for slo-aware llm serving","venue":null,"work_id":"8f07f1e8-04fe-49b7-b64c-ce1fdd3001a3","year":2024},"citing_paper":{"arxiv_id":"2602.18755","last_updated":"2026-04-03T22:02:59Z","snapshot_observed_at":"2026-08-13T05:36:51.312686Z","submitted_at":"2026-02-21T08:31:49Z","title":"DualScale: Energy-Efficient Disaggregated LLM Serving via Phase-Aware Placement and DVFS","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-15T20:54:30.486885Z"},"links":{"cited_paper":"/paper/2405.06856","citing_paper":"/paper/2602.18755"},"observation_digest":"sha256:e49005a68670ac4816c4e9ad321fcaf611a70ca1a245310132613ec00feb7298","observation_id":"b2493590-e96c-467e-921b-b1d2e9d859bf","resolution":{"observed_at":"2026-05-15T20:56:37.081449Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2405.06856/citation-record","integrity":"/paper/2405.06856/integrity","json":"/paper/2405.06856/citation-record.json","paper":"/paper/2405.06856"},"outbound":[],"paper":{"arxiv_id":"2405.06856","last_updated":"2024-05-11T00:00:23Z","latest_version":1,"primary_category":"cs.DC","snapshot_observed_at":"2026-08-18T13:43:37.657055Z","submitted_at":"2024-05-11T00:00:23Z","title":"Aladdin: Joint Placement and Scaling for SLO-Aware LLM Serving"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 8 inbound Pith citation observations for arXiv:2405.06856."}