{"as_of":"2026-08-09T21:36:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c938d8bf0d12c8c30ec2dab3180b0e0d1c6e109484219342b5118ddae4f98a24","coverage":[{"denominator":15,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":15,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T18:45:05.315427Z","state":"measured"},{"denominator":15,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":15,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2507.08877/citation-record","integrity":"/paper/2507.08877/integrity","json":"/paper/2507.08877/citation-record.json","paper":"/paper/2507.08877"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:45:06.298999Z","title":null,"venue":null,"work_id":"01b47969-2e86-4b0e-9537-c1b59eb73514","year":2023},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:03.763685Z"},"links":{"citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:ce11de3b4d66c371d60042cb6b066d101c96f58140b9c64bdb41a55e37b2d736","observation_id":"538ad4a7-3eac-4406-be64-a9055c1df82c","resolution":{"observed_at":"2026-08-06T18:45:06.397855Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.08590","last_updated":"2022-10-18T08:28:20Z","snapshot_observed_at":"2026-08-08T01:16:39.687991Z","submitted_at":"2022-10-16T17:24:06Z","title":"Zero-Shot Learners for Natural Language Understanding via a Unified Multiple Choice Perspective","version":2},"cited_work":{"arxiv_id":"2210.08590","doi":null,"metadata_source":"pith","pith_arxiv_id":"2210.08590","snapshot_observed_at":"2026-08-06T18:45:05.841649Z","title":"Zero-Shot Learners for Natural Language Understanding via a Unified Multiple Choice Perspective","venue":"cs.CL","work_id":"196c0b14-e494-4c77-8723-3af6b1bc0e89","year":2022},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:03.876769Z"},"links":{"cited_paper":"/paper/2210.08590","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:96c932e494afdada7bcc4827c8226f9cf298845910c9ada4a70528ed0680b53d","observation_id":"6ca7e70b-5d9d-442b-8380-0454ab2acc6d","resolution":{"observed_at":"2026-08-06T18:45:05.925584Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1503.02531","last_updated":"2015-03-09T15:44:49Z","snapshot_observed_at":"2026-07-06T04:11:24.157003Z","submitted_at":"2015-03-09T15:44:49Z","title":"Distilling the Knowledge in a Neural Network","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1503.02531","snapshot_observed_at":"2026-08-06T18:45:03.993003Z","title":null,"venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:03.993003Z"},"links":{"cited_paper":"/paper/1503.02531","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:389d12396fc9f02ffc23e8eb7b2cebf7c748be7e337d5cb732bd820389640299","observation_id":"cf33eb8b-5f02-4dab-9d67-e84d2b6b3cef","resolution":{"observed_at":"2026-08-06T18:45:03.993003Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:45:06.137063Z","title":"B., Mann, B., Ryder, N., Subbiah, M., Kaplan, J., Dhariwal, P.,","venue":null,"work_id":"57d99208-e7a9-4369-a6a7-a8d22ce1fea9","year":2020},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:04.137058Z"},"links":{"citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:2e460b436c6ac02c0338b9094bfa9a584191809c742d8e858f8bf248fe903144","observation_id":"ff1a775b-9e2d-4359-a0cd-8679b91733d7","resolution":{"observed_at":"2026-08-06T18:45:06.215793Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T18:45:06.006174Z","title":null,"venue":null,"work_id":"552174f9-1413-433c-98c6-e8b3c69baeb2","year":2019},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:04.256735Z"},"links":{"citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:39722eaedaf843b12065d15c5ca7904a4963f55d6f023fb909fc8dbad9e5d968","observation_id":"32c81672-fa4b-4c24-af0d-b4d5f128e2ec","resolution":{"observed_at":"2026-08-06T18:45:06.070385Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.04623","last_updated":"2023-08-08T23:29:55Z","snapshot_observed_at":"2026-08-09T00:36:07.880808Z","submitted_at":"2023-08-08T23:29:55Z","title":"Accelerating LLM Inference with Staged Speculative Decoding","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.04623","snapshot_observed_at":"2026-08-06T18:45:04.310197Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:04.310197Z"},"links":{"cited_paper":"/paper/2308.04623","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:b1e8965e2ab09d77a49c8169dad345157d1aa7b13027a19657f50bc084542625","observation_id":"e4e9bb28-4121-491e-b663-aa79566a6431","resolution":{"observed_at":"2026-08-06T18:45:04.310197Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.15766","last_updated":"2025-02-26T11:47:58Z","snapshot_observed_at":"2026-08-08T01:15:47.353870Z","submitted_at":"2024-08-28T12:59:12Z","title":"Learning Harmonized Representations for Speculative Sampling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.15766","snapshot_observed_at":"2026-08-06T18:45:04.421303Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:04.421303Z"},"links":{"cited_paper":"/paper/2408.15766","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:8492b3a4b37519444eac4691889c99cfc62d44e2582c9e3f3b55fc3436dfcd2a","observation_id":"01f2a3c3-9e6b-4c52-8a6e-0c09e258d9af","resolution":{"observed_at":"2026-08-06T18:45:04.421303Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.11309","last_updated":"2025-06-12T21:15:58Z","snapshot_observed_at":"2026-08-08T01:16:03.022130Z","submitted_at":"2025-06-12T21:15:58Z","title":"SwiftSpec: Ultra-Low Latency LLM Decoding by Scaling Asynchronous Speculative Decoding","version":1},"cited_work":{"arxiv_id":"2506.11309","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.11309","snapshot_observed_at":"2026-08-06T18:45:05.618170Z","title":"SwiftSpec: Ultra-Low Latency LLM Decoding by Scaling Asynchronous Speculative Decoding","venue":"cs.DC","work_id":"b97de753-2dfb-4076-b1f7-2a3bb28af0ba","year":2025},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:04.545787Z"},"links":{"cited_paper":"/paper/2506.11309","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:1bf918cdfa04809b6d89803d794d8e0a4d964587495581fe9f73e32ba754a9fc","observation_id":"73e76fe6-46b5-44f6-a9f2-eec8c7d42f7f","resolution":{"observed_at":"2026-08-06T18:45:05.693179Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.15921","last_updated":"2025-03-20T07:57:57Z","snapshot_observed_at":"2026-08-07T16:49:52.691330Z","submitted_at":"2025-03-20T07:57:57Z","title":"SPIN: Accelerating Large Language Model Inference with Heterogeneous Speculative Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.15921","snapshot_observed_at":"2026-08-06T18:45:04.626816Z","title":"H., Su, Z., & Deng, J","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:04.626816Z"},"links":{"cited_paper":"/paper/2503.15921","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:08dabf5fbb86360aa28d8764cda7f2392d1ce91060a3db80e48dbc42558881db","observation_id":"48d3e85d-8e87-428f-ab41-7c385ccae56f","resolution":{"observed_at":"2026-08-06T18:45:04.626816Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.05276","last_updated":"2024-12-09T01:44:10Z","snapshot_observed_at":"2026-08-02T07:37:21.335822Z","submitted_at":"2024-11-08T02:21:19Z","title":"GPT Semantic Cache: Reducing LLM Costs and Latency via Semantic Embedding Caching","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.05276","snapshot_observed_at":"2026-08-06T18:45:04.721093Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:04.721093Z"},"links":{"cited_paper":"/paper/2411.05276","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:6acab1f02765deb47f590d1633190df96bc9bf775198cc0a9f6521389a259d01","observation_id":"37e2e66c-f579-4cd9-b8ee-9d21bc856520","resolution":{"observed_at":"2026-08-06T18:45:04.721093Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.17603","last_updated":"2025-03-22T01:17:56Z","snapshot_observed_at":"2026-08-07T16:44:44.142130Z","submitted_at":"2025-03-22T01:17:56Z","title":"A Generative Caching System for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.17603","snapshot_observed_at":"2026-08-06T18:45:04.856861Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:04.856861Z"},"links":{"cited_paper":"/paper/2503.17603","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:03742a3a7ab4f75778d22c6c83da0cb871c111c5e4923ee905e4c7c3d41e3064","observation_id":"6df44339-6fa6-4348-8bc0-01fdb1654847","resolution":{"observed_at":"2026-08-06T18:45:04.856861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.01173","last_updated":"2024-02-02T06:34:11Z","snapshot_observed_at":"2026-08-09T19:48:04.723755Z","submitted_at":"2024-02-02T06:34:11Z","title":"Efficient Prompt Caching via Embedding Similarity","version":1},"cited_work":{"arxiv_id":"2402.01173","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.01173","snapshot_observed_at":"2026-08-06T18:45:05.427249Z","title":"Efficient Prompt Caching via Embedding Similarity","venue":"cs.CL","work_id":"25d42777-f976-4b6d-8ad3-46d279f48867","year":2024},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:04.963570Z"},"links":{"cited_paper":"/paper/2402.01173","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:138ed27e42fc3eab2aba6a39ca4553e63595a2bf17e3927c9754cc5634e923c1","observation_id":"5f9e67e7-ee05-4147-acc0-a1df1d10d38c","resolution":{"observed_at":"2026-08-06T18:45:05.479876Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.07017","last_updated":"2024-12-09T21:53:10Z","snapshot_observed_at":"2026-07-06T20:04:15.943803Z","submitted_at":"2024-12-09T21:53:10Z","title":"Asynchronous LLM Function Calling","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.07017","snapshot_observed_at":"2026-08-06T18:45:05.106051Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:05.106051Z"},"links":{"cited_paper":"/paper/2412.07017","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:3ea61a79b7f8398ffa2651eb9d21f5c6bbc2021a5e9ebe138617d2d1602d3708","observation_id":"32d19d2f-13a5-417f-8f8f-bb368c7008d3","resolution":{"observed_at":"2026-08-06T18:45:05.106051Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.20438","last_updated":"2025-05-26T18:34:07Z","snapshot_observed_at":"2026-08-07T13:52:30.019778Z","submitted_at":"2025-05-26T18:34:07Z","title":"HAMburger: Accelerating LLM Inference via Token Smashing","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.20438","snapshot_observed_at":"2026-08-06T18:45:05.226057Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:05.226057Z"},"links":{"cited_paper":"/paper/2505.20438","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:30f9ca237511087bc9b0f5bd7a4e7676382e7a4ff369903ffb30da906b4a42fc","observation_id":"36fd8cbc-ff44-4fcf-97d2-6be966b45356","resolution":{"observed_at":"2026-08-06T18:45:05.226057Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.20330","last_updated":"2025-06-23T03:05:26Z","snapshot_observed_at":"2026-08-07T17:42:07.783822Z","submitted_at":"2025-02-27T17:59:36Z","title":"RAPID: Long-Context Inference with Retrieval-Augmented Speculative Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.20330","snapshot_observed_at":"2026-08-06T18:45:05.315427Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T18:45:05.315427Z"},"links":{"cited_paper":"/paper/2502.20330","citing_paper":"/paper/2507.08877"},"observation_digest":"sha256:8314a00791c942057e7e0304bc76ed568e37b38833bf271b180cb16537222ece","observation_id":"47064d09-b258-48ab-9e03-b1e2c6ac7ac2","resolution":{"observed_at":"2026-08-06T18:45:05.315427Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.08877","last_updated":"2025-07-10T04:44:47Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-09T14:45:42.221113Z","submitted_at":"2025-07-10T04:44:47Z","title":"ODIA: Oriented Distillation for Inline Acceleration of LLM-based Function Calling"},"reference_resolution":{"displayed":15,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":11,"verified_exact":3,"verified_fuzzy":1},"total_outbound_references":15},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 15 of 15 outbound references and 0 inbound Pith citation observations for arXiv:2507.08877."}