{"as_of":"2026-08-12T11:52:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:64526abfb69076cbb0e20f0c1b5e6392b3053328956d563118ffe24b24ee0dcd","coverage":[{"denominator":40,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":40,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T15:19:49.664592Z","state":"measured"},{"denominator":42,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":42,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-01T18:53:34.664986Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T16:58:42.766816Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"cited_work":{"arxiv_id":"2501.14205","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.14205","snapshot_observed_at":"2026-07-03T16:58:42.766816Z","title":"Serving long-context LLMs at the mobile edge: Test-time reinforcement learning-based model caching and inference offloading.arXiv preprint arXiv:2501.14205, 2025","venue":null,"work_id":"36b0fbf0-4470-4b3c-87ef-5ddc9db0f115","year":2025},"citing_paper":{"arxiv_id":"2606.17081","last_updated":"2026-06-11T21:03:04Z","snapshot_observed_at":"2026-07-06T23:52:42.390632Z","submitted_at":"2026-06-11T21:03:04Z","title":"The Price of Anarchy in Disaggregated Inference","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-27T04:52:48.847661Z"},"links":{"cited_paper":"/paper/2501.14205","citing_paper":"/paper/2606.17081"},"observation_digest":"sha256:102883e40359f6b33001b4c98a130cce0eb9c8a3151cb8a35bfebd0caab3d3e7","observation_id":"e0444af8-1be3-4fcc-8942-70a4602b0f22","resolution":{"observed_at":"2026-07-03T16:58:42.768430Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.14205","snapshot_observed_at":"2026-08-01T18:53:34.664986Z","title":"arXiv preprint arXiv:2501.14205 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.17175","last_updated":"2026-07-19T10:21:30Z","snapshot_observed_at":"2026-08-12T08:55:38.150680Z","submitted_at":"2026-07-19T10:21:30Z","title":"LMEdge: QoS-Aware LLM Inference Orchestration on Edge Clusters","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-01T18:53:34.664986Z"},"links":{"cited_paper":"/paper/2501.14205","citing_paper":"/paper/2607.17175"},"observation_digest":"sha256:03471da85d3599bb95991c520fc45fd0452c623346abcb96e232daf2052aef7b","observation_id":"aae32842-d1be-4de0-846a-af0918afdbb6","resolution":{"observed_at":"2026-08-01T18:53:34.664986Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2501.14205/citation-record","integrity":"/paper/2501.14205/integrity","json":"/paper/2501.14205/citation-record.json","paper":"/paper/2501.14205"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2401.07764","last_updated":"2024-02-16T19:15:31Z","snapshot_observed_at":"2026-07-06T17:15:42.343063Z","submitted_at":"2024-01-15T15:20:59Z","title":"When Large Language Model Agents Meet 6G Networks: Perception, Grounding, and Alignment","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.07764","snapshot_observed_at":"2026-08-10T15:19:49.515453Z","title":"When large language model agents meet 6g networks: Percep- tion, grounding, and alignment,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.515453Z"},"links":{"cited_paper":"/paper/2401.07764","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:c062a9e2b9e35b897a4235a49904fef0be31b09ae8ba75dc93f3cc0f6f86f13c","observation_id":"d84a032f-a3d7-4da2-bffa-b46f63043aba","resolution":{"observed_at":"2026-08-10T15:19:49.515453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.519977Z","title":"Language mod- els are few-shot learners,","venue":null,"work_id":null,"year":1901},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.519977Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:7e69d635ffb86a1a70a43b8e9aed8ef480e7b026994dea4a0a97a18f8f46e3d6","observation_id":"10c0479d-77f9-42de-b2dd-70ba141c59fe","resolution":{"observed_at":"2026-08-10T15:19:49.519977Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.12712","last_updated":"2023-04-13T20:41:31Z","snapshot_observed_at":"2026-08-03T04:49:15.195814Z","submitted_at":"2023-03-22T16:51:28Z","title":"Sparks of Artificial General Intelligence: Early experiments with GPT-4","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.12712","snapshot_observed_at":"2026-08-10T15:19:49.523574Z","title":"Sparks of artificial general intelligence: Early experiments with gpt-4,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.523574Z"},"links":{"cited_paper":"/paper/2303.12712","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:40cbab47805a912cdcb8e7e02a917a63796ddaf1f89a9e1071974612824366e7","observation_id":"ba4d54c5-4232-4607-8a4d-db3262dcc5e0","resolution":{"observed_at":"2026-08-10T15:19:49.523574Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.10825","last_updated":"2024-09-16T05:09:57Z","snapshot_observed_at":"2026-07-06T18:15:50.255628Z","submitted_at":"2024-05-17T14:46:13Z","title":"Large Language Model (LLM) for Telecommunications: A Comprehensive Survey on Principles, Key Techniques, and Opportunities","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.10825","snapshot_observed_at":"2026-08-10T15:19:49.528250Z","title":"Large language model (llm) for telecommu- nications: A comprehensive survey on principles, key techniques, and opportunities,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.528250Z"},"links":{"cited_paper":"/paper/2405.10825","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:5eae602d6f23b1798f99a2859f11f0159b99125d11b403ed78824ec7ee3006d6","observation_id":"07bcb665-a1fa-4fa8-bcd8-b820ba246c89","resolution":{"observed_at":"2026-08-10T15:19:49.528250Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.11299","last_updated":"2024-05-27T01:09:07Z","snapshot_observed_at":"2026-08-12T08:55:22.960994Z","submitted_at":"2024-05-18T14:00:04Z","title":"The CAP Principle for LLM Serving: A Survey of Long-Context Large Language Model Serving","version":2},"cited_work":{"arxiv_id":"2405.11299","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.11299","snapshot_observed_at":"2026-08-10T15:19:49.888392Z","title":"The CAP Principle for LLM Serving: A Survey of Long-Context Large Language Model Serving","venue":"cs.DB","work_id":"448382b1-4791-4728-8130-5247f0747f87","year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.532652Z"},"links":{"cited_paper":"/paper/2405.11299","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:bd2f59b3c741b2f11358d88185c13152b7b429d6b35c55e82abef986a2ab023d","observation_id":"7f4ee1f1-789c-46ed-873a-0ec185a40574","resolution":{"observed_at":"2026-08-10T15:19:49.892687Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.17245","last_updated":"2024-05-27T15:01:04Z","snapshot_observed_at":"2026-08-04T21:01:51.439514Z","submitted_at":"2024-05-27T15:01:04Z","title":"Galaxy: A Resource-Efficient Collaborative Edge AI System for In-situ Transformer Inference","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.17245","snapshot_observed_at":"2026-08-10T15:19:49.536634Z","title":"Galaxy: A resource-efficient collaborative edge ai system for in-situ transformer inference,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.536634Z"},"links":{"cited_paper":"/paper/2405.17245","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:6b342a1b111be7b49fdbb39059d46c7bee07b9c194e492d01a2572535bce0d87","observation_id":"324835ab-36ac-41f0-a772-95bd158246ae","resolution":{"observed_at":"2026-08-10T15:19:49.536634Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:50.116496Z","title":"Retention-aware container caching for serverless edge computing,","venue":null,"work_id":"f9678b39-dc27-4cfd-9be9-70fa5e3903fd","year":2022},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.541162Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:0db67826144b4de6b1d73d60b1f57bbcdf85f14f2a4da05526cac4fdaf8b3532","observation_id":"c30e738c-4fea-4cca-8947-4a9a28a11baa","resolution":{"observed_at":"2026-08-10T15:19:50.120400Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.544636Z","title":"Cache- enabled federated learning systems,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.544636Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:790fadd81c997e932e9e62f7f86c6048f3980c9d77e79c8e6c6909097df0a1a2","observation_id":"f5daf523-da96-4009-b6ea-729c2d61fe2e","resolution":{"observed_at":"2026-08-10T15:19:49.544636Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:50.098132Z","title":"Edgeadaptor: Online configuration adaption, model selection and re- source provisioning for edge dnn inference serving at scale,","venue":null,"work_id":"ff5cea07-483f-4ba7-a401-8eb15bd16e26","year":2022},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.548117Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:2156f4d441fed8c81123be7c30f16c89ed0e23ebb6c052952094e8d2f5e9851c","observation_id":"be5ebe3d-bac9-4f38-9898-d668c14f884f","resolution":{"observed_at":"2026-08-10T15:19:50.102086Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:50.087340Z","title":"Cooperative service caching and workload scheduling in mobile edge computing,","venue":null,"work_id":"46c2f2f4-3995-4ae9-b8ce-9ec8209872f8","year":2020},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.551751Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:c6d0e414f5ade3a2a0f332f4222e792055989d497353572197f891732448fc09","observation_id":"bafd9c87-ef0f-48d7-a912-e7242e487e42","resolution":{"observed_at":"2026-08-10T15:19:50.091182Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.17565","last_updated":"2024-12-21T13:55:49Z","snapshot_observed_at":"2026-08-11T06:18:26.517803Z","submitted_at":"2024-06-25T14:02:08Z","title":"MemServe: Context Caching for Disaggregated LLM Serving with Elastic Memory Pool","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.17565","snapshot_observed_at":"2026-08-10T15:19:49.555305Z","title":"Memserve: Context caching for disaggregated llm serving with elastic memory pool,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.555305Z"},"links":{"cited_paper":"/paper/2406.17565","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:248fb283e9d079f604e6bd688e95d643e9b2aa18e58fa3dd5034e2054b5e4f38","observation_id":"4f9d75c6-b1d1-42e9-8a1b-d87f9a6764be","resolution":{"observed_at":"2026-08-10T15:19:49.555305Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.06799","last_updated":"2024-09-21T09:10:02Z","snapshot_observed_at":"2026-08-06T08:22:33.331966Z","submitted_at":"2024-06-10T21:08:39Z","title":"LLM-dCache: Improving Tool-Augmented LLMs with GPT-Driven Localized Data Caching","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.06799","snapshot_observed_at":"2026-08-10T15:19:49.560249Z","title":"Llm-dcache: Improving tool- augmented llms with gpt-driven localized data caching,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.560249Z"},"links":{"cited_paper":"/paper/2406.06799","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:403e82f0a3c2b59cb1a1199f7cb869d779b2b242a0aaa0815443d36b017cfde9","observation_id":"b926fd4d-4dd1-4781-bf88-329a2c568c07","resolution":{"observed_at":"2026-08-10T15:19:49.560249Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.564040Z","title":"Chain-of-thought prompting elicits reasoning in large language models,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.564040Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:a9cd360c9278328225aa3a992ebde451fca5fe5be8e6293d2ed80177618139ae","observation_id":"5c5236eb-bcb1-4540-9a57-1cc26038a826","resolution":{"observed_at":"2026-08-10T15:19:49.564040Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2203.11171","last_updated":"2023-03-07T17:57:37Z","snapshot_observed_at":"2026-07-06T12:50:22.773056Z","submitted_at":"2022-03-21T17:48:52Z","title":"Self-Consistency Improves Chain of Thought Reasoning in Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2203.11171","snapshot_observed_at":"2026-08-10T15:19:49.567437Z","title":"Self-consistency improves chain of thought reasoning in language models,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.567437Z"},"links":{"cited_paper":"/paper/2203.11171","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:1173574bced9c7560bc52e4f1150199bc21d19197aaa3740d4d3504af99f76f2","observation_id":"705d8997-317d-498a-9e27-59e210c9deb9","resolution":{"observed_at":"2026-08-10T15:19:49.567437Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-10T15:19:49.571013Z","title":"Prox- imal policy optimization algorithms,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.571013Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:195d0cd1f4a2bd1426aae47c42ad3eba8c88936f2784e9a5eaf25bb11ea54b07","observation_id":"c18bf34f-319d-459a-b120-e013e31636f6","resolution":{"observed_at":"2026-08-10T15:19:49.571013Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.04620","last_updated":"2025-08-31T18:32:59Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-07-05T16:23:20Z","title":"Learning to (Learn at Test Time): RNNs with Expressive Hidden States","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.04620","snapshot_observed_at":"2026-08-10T15:19:49.574849Z","title":"Learning to (learn at test time): Rnns with expressive hidden states,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.574849Z"},"links":{"cited_paper":"/paper/2407.04620","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:703e97154974cc4b878e93091a28b0ece3e59f78c5b12ac5e7090cb2c1be428a","observation_id":"0d234bdd-393a-4d3f-971b-85005cbf12ef","resolution":{"observed_at":"2026-08-10T15:19:49.574849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.16739","last_updated":"2025-06-04T07:22:47Z","snapshot_observed_at":"2026-08-12T08:54:25.054691Z","submitted_at":"2023-09-28T06:22:59Z","title":"Pushing Large Language Models to the 6G Edge: Vision, Challenges, and Opportunities","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.16739","snapshot_observed_at":"2026-08-10T15:19:49.578508Z","title":"Pushing large language models to the 6g edge: Vision, challenges, and opportunities,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.578508Z"},"links":{"cited_paper":"/paper/2309.16739","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:b5cc197d5d09b2ac2da97c94d5194ca3ef303e8197da90f05cb297700acddf01","observation_id":"dfe2d729-d1ce-43e7-80a6-73929e45fe0a","resolution":{"observed_at":"2026-08-10T15:19:49.578508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14636","last_updated":"2024-05-23T14:41:22Z","snapshot_observed_at":"2026-08-10T08:07:18.214968Z","submitted_at":"2024-05-23T14:41:22Z","title":"PerLLM: Personalized Inference Scheduling with Edge-Cloud Collaboration for Diverse LLM Services","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.14636","snapshot_observed_at":"2026-08-10T15:19:49.582503Z","title":"Perllm: Personalized inference scheduling with edge-cloud collaboration for diverse llm services,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.582503Z"},"links":{"cited_paper":"/paper/2405.14636","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:5a20b568098369ea42fcf872c85aff505656d87688acbbe738248fd4f5f8d4fb","observation_id":"3d87c17d-e2a4-4fdd-8215-d8abd3b885b2","resolution":{"observed_at":"2026-08-10T15:19:49.582503Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:50.069761Z","title":"Titanic: Towards production federated learning with large language models,","venue":null,"work_id":"f71923fe-1c17-4270-baa7-e76c527b4274","year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.586311Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:631bdd46edcdd9422d23cdf126a6ba98ca2b3613ec9df6f89d7ef3d54e08096a","observation_id":"558a7db0-9503-4bee-8958-d1bb52c45b96","resolution":{"observed_at":"2026-08-10T15:19:50.073413Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:50.058241Z","title":"Generative inference of large language models in edge computing: An energy efficient approach,","venue":null,"work_id":"49ccbfe6-f60d-41a2-a239-cf00eefb48e2","year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.590170Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:6aae4db9470fc2bcb08005a3cb46a50652db76bef313b63ea3625f39c5c07b32","observation_id":"d75add3e-0cf4-4cfb-8a7d-7cfc7c88d351","resolution":{"observed_at":"2026-08-10T15:19:50.062046Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:50.047297Z","title":"Two time-scale joint service caching and task offloading for uav-assisted mobile edge computing,","venue":null,"work_id":"52861e28-a68e-4429-a7b7-015c98c059c2","year":2022},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.593639Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:05e3ff2954ea005f74910bd20a57355409bb17749e925155c5233687324760f8","observation_id":"ddf08fb1-3e2e-4b54-96ff-e77a2abb545e","resolution":{"observed_at":"2026-08-10T15:19:50.051289Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.14204","last_updated":"2026-05-20T02:47:49Z","snapshot_observed_at":"2026-07-06T18:03:42.703304Z","submitted_at":"2024-04-22T14:13:36Z","title":"TrimCaching: Parameter-sharing Edge Caching for AI Model Downloading","version":5},"cited_work":{"arxiv_id":"2404.14204","doi":null,"metadata_source":"pith","pith_arxiv_id":"2404.14204","snapshot_observed_at":"2026-08-10T15:19:49.783689Z","title":"TrimCaching: Parameter-sharing Edge Caching for AI Model Downloading","venue":"cs.NI","work_id":"c544a27a-f5da-4e44-ab1f-39214a7b65ce","year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.596984Z"},"links":{"cited_paper":"/paper/2404.14204","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:47aad159e230e8403946c7ae24bfa3bc9228cd48295ab4776f0ad17072d299b2","observation_id":"20ac9efd-3b09-4f68-b21c-f5e6764c2d81","resolution":{"observed_at":"2026-08-10T15:19:49.790517Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:50.035262Z","title":"A3c-based computation offloading and service caching in cloud-edge computing networks,","venue":null,"work_id":"896ddb85-db73-4ca9-be99-07e1e75ca750","year":2022},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.600715Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:58ec3f7533998a6c329e816a475dc7f83858f6395fff90e9b738284b32d2640c","observation_id":"c27acdd5-8e20-49ec-aa36-8258fa3e4b5e","resolution":{"observed_at":"2026-08-10T15:19:50.040012Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:50.022845Z","title":"Deepcache: A deep learning based framework for content caching,","venue":null,"work_id":"a0ae8d69-04b4-47e5-bb73-66d31fc1c3d7","year":2018},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.604274Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:9e5e667f68f53ca3ef3dd2820962bab8518d12ab456d5fce22af92838c723e64","observation_id":"e482c0ba-9f64-42f9-ad97-5883a39ea08f","resolution":{"observed_at":"2026-08-10T15:19:50.027327Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:50.010558Z","title":"Deep reinforcement learning-based computation offloading and distributed edge service caching for mobile edge computing,","venue":null,"work_id":"0d4e2cf6-21e8-441f-a64b-bb756533b55f","year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.607789Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:52028afcaa7e087148202b463b50c95736bb3ebd6ca5a45901a1553c138530d5","observation_id":"73400278-8708-4782-9b07-d901433e05c5","resolution":{"observed_at":"2026-08-10T15:19:50.015269Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.995327Z","title":"Neighboring- aware caching in heterogeneous edge networks by actor-attention-critic learning,","venue":null,"work_id":"65b7a91f-3e87-4ed4-b490-0c85df59155c","year":2021},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.611078Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:34c7bf8c6dce119aeb03b97e5370bfb4990eda8b6e5ac908fcca7fed03145d9f","observation_id":"d1b9ca53-b7dd-4cb3-ab6e-296d6691b3ad","resolution":{"observed_at":"2026-08-10T15:19:50.000141Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.979716Z","title":"Cooperative task offloading and service caching for digital twin edge networks: A graph attention multi- agent reinforcement learning approach,","venue":null,"work_id":"76d376f2-03c6-415b-8dad-d3a6af88bc7e","year":2023},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.614544Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:10eeeeca0a8215429cb764b911e40221eb6824055c21474d6c8ed1ae661aba03","observation_id":"febaaaee-ece9-4f37-a2ca-6b25a529ef91","resolution":{"observed_at":"2026-08-10T15:19:49.984386Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.968106Z","title":"Large language models (llms) inference offloading and resource allocation in cloud-edge com- puting: An active inference approach,","venue":null,"work_id":"38c913a1-e7c7-4fd5-8383-79cafdddadfb","year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.618260Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:f3ea927fff50ef322c9caf75f290c1fc07ce3dcbe8a98fc1e3b45b6e16a76e5c","observation_id":"14ba85ae-987e-4d82-926f-cf3bb749626d","resolution":{"observed_at":"2026-08-10T15:19:49.972042Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.956407Z","title":"Are transformers universal approximators of sequence-to-sequence func- tions?","venue":null,"work_id":"d1812c66-24e8-4d51-911d-be9744efc789","year":2020},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.621562Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:60e03cf337a8691ed31fbbe40ce7d7b275a9af67105f11c74583efdf8ddeac43","observation_id":"4b35c3fb-0782-46ce-a0cc-d39fc95bd21c","resolution":{"observed_at":"2026-08-10T15:19:49.960600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13571","last_updated":"2024-06-06T12:18:56Z","snapshot_observed_at":"2026-08-10T03:44:23.078835Z","submitted_at":"2023-10-20T15:09:46Z","title":"Why Can Large Language Models Generate Correct Chain-of-Thoughts?","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13571","snapshot_observed_at":"2026-08-10T15:19:49.624860Z","title":"Why can large language models generate correct chain-of-thoughts?","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.624860Z"},"links":{"cited_paper":"/paper/2310.13571","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:67f3a4d503df46c47924430f99aacf903ba2ed969d79f06589219273949826c1","observation_id":"39190684-d285-4a76-9702-86ce715496c3","resolution":{"observed_at":"2026-08-10T15:19:49.624860Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.09960","last_updated":"2023-09-13T18:52:02Z","snapshot_observed_at":"2026-08-10T10:01:29.858506Z","submitted_at":"2023-04-19T20:45:01Z","title":"A Latent Space Theory for Emergent Abilities in Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.09960","snapshot_observed_at":"2026-08-10T15:19:49.628482Z","title":"A latent space theory for emergent abilities in large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.628482Z"},"links":{"cited_paper":"/paper/2304.09960","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:703d231abafd809cb142ca30a78ef243d820ba9f4711b2599fc82dcb8c41199c","observation_id":"1138a40f-56d7-45f9-aaa3-3df80e5f49cb","resolution":{"observed_at":"2026-08-10T15:19:49.628482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05826","last_updated":"2024-05-31T14:14:00Z","snapshot_observed_at":"2026-07-06T17:41:59.818493Z","submitted_at":"2024-03-09T07:37:13Z","title":"Cached Model-as-a-Resource: Provisioning Large Language Model Agents for Edge Intelligence in Space-air-ground Integrated Networks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05826","snapshot_observed_at":"2026-08-10T15:19:49.632041Z","title":"Cached model-as-a-resource: Provisioning large language model agents for edge intelligence in space-air-ground integrated networks,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.632041Z"},"links":{"cited_paper":"/paper/2403.05826","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:fed7f6ae7cb569e382a6d6b124cd85396856ec2faf89d4d81a7c642089249ee6","observation_id":"7abc262e-0c74-49b5-b4b3-2dbab7d53f38","resolution":{"observed_at":"2026-08-10T15:19:49.632041Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.635629Z","title":"Imagebind: One embedding space to bind them all,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.635629Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:1086dee465ea387eaebe2a0c6d4979eb3e31337d6bd90a03882ddea99ae25f1c","observation_id":"6bff87d8-a937-46a8-afd5-3eac5f51219f","resolution":{"observed_at":"2026-08-10T15:19:49.635629Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1608.01413","last_updated":"2016-08-20T11:50:41Z","snapshot_observed_at":"2026-08-09T10:12:36.401845Z","submitted_at":"2016-08-04T01:47:23Z","title":"Solving General Arithmetic Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1608.01413","snapshot_observed_at":"2026-08-10T15:19:49.639218Z","title":"Solving general arithmetic word problems,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.639218Z"},"links":{"cited_paper":"/paper/1608.01413","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:c3e4dfc4f76343e72e9247eceab51c86faae4074e204f136ce5b28b2291fae6b","observation_id":"cb3af874-46ee-4c30-8ae5-3dc7885dc6b0","resolution":{"observed_at":"2026-08-10T15:19:49.639218Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1803.05457","last_updated":"2018-03-14T18:04:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2018-03-14T18:04:21Z","title":"Think you have Solved Question Answering? Try ARC, the AI2 Reasoning Challenge","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1803.05457","snapshot_observed_at":"2026-08-10T15:19:49.643307Z","title":"Think you have solved question answering? try arc, the ai2 reasoning challenge,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.643307Z"},"links":{"cited_paper":"/paper/1803.05457","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:49e4ebdbf4cd48f4bc2c775ec5556d6116a8b343d1d1298368154f3282d061ec","observation_id":"9625d1c3-274b-4f21-b68a-0c4480f8d821","resolution":{"observed_at":"2026-08-10T15:19:49.643307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.648085Z","title":"Did aristotle use a laptop? a question answering benchmark with implicit reasoning strategies,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.648085Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:d03a7e66cc32e6b70410435f6800c327a728893559a0d5780a44de8924e56f4f","observation_id":"77168914-eb45-40cd-8636-c0a7241f0c22","resolution":{"observed_at":"2026-08-10T15:19:49.648085Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1811.00937","last_updated":"2019-03-15T18:02:58Z","snapshot_observed_at":"2026-08-10T07:37:55.457401Z","submitted_at":"2018-11-02T15:34:29Z","title":"CommonsenseQA: A Question Answering Challenge Targeting Commonsense Knowledge","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1811.00937","snapshot_observed_at":"2026-08-10T15:19:49.651683Z","title":"Commonsenseqa: A question answering challenge targeting commonsense knowledge,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.651683Z"},"links":{"cited_paper":"/paper/1811.00937","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:d2e6290fa4744991521b11d94d5296c934246ca8f10d8b6a9d3600cb11bd708b","observation_id":"4b0c23f5-f14f-4c7a-9454-88602f16b0ba","resolution":{"observed_at":"2026-08-10T15:19:49.651683Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2111.02840","last_updated":"2022-01-10T06:05:16Z","snapshot_observed_at":"2026-07-06T12:05:25.336577Z","submitted_at":"2021-11-04T12:59:55Z","title":"Adversarial GLUE: A Multi-Task Benchmark for Robustness Evaluation of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2111.02840","snapshot_observed_at":"2026-08-10T15:19:49.656199Z","title":"Adversarial glue: A multi-task benchmark for robustness evaluation of language models,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.656199Z"},"links":{"cited_paper":"/paper/2111.02840","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:e3d3f765410d8f1debd4110a3ed6092e2d8b103c9d2a3ae2480864b36d0a0455","observation_id":"81301fb8-5a9e-45f0-b363-1aa1c5868a03","resolution":{"observed_at":"2026-08-10T15:19:49.656199Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T15:19:49.932083Z","title":"Improve diverse text generation by self labeling conditional variational auto encoder,","venue":null,"work_id":"381b9b66-477e-438b-8e3e-2c032175f98d","year":2019},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.660633Z"},"links":{"citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:70d53bee085edbf5914f830e306f978529841e74db0021c64a5d3665dcc30bb3","observation_id":"305919c9-70f4-428d-8177-f427c898aea5","resolution":{"observed_at":"2026-08-10T15:19:49.936143Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-08-07T01:45:38.840969Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-08-10T15:19:49.664592Z","title":"Training verifiers to solve math word problems,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-10T15:19:49.664592Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2501.14205"},"observation_digest":"sha256:ea68cbe0ce810b2e288962ee20cf661c5aacd3ea9c788d998ead2c4e5d76c24b","observation_id":"7234443e-7d09-4bb8-b0cb-9494d4685cf6","resolution":{"observed_at":"2026-08-10T15:19:49.664592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2501.14205","last_updated":"2025-01-24T03:21:20Z","latest_version":1,"primary_category":"cs.NI","snapshot_observed_at":"2026-08-12T08:55:23.913616Z","submitted_at":"2025-01-24T03:21:20Z","title":"Serving Long-Context LLMs at the Mobile Edge: Test-Time Reinforcement Learning-based Model Caching and Inference Offloading"},"reference_resolution":{"displayed":40,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":24,"verified_exact":2,"verified_fuzzy":14},"total_outbound_references":40},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 40 of 40 outbound references and 2 inbound Pith citation observations for arXiv:2501.14205."}