{"as_of":"2026-08-09T19:08:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2524b6f9abf4002037955cc963d46ee6b077363b3beb9f3b7b1f17ac88ed220c","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":23,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":23,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":23,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":23,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T22:33:09.714455Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T10:09:44.136682Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2410.15761","last_updated":"2026-06-03T09:17:37Z","snapshot_observed_at":"2026-08-01T14:18:40.066269Z","submitted_at":"2024-10-21T08:21:00Z","title":"Optimal Query Allocation in Extractive QA with LLMs: A Learning-to-Defer Framework with Theoretical Guarantees","version":4},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-05-23T18:49:39.718108Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2410.15761"},"observation_digest":"sha256:804c1f0ae355a9b11b8052cc5e346433eefe2181b559e47b66e1e0dd591f2f8e","observation_id":"19bbf7df-5b0e-49c9-833e-5c1a8658f84b","resolution":{"observed_at":"2026-05-23T18:53:21.433797Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-08-08T22:33:09.714455Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.04506","last_updated":"2025-02-06T21:13:44Z","snapshot_observed_at":"2026-08-09T04:44:07.719948Z","submitted_at":"2025-02-06T21:13:44Z","title":"When One LLM Drools, Multi-LLM Collaboration Rules","version":1},"reference_index":129,"source":"arxiv_source","source_observed_at":"2026-08-08T22:33:09.714455Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2502.04506"},"observation_digest":"sha256:35d6f9b66102ab26f283c46859d261020ad923f045e889eb4493f1122eb7c94f","observation_id":"7c833d01-c69a-4345-87d9-481d769f9399","resolution":{"observed_at":"2026-08-08T22:33:09.714455Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-08-07T04:30:43.727123Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.10525","last_updated":"2025-06-12T09:43:48Z","snapshot_observed_at":"2026-08-08T21:26:08.789047Z","submitted_at":"2025-06-12T09:43:48Z","title":"AdaptiveLLM: A Framework for Selecting Optimal Cost-Efficient LLM for Code-Generation Based on CoT Length","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T04:30:43.727123Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2506.10525"},"observation_digest":"sha256:f13732f0ab45b002c0a8cba9ce2ce243e88af597763955b4daa4400f41af7d54","observation_id":"eae85b95-1cc1-44ad-964a-a3da894f18e1","resolution":{"observed_at":"2026-08-07T04:30:43.727123Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-08-06T22:05:29.245616Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.22716","last_updated":"2025-06-28T01:52:50Z","snapshot_observed_at":"2026-08-06T21:57:22.791105Z","submitted_at":"2025-06-28T01:52:50Z","title":"BEST-Route: Adaptive LLM Routing with Test-Time Optimal Compute","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-06T22:05:29.245616Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2506.22716"},"observation_digest":"sha256:40a3eea793a1facb1e9ecf5e070ac687fc4b1a883d4d2be086decfdc412faeb4","observation_id":"e511c90b-13a8-4ed9-984b-4cff45cb6ba5","resolution":{"observed_at":"2026-08-06T22:05:29.245616Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2603.27098","last_updated":"2026-04-04T02:18:57Z","snapshot_observed_at":"2026-08-02T12:17:19.632566Z","submitted_at":"2026-03-28T02:37:36Z","title":"Ensemble-Based Uncertainty Estimation for Code Correctness Estimation","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-14T22:52:58.524934Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2603.27098"},"observation_digest":"sha256:a018ffbef9026f01585a1e8ba720032af935d57b6c3dd8b0c7703d349a036d89","observation_id":"cb82fa50-a591-4d89-b43f-a0422c9afd68","resolution":{"observed_at":"2026-05-14T22:53:14.009570Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2604.02367","last_updated":"2026-03-26T15:57:46Z","snapshot_observed_at":"2026-07-06T22:51:44.663367Z","submitted_at":"2026-03-26T15:57:46Z","title":"Evaluating Small Language Models for Front-Door Routing: A Harmonized Benchmark and Synthetic-Traffic Experiment","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-15T00:38:17.474611Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2604.02367"},"observation_digest":"sha256:eb42a72f347ac77b31dcbc563aab4d46f3d13bcab377edab4d80976fda9f4fc7","observation_id":"b8eb085c-a640-47c5-aae6-f94e5d34b11b","resolution":{"observed_at":"2026-05-15T00:38:23.073475Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2604.15728","last_updated":"2026-04-17T06:02:27Z","snapshot_observed_at":"2026-07-06T23:03:12.000381Z","submitted_at":"2026-04-17T06:02:27Z","title":"Privacy-Preserving LLMs Routing","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-10T08:32:43.015370Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2604.15728"},"observation_digest":"sha256:ee45bd22f527409479b0cd69da0bda983e3127ff09a3f6915452795150b19e7e","observation_id":"9e2a85f5-e476-4d01-be07-0cd8f57165a7","resolution":{"observed_at":"2026-05-10T08:32:51.625602Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2604.18612","last_updated":"2026-04-14T07:35:37Z","snapshot_observed_at":"2026-07-31T08:42:18.257117Z","submitted_at":"2026-04-14T07:35:37Z","title":"Agent-GWO: Collaborative Agents for Dynamic Prompt Optimization in Large Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T14:24:55.735201Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2604.18612"},"observation_digest":"sha256:fb01a7640ecf4fe45fc660cd776779ebfd97fd3547f582607ed453c8688d6220","observation_id":"fedde61c-d6d3-490d-bcfa-a014536387a7","resolution":{"observed_at":"2026-05-10T14:25:29.196389Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2604.23477","last_updated":"2026-05-14T05:59:30Z","snapshot_observed_at":"2026-07-06T23:09:43.659053Z","submitted_at":"2026-04-26T00:05:53Z","title":"SEMA-SQL: Beyond Traditional Relational Querying with Large Language Models","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-08T05:11:46.991452Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2604.23477"},"observation_digest":"sha256:0c0f2b8df03459f0add703088ddba2f956269d78c49ad36038ab7df216665a39","observation_id":"0ec3e192-c478-4372-8ad1-d81ee46e003d","resolution":{"observed_at":"2026-05-11T21:31:17.421068Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2604.23477","last_updated":"2026-05-14T05:59:30Z","snapshot_observed_at":"2026-07-06T23:09:43.659053Z","submitted_at":"2026-04-26T00:05:53Z","title":"SEMA-SQL: Beyond Traditional Relational Querying with Large Language Models","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-15T06:53:11.604043Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2604.23477"},"observation_digest":"sha256:17e91c85741f35f1c74155402013d3bc9f6b3e29976588ad2d18c7ded8b141ba","observation_id":"6e7fee11-b8dd-4822-8840-2a525bf18d90","resolution":{"observed_at":"2026-05-15T06:55:10.577413Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2605.12340","last_updated":"2026-05-29T07:49:08Z","snapshot_observed_at":"2026-08-02T07:56:41.241469Z","submitted_at":"2026-05-12T16:19:44Z","title":"Online Learning-to-Defer with Varying Experts","version":1},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-05-13T04:05:52.853863Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2605.12340"},"observation_digest":"sha256:bedccaf192cda9ea6d88811b5a1ac3dc2275eb1675b9d9d1a78acb60b2e870b8","observation_id":"6b2207cc-b58c-4b5b-83e5-b26aefca08f3","resolution":{"observed_at":"2026-05-13T04:07:13.156453Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2605.12340","last_updated":"2026-05-29T07:49:08Z","snapshot_observed_at":"2026-08-02T07:56:41.241469Z","submitted_at":"2026-05-12T16:19:44Z","title":"Online Learning-to-Defer with Varying Experts","version":2},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-05-21T08:07:35.959083Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2605.12340"},"observation_digest":"sha256:aef0e13783b35c34c7ba6dd789b83e194d034f0f630b30a139fef0696ac72887","observation_id":"f2a63b40-57a6-4ff6-a951-03a2841e180d","resolution":{"observed_at":"2026-05-21T08:09:51.328157Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2606.06924","last_updated":"2026-06-05T05:42:00Z","snapshot_observed_at":"2026-08-06T11:49:12.989321Z","submitted_at":"2026-06-05T05:42:00Z","title":"From Sampled Outcomes to Capability Distributions: Rethinking Supervision for LLM Routing","version":1},"reference_index":151,"source":"arxiv_source","source_observed_at":"2026-06-27T22:54:28.452796Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2606.06924"},"observation_digest":"sha256:98418c26d77a04674bb7fe267a5c4ff9f17bcfc6cb965eaca2e9cba5a39507ea","observation_id":"b78d259c-7345-4ac3-b67c-ed8d22174803","resolution":{"observed_at":"2026-07-02T16:17:08.804014Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2606.18774","last_updated":"2026-06-19T09:33:33Z","snapshot_observed_at":"2026-08-06T18:36:21.111914Z","submitted_at":"2026-06-17T07:35:10Z","title":"RouteJudge: An Open Platform for Reproducible and Preference-Aware LLM Routing","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-26T21:17:29.543901Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2606.18774"},"observation_digest":"sha256:a58d75eeb6cd9ccbeb654cebf4251795c084bc3603f0ae1f73d2157bef6e0486","observation_id":"1c332ccd-0c19-4907-aa8b-53723d7de47d","resolution":{"observed_at":"2026-07-04T00:19:13.488572Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2606.20295","last_updated":"2026-07-24T07:54:17Z","snapshot_observed_at":"2026-08-02T10:48:55.986116Z","submitted_at":"2026-06-18T14:33:09Z","title":"Token-Operations-Oriented Inference Optimization Techniques for Large Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-26T16:15:22.543601Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2606.20295"},"observation_digest":"sha256:c94e16fc5cb257c3a0e5d16b3f6be484f64f4b55b285167838db3972cc3dc2c8","observation_id":"5dfd33e0-ec31-465d-bd4c-cccfddcbbd51","resolution":{"observed_at":"2026-07-04T05:09:36.813433Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-08-02T10:49:00.720574Z","title":"Large Language Model Cascades with Mixture of Thoughts Rep- resentations for Cost-efficient Reasoning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.20295","last_updated":"2026-07-24T07:54:17Z","snapshot_observed_at":"2026-08-02T10:48:55.986116Z","submitted_at":"2026-06-18T14:33:09Z","title":"Token-Operations-Oriented Inference Optimization Techniques for Large Models","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-02T10:49:00.720574Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2606.20295"},"observation_digest":"sha256:87ad92753865efca25a47d519168080b942f701a4b4ced8d5445adfd00651216","observation_id":"79bf9a32-a2f2-425a-9005-ddfeea326fb8","resolution":{"observed_at":"2026-08-02T10:49:00.720574Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":"2310.03094","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-04T10:09:44.136682Z","title":"Large language model cascades with mixture of thoughts representations for cost-efficient reasoning","venue":null,"work_id":"29ee2d8e-d434-44e2-b30f-f7b9198ee961","year":2023},"citing_paper":{"arxiv_id":"2606.22840","last_updated":"2026-06-22T04:27:45Z","snapshot_observed_at":"2026-08-08T08:58:39.390391Z","submitted_at":"2026-06-22T04:27:45Z","title":"RLM-Cascade: Response-Level Speculative Decoding for Cost-Efficient LLM API Serving","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-26T09:10:47.370434Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2606.22840"},"observation_digest":"sha256:c922c75473400f1531ad47d0de5b65e95f3c831f8cc0ab27296c6809311c17a8","observation_id":"d116d5c0-8419-4e0c-ac75-4b9ccc5ee714","resolution":{"observed_at":"2026-07-04T10:09:44.138484Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-08-01T13:11:02.437028Z","title":"Large language model cas- cades with mixture of thoughts representations for cost-efficient reasoning.arXiv preprint arXiv:2310.03094, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.19215","last_updated":"2026-07-21T15:47:35Z","snapshot_observed_at":"2026-08-07T09:39:38.164665Z","submitted_at":"2026-07-21T15:47:35Z","title":"HACO: Hedged Agent Computing for Reliable LLM Systems","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-01T13:11:02.437028Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2607.19215"},"observation_digest":"sha256:8805eda1bdcc10a4aac48377eb4183b017db7c11be0788975cd6852972205127","observation_id":"7ef6c68d-2314-4cae-ab4f-1cefe6c15f45","resolution":{"observed_at":"2026-08-01T13:11:02.437028Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-31T03:38:02.101711Z","title":"Large language model cas- cades with mixture of thought representations for cost-efficient reasoning.arXiv preprint arXiv:2310.03094,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.25018","last_updated":"2026-07-31T02:55:33Z","snapshot_observed_at":"2026-08-06T02:14:46.802306Z","submitted_at":"2026-07-27T19:19:58Z","title":"Conformal Cascade: Distribution-Free Accuracy Guarantees for Multi-Tier LLM Inference","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-31T03:38:02.101711Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2607.25018"},"observation_digest":"sha256:e51586ea4018a561ea9bea21cf6db09a0ce89a566b83d10953b0e5487e1f8983","observation_id":"09860c77-6b1f-4a43-bc81-491d1e7be16f","resolution":{"observed_at":"2026-07-31T03:38:02.101711Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-08-03T01:49:19.070875Z","title":"Large language model cas- cades with mixture of thought representations for cost-efficient reasoning.arXiv preprint arXiv:2310.03094,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.25018","last_updated":"2026-07-31T02:55:33Z","snapshot_observed_at":"2026-08-06T02:14:46.802306Z","submitted_at":"2026-07-27T19:19:58Z","title":"Conformal Cascade: Distribution-Free Accuracy Guarantees for Multi-Tier LLM Inference","version":3},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-03T01:49:19.070875Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2607.25018"},"observation_digest":"sha256:0d20ee7df45f93217756721bd29a467bc13dae2f5563e4ecda9855b7629edd2c","observation_id":"7631a5a5-0a6e-492a-8c90-ce0746b4dee1","resolution":{"observed_at":"2026-08-03T01:49:19.070875Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-31T02:16:44.828313Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.25068","last_updated":"2026-07-27T20:52:15Z","snapshot_observed_at":"2026-08-02T23:45:34.427712Z","submitted_at":"2026-07-27T20:52:15Z","title":"How Often Should a Recommender Call an LLM? Value-Weighted Routing, Monitoring, and Seasonal Robustness","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-31T02:16:44.828313Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2607.25068"},"observation_digest":"sha256:eb4e5d4ae2046ae1d1ca38fe1a6a2b4647b5254f926342d44f0c7556dae1d57c","observation_id":"a1ab2906-8665-43d2-8239-46f0fdb9c1d1","resolution":{"observed_at":"2026-07-31T02:16:44.828313Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-07-30T18:50:46.197482Z","title":"Large language model cascades with mixture of thought representations for cost-efficient reasoning.arXiv preprint arXiv:2310.03094,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.26865","last_updated":"2026-07-29T12:47:05Z","snapshot_observed_at":"2026-08-02T19:28:19.003719Z","submitted_at":"2026-07-29T12:47:05Z","title":"Think Short, Defer Smart, Act, and Repeat: Calibrated Reasoning and Uncertainty-Aware Deferral for Edge LLM Agents","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-30T18:50:46.197482Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2607.26865"},"observation_digest":"sha256:c40f20e63ed528e9083628fc03152d6166fe98e5506e9e8eb684259e7d2560be","observation_id":"7aa18767-100d-4500-9cd4-eaa65344a713","resolution":{"observed_at":"2026-07-30T18:50:46.197482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03094","snapshot_observed_at":"2026-08-04T07:49:39.987618Z","title":"arXiv preprint arXiv:2310.03094 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.02415","last_updated":"2026-08-03T15:53:49Z","snapshot_observed_at":"2026-08-08T15:56:39.139618Z","submitted_at":"2026-08-03T15:53:49Z","title":"Training-Free versus Training-Based Intent Classification in LLMs: Accuracy, Robustness, and Failure Modes","version":1},"reference_index":103,"source":"arxiv_source","source_observed_at":"2026-08-04T07:49:39.987618Z"},"links":{"cited_paper":"/paper/2310.03094","citing_paper":"/paper/2608.02415"},"observation_digest":"sha256:b86cf31dff9f3db19d6abaad24b249865941e686a2ef2d10d92e439ad0f075c8","observation_id":"c2151bb0-53f0-4594-a0fd-67999d38183f","resolution":{"observed_at":"2026-08-04T07:49:39.987618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2310.03094/citation-record","integrity":"/paper/2310.03094/integrity","json":"/paper/2310.03094/citation-record.json","paper":"/paper/2310.03094"},"outbound":[],"paper":{"arxiv_id":"2310.03094","last_updated":"2024-02-08T22:02:22Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-06T02:17:36.856108Z","submitted_at":"2023-10-04T18:21:17Z","title":"Large Language Model Cascades with Mixture of Thoughts Representations for Cost-efficient Reasoning"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 23 inbound Pith citation observations for arXiv:2310.03094."}