{"as_of":"2026-08-19T05:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:f551d4573f48eea0204e8fed2d0dc4049297c9ffbcf2504c4ef74da06884246d","coverage":[{"denominator":51,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":51,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T22:49:51.844173Z","state":"measured"},{"denominator":60,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":60,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-18T06:34:40.430872+00:00","state":"measured"},{"denominator":9,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":9,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T00:57:00.320706Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.04428","snapshot_observed_at":"2026-08-07T00:57:00.320706Z","title":"Confident or seek stronger: Exploring uncertainty-based on-device Ilm routing from benchmarking to generalization","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.12353","last_updated":"2025-06-14T05:30:09Z","snapshot_observed_at":"2026-08-07T08:59:45.729814Z","submitted_at":"2025-06-14T05:30:09Z","title":"Efficient Reasoning Through Suppression of Self-Affirmation Reflections in Large Reasoning Models","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-07T00:57:00.320706Z"},"links":{"cited_paper":"/paper/2502.04428","citing_paper":"/paper/2506.12353"},"observation_digest":"sha256:cf47e6676b92110e37e3b5209be364476955a4b9b7704de4969e950de881c942","observation_id":"c37b93d1-a14e-4fe0-9281-4a6480356d16","resolution":{"observed_at":"2026-08-07T00:57:00.320706Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.04428","snapshot_observed_at":"2026-08-06T15:06:48.156908Z","title":"Confident or seek stronger: Exploring uncertainty-based on-device llm routing from benchmarking to generalization","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.16731","last_updated":"2025-07-22T16:13:43Z","snapshot_observed_at":"2026-08-16T02:08:34.366113Z","submitted_at":"2025-07-22T16:13:43Z","title":"Collaborative Inference and Learning between Edge SLMs and Cloud LLMs: A Survey of Algorithms, Execution, and Open Challenges","version":1},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-06T15:06:48.156908Z"},"links":{"cited_paper":"/paper/2502.04428","citing_paper":"/paper/2507.16731"},"observation_digest":"sha256:f16c464ff43024dfdf7a4255fa54032bac88a20acf36e0f7302aa3463cd8973e","observation_id":"e16f40c1-1915-40bf-a9e3-ffc0c2c0c6b0","resolution":{"observed_at":"2026-08-06T15:06:48.156908Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"cited_work":{"arxiv_id":"2502.04428","doi":"10.48550/arxiv.2502.04428","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.04428","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Confident or seek stronger: Exploring uncertainty-based on-device llm routing from benchmark- ing to generalization.arXiv preprint arXiv:2502.04428","venue":"ArXiv.org","work_id":"e2408de0-81dd-4ea4-ba03-7f6ed074b853","year":2025},"citing_paper":{"arxiv_id":"2509.24814","last_updated":"2026-05-07T06:47:05Z","snapshot_observed_at":"2026-08-13T11:27:33.387352Z","submitted_at":"2025-09-29T14:02:27Z","title":"A Greedy PDE Router for Blending Neural Operators and Classical Methods","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-18T12:32:13.577474Z"},"links":{"cited_paper":"/paper/2502.04428","citing_paper":"/paper/2509.24814"},"observation_digest":"sha256:25e8df2edd7d18dcbe8a2982fc02015b50705fccf03eca37e19a38a99806b607","observation_id":"cda870a6-3693-44f1-bec5-e9b615073d49","resolution":{"observed_at":"2026-05-18T12:32:36.204473Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.04428","snapshot_observed_at":"2026-08-03T07:14:24.229128Z","title":"Confident or seek stronger: Exploring uncertainty-based on-device llm routing from benchmarking to generalization","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.21003","last_updated":"2026-07-25T01:23:26Z","snapshot_observed_at":"2026-08-16T16:39:52.630416Z","submitted_at":"2026-01-28T19:54:31Z","title":"Bayesian-LoRA: Probabilistic Low-Rank Adaptation of Large Language Models","version":3},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-03T07:14:24.229128Z"},"links":{"cited_paper":"/paper/2502.04428","citing_paper":"/paper/2601.21003"},"observation_digest":"sha256:bd663f6182430080a22d4545699d5526e7f51b73173050421302fe38db77fd93","observation_id":"905ca03c-2dfc-488e-b02c-4fae425b8b9b","resolution":{"observed_at":"2026-08-03T07:14:24.229128Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"cited_work":{"arxiv_id":"2502.04428","doi":"10.48550/arxiv.2502.04428","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.04428","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Confident or seek stronger: Exploring uncertainty-based on-device llm routing from benchmark- ing to generalization.arXiv preprint arXiv:2502.04428","venue":"ArXiv.org","work_id":"e2408de0-81dd-4ea4-ba03-7f6ed074b853","year":2025},"citing_paper":{"arxiv_id":"2604.19781","last_updated":"2026-03-29T20:28:00Z","snapshot_observed_at":"2026-08-11T07:02:17.485746Z","submitted_at":"2026-03-29T20:28:00Z","title":"Do Small Language Models Know When They're Wrong? Confidence-Based Cascade Scoring for Educational Assessment","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-14T21:13:29.201794Z"},"links":{"cited_paper":"/paper/2502.04428","citing_paper":"/paper/2604.19781"},"observation_digest":"sha256:d9632e66dc749ad2c2aa163eea2c225deb547ef4be840a7dd10b5ba9fc1cacec","observation_id":"94b7cb36-fe5c-4ff2-a420-0222f3863743","resolution":{"observed_at":"2026-05-14T21:17:58.436227Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"cited_work":{"arxiv_id":"2502.04428","doi":"10.48550/arxiv.2502.04428","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.04428","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Confident or seek stronger: Exploring uncertainty-based on-device llm routing from benchmark- ing to generalization.arXiv preprint arXiv:2502.04428","venue":"ArXiv.org","work_id":"e2408de0-81dd-4ea4-ba03-7f6ed074b853","year":2025},"citing_paper":{"arxiv_id":"2605.02241","last_updated":"2026-05-07T16:11:01Z","snapshot_observed_at":"2026-08-14T10:16:32.913957Z","submitted_at":"2026-05-04T05:33:03Z","title":"Zero-Shot Confidence Estimation for Small LLMs: When Supervised Baselines Aren't Worth Training","version":3},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-09T16:32:51.205764Z"},"links":{"cited_paper":"/paper/2502.04428","citing_paper":"/paper/2605.02241"},"observation_digest":"sha256:4f9026030ca877884c8702168a88f4a8e44311a33328a5212dcae2d8b33e3914","observation_id":"fafb6711-bb17-4ee8-9906-e6ec68d14b2c","resolution":{"observed_at":"2026-05-11T16:31:07.472454Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"cited_work":{"arxiv_id":"2502.04428","doi":"10.48550/arxiv.2502.04428","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.04428","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Confident or seek stronger: Exploring uncertainty-based on-device llm routing from benchmark- ing to generalization.arXiv preprint arXiv:2502.04428","venue":"ArXiv.org","work_id":"e2408de0-81dd-4ea4-ba03-7f6ed074b853","year":2025},"citing_paper":{"arxiv_id":"2605.06165","last_updated":"2026-05-07T12:51:49Z","snapshot_observed_at":"2026-08-17T00:27:13.360022Z","submitted_at":"2026-05-07T12:51:49Z","title":"Post Reasoning: Improving the Performance of Non-Thinking Models at No Cost","version":1},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-05-08T10:19:08.451445Z"},"links":{"cited_paper":"/paper/2502.04428","citing_paper":"/paper/2605.06165"},"observation_digest":"sha256:6e4028a1fc205438e24b5ffa7fb933e58cfedb58dad2712475769984a0137386","observation_id":"4226c244-b079-491f-81cc-eb3abf4cc3d6","resolution":{"observed_at":"2026-05-11T20:06:09.079876Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"cited_work":{"arxiv_id":"2502.04428","doi":"10.48550/arxiv.2502.04428","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.04428","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Confident or seek stronger: Exploring uncertainty-based on-device llm routing from benchmark- ing to generalization.arXiv preprint arXiv:2502.04428","venue":"ArXiv.org","work_id":"e2408de0-81dd-4ea4-ba03-7f6ed074b853","year":2025},"citing_paper":{"arxiv_id":"2606.30217","last_updated":"2026-06-29T12:30:24Z","snapshot_observed_at":"2026-08-03T23:12:01.383471Z","submitted_at":"2026-06-29T12:30:24Z","title":"Before Thinking, Learn to Decide: Proactive Routing for Efficient Visual Reasoning","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-30T05:55:19.083517Z"},"links":{"cited_paper":"/paper/2502.04428","citing_paper":"/paper/2606.30217"},"observation_digest":"sha256:f248f29c40087490fc9699a6931a03cf167a72616d4ff02da10e7dd60d79814d","observation_id":"644b0a1f-8c50-48fd-9e30-98f0e604dcbd","resolution":{"observed_at":"2026-06-30T13:34:41.164791Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"cited_work":{"arxiv_id":"2502.04428","doi":"10.48550/arxiv.2502.04428","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.04428","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Confident or seek stronger: Exploring uncertainty-based on-device llm routing from benchmark- ing to generalization.arXiv preprint arXiv:2502.04428","venue":"ArXiv.org","work_id":"e2408de0-81dd-4ea4-ba03-7f6ed074b853","year":2025},"citing_paper":{"arxiv_id":"2607.00862","last_updated":"2026-07-01T12:27:14Z","snapshot_observed_at":"2026-08-06T23:19:28.857673Z","submitted_at":"2026-07-01T12:27:14Z","title":"CAT: Confidence-Adaptive Thinking for Efficient Reasoning of Large Reasoning Models","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-07-02T13:22:27.566432Z"},"links":{"cited_paper":"/paper/2502.04428","citing_paper":"/paper/2607.00862"},"observation_digest":"sha256:c6d2a7577671fa53e2da653d82c7e99ade52ecf43e4adf3e61954d9888b02c0c","observation_id":"82b2bf88-ca9d-45db-8019-c8bfce3ac5fd","resolution":{"observed_at":"2026-07-02T13:26:58.122547Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2502.04428/citation-record","integrity":"/paper/2502.04428/integrity","json":"/paper/2502.04428/citation-record.json","paper":"/paper/2502.04428"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.14219","last_updated":"2024-08-30T21:17:17Z","snapshot_observed_at":"2026-08-17T03:25:04.404839Z","submitted_at":"2024-04-22T14:32:33Z","title":"Phi-3 Technical Report: A Highly Capable Language Model Locally on Your Phone","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.14219","snapshot_observed_at":"2026-08-08T22:49:51.685048Z","title":"A., Bach, N., Bahree, A., Bakhtiari, A., Bao, J., Behl, H., et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.685048Z"},"links":{"cited_paper":"/paper/2404.14219","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:842f75e2d004b87f688a3a8e97465663acff119411b2f892bc7343587af59bb6","observation_id":"d2a2bc24-c1a4-4b2c-9210-0a443b2e4473","resolution":{"observed_at":"2026-08-08T22:49:51.685048Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.198433Z","title":"FF: Free-form question answering (including numerical answers for math tasks); MCQ: Multiple-choice question answering; TF: True/False question answering","venue":null,"work_id":"68204779-e240-466b-99b1-f8b1b2fbc86d","year":1954},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.844173Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:e586cbf786794b7f67b7c1a0351e84c47bf4700f4cc6296e1ddd57f06e0af012","observation_id":"159af495-20ff-4067-92e5-a1fda59f0590","resolution":{"observed_at":"2026-08-08T22:49:52.203211Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.336689Z","title":"and Mitchell, T","venue":null,"work_id":"956c85df-bd89-4c76-b0a3-d42e8df4af25","year":2023},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.692457Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:42ff641660ca790f5e6a92bd2e7d83e7785723c80fa2e76fc7489991496db34f","observation_id":"4d0e325b-6477-4720-99fb-84d08ef23c15","resolution":{"observed_at":"2026-08-08T22:49:52.339567Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03766","last_updated":"2024-02-06T07:16:36Z","snapshot_observed_at":"2026-08-09T19:28:34.281681Z","submitted_at":"2024-02-06T07:16:36Z","title":"MobileVLM V2: Faster and Stronger Baseline for Vision Language Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03766","snapshot_observed_at":"2026-08-08T22:49:51.705123Z","title":"Mobilevlm v2: Faster and stronger baseline for vision language model","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.705123Z"},"links":{"cited_paper":"/paper/2402.03766","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:1a34fae5ac078e0f3a9f54655a3b018157cda257ae531d4ce50dd3d3df342819","observation_id":"3eb0df70-9763-4952-b827-2fb58d1ae35f","resolution":{"observed_at":"2026-08-08T22:49:51.705123Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.13284","last_updated":"2025-06-19T22:52:45Z","snapshot_observed_at":"2026-08-16T13:08:32.519541Z","submitted_at":"2024-10-17T07:28:18Z","title":"Learning to Route LLMs with Confidence Tokens","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.13284","snapshot_observed_at":"2026-08-08T22:49:51.708344Z","title":"K., Gopalan, P., Boccio, J., Bolouki, S., and Hu, X","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.708344Z"},"links":{"cited_paper":"/paper/2410.13284","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:4e44f51bee01e1529cc73994a197e34a6be01cf654a64a958846035d9708e650","observation_id":"4e187f49-85f0-4922-98d3-38f5a1187581","resolution":{"observed_at":"2026-08-08T22:49:51.708344Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.328592Z","title":"Boolq: Exploring the surprising difficulty of natural yes/no questions","venue":null,"work_id":"9c5a6a8e-f36e-4631-b071-40477ff91745","year":2019},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.711331Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:e96ac7ee154a3bd9d5f1a9b6d912ba5a73a08d2f6c0aabf6588c4c1852d76474","observation_id":"dc6aa7aa-5e90-4843-88ea-46b0d11438d0","resolution":{"observed_at":"2026-08-08T22:49:52.331472Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.04689","last_updated":"2024-04-06T17:33:37Z","snapshot_observed_at":"2026-08-16T14:03:06.688327Z","submitted_at":"2024-04-06T17:33:37Z","title":"Multicalibration for Confidence Scoring in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.04689","snapshot_observed_at":"2026-08-08T22:49:51.721795Z","title":"Multicalibration for confidence scoring in llms","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.721795Z"},"links":{"cited_paper":"/paper/2404.04689","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:a934820a1fc309b24030c63056f1086f114e68a3728672e5edaf0305f543f9dd","observation_id":"a198a7e6-51a1-4f31-bdbc-0d9414650b8c","resolution":{"observed_at":"2026-08-08T22:49:51.721795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.320028Z","title":"A survey on in- context learning","venue":null,"work_id":"9da9387d-b6a1-43b8-9972-4a5b50bf7851","year":2024},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.725638Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:afa7ce833863ab55b30645fd6fbbcd7c0e0bbabf49cf168275e68f61621b9cbd","observation_id":"cad77890-a070-4ae5-a4e8-d3266f39f776","resolution":{"observed_at":"2026-08-08T22:49:52.323112Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-13T17:20:44.002518Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-08T22:49:51.729047Z","title":"The llama 3 herd of models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.729047Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:d2ec478186208f8c43d195effc3d213de1792da808aefe29172cd90648651448","observation_id":"c2897377-745b-45d4-8103-cd1653a7f600","resolution":{"observed_at":"2026-08-08T22:49:51.729047Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.310723Z","title":"Lm-polygraph: Uncer- tainty estimation for language models","venue":null,"work_id":"fd91b2d2-5884-464e-b04e-048e6232c3b5","year":2023},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.732044Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:016967eb190da950a8b41f5f95c9af7203679f774a39eb5c029e968c48e50d6b","observation_id":"58967096-10cd-4262-88ce-18dce25411fd","resolution":{"observed_at":"2026-08-08T22:49:52.313590Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.12031","last_updated":"2024-03-28T17:56:28Z","snapshot_observed_at":"2026-08-16T16:38:21.205977Z","submitted_at":"2024-03-18T17:59:04Z","title":"RouterBench: A Benchmark for Multi-LLM Routing System","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.12031","snapshot_observed_at":"2026-08-08T22:49:51.737894Z","title":"J., Bieker, J., Li, X., Jiang, N., Keigwin, B., Ran- ganath, G., Keutzer, K., and Upadhyay, S","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.737894Z"},"links":{"cited_paper":"/paper/2403.12031","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:91dcf9f40dcb5921cf35989f1665b03743195e127369c0b8eab1a99052611231","observation_id":"e2892b9b-cbfb-40a7-9cdf-aa8c7740a2d3","resolution":{"observed_at":"2026-08-08T22:49:51.737894Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.10236","last_updated":"2025-01-05T06:15:04Z","snapshot_observed_at":"2026-08-16T15:15:56.313771Z","submitted_at":"2023-07-16T08:28:04Z","title":"Look Before You Leap: An Exploratory Study of Uncertainty Measurement for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.10236","snapshot_observed_at":"2026-08-08T22:49:51.741605Z","title":"Look before you leap: An exploratory study of uncertainty measurement for large language mod- els","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.741605Z"},"links":{"cited_paper":"/paper/2307.10236","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:21237c46237e7b673d8837e09f67ab24a64ca6efb4f5c551a1fa98e2e32ffb18","observation_id":"8c9985b1-0ac7-42d2-844b-922eff5d1786","resolution":{"observed_at":"2026-08-08T22:49:51.741605Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-08-15T14:02:47.366139Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-08-08T22:49:51.745213Z","title":"P., Perelman, A., Ramesh, A., Clark, A., Ostrow, A., Welihinda, A., Hayes, A., Radford, A., et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.745213Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:1ecf3f7fb0a8e2a1e749ef310330cc16b71f9757a70dde3f3d0754fa3db0eeb2","observation_id":"66d44b64-e23e-4be6-b590-342b06263088","resolution":{"observed_at":"2026-08-08T22:49:51.745213Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04088","last_updated":"2024-01-08T18:47:34Z","snapshot_observed_at":"2026-08-13T19:43:49.936776Z","submitted_at":"2024-01-08T18:47:34Z","title":"Mixtral of Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04088","snapshot_observed_at":"2026-08-08T22:49:51.748563Z","title":"Q., Sablayrolles, A., Roux, A., Mensch, A., Savary, B., Bamford, C., Chaplot, D","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.748563Z"},"links":{"cited_paper":"/paper/2401.04088","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:d59d6070958544a5a8fe8a87e9af3f552f2305a8a35f1cf22b27fd4c32ecd121","observation_id":"fbfdd575-6a44-4240-bf50-d2e37c3ef7cb","resolution":{"observed_at":"2026-08-08T22:49:51.748563Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.05221","last_updated":"2022-11-21T16:38:35Z","snapshot_observed_at":"2026-08-14T05:42:29.067319Z","submitted_at":"2022-07-11T22:59:39Z","title":"Language Models (Mostly) Know What They Know","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2207.05221","snapshot_observed_at":"2026-08-08T22:49:51.751528Z","title":"Language models (mostly) know what they know","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.751528Z"},"links":{"cited_paper":"/paper/2207.05221","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:4f740631c2e8688894bf03662a3a1b6a9aead431c19e2c25ce55d298def5463e","observation_id":"d9321a90-20f2-446e-bc06-434f17606ab9","resolution":{"observed_at":"2026-08-08T22:49:51.751528Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.09664","last_updated":"2023-04-15T12:55:45Z","snapshot_observed_at":"2026-07-06T14:53:27.667483Z","submitted_at":"2023-02-19T20:10:07Z","title":"Semantic Uncertainty: Linguistic Invariances for Uncertainty Estimation in Natural Language Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.09664","snapshot_observed_at":"2026-08-08T22:49:51.755030Z","title":"Semantic uncer- tainty: Linguistic invariances for uncertainty estima- tion in natural language generation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.755030Z"},"links":{"cited_paper":"/paper/2302.09664","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:95deb0563e83efb74cecc0ff547b5cd7449db338ae1a3ef7091169e1f397103a","observation_id":"b0e2ab73-70ba-446f-a1be-035a60e79861","resolution":{"observed_at":"2026-08-08T22:49:51.755030Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2109.07958","last_updated":"2022-05-08T02:43:02Z","snapshot_observed_at":"2026-08-03T16:26:48.747700Z","submitted_at":"2021-09-08T17:15:27Z","title":"TruthfulQA: Measuring How Models Mimic Human Falsehoods","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.07958","snapshot_observed_at":"2026-08-08T22:49:51.757706Z","title":"Truthfulqa: Measuring how models mimic human falsehoods","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.757706Z"},"links":{"cited_paper":"/paper/2109.07958","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:40ab00ae3f80fa8949012b48fefa9cb0756d8d90e863d8abe2cbd0387f1b0880","observation_id":"bd64bf2e-8f9f-474f-bb9d-7d4199e673fc","resolution":{"observed_at":"2026-08-08T22:49:51.757706Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.14334","last_updated":"2022-06-13T05:04:53Z","snapshot_observed_at":"2026-08-09T23:45:38.167954Z","submitted_at":"2022-05-28T05:02:31Z","title":"Teaching Models to Express Their Uncertainty in Words","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.14334","snapshot_observed_at":"2026-08-08T22:49:51.760366Z","title":"Teaching models to express their uncertainty in words","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.760366Z"},"links":{"cited_paper":"/paper/2205.14334","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:f72f9c02dbe9fb8cbd2cc90efc05e589929c022bb94c12b58574aaa1d279733d","observation_id":"b1189059-4a03-498c-a570-017b44682832","resolution":{"observed_at":"2026-08-08T22:49:51.760366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.19187","last_updated":"2024-05-20T01:53:36Z","snapshot_observed_at":"2026-08-18T09:28:29.473363Z","submitted_at":"2023-05-30T16:31:26Z","title":"Generating with Confidence: Uncertainty Quantification for Black-box Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.19187","snapshot_observed_at":"2026-08-08T22:49:51.763242Z","title":"Generating with confidence: Uncertainty quantification for black-box large language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.763242Z"},"links":{"cited_paper":"/paper/2305.19187","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:cbd043500d6849ab41ecf5729d40a92612bbf47ca7cf0ba31d3c4034d07bf7e0","observation_id":"b253e4b8-bf55-428d-96bf-d7e265570232","resolution":{"observed_at":"2026-08-08T22:49:51.763242Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.13415","last_updated":"2024-06-19T10:11:37Z","snapshot_observed_at":"2026-08-16T13:41:30.954804Z","submitted_at":"2024-06-19T10:11:37Z","title":"Factual Confidence of LLMs: on Reliability and Robustness of Current Estimators","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.13415","snapshot_observed_at":"2026-08-08T22:49:51.769405Z","title":"Factual confidence of llms: on reliability and robustness of current estimators","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.769405Z"},"links":{"cited_paper":"/paper/2406.13415","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:317e83abd62aeb7ea05a1db78e1246ec1a15ad632e4793711ad0582a27280fcb","observation_id":"3f90f238-7bbe-463b-9304-c76eb7148589","resolution":{"observed_at":"2026-08-08T22:49:51.769405Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.291822Z","title":"Selfcheckgpt: Zero- resource black-box hallucination detection for genera- tive large language models","venue":null,"work_id":"fa068e2d-0daa-4e53-87e5-e509175504dd","year":2023},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.772217Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:ecbe0c482ea591cc9394286804ee10c348c849574677111557c895d7529218e1","observation_id":"c318a8ca-cc00-4450-b448-de7e3cfae4ce","resolution":{"observed_at":"2026-08-08T22:49:52.295701Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.280494Z","title":"Llama 3.2: Revolutionizing edge ai and vision with open, customizable models","venue":null,"work_id":"40c3f947-5e5d-4b19-b3ed-61566b5245eb","year":2024},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.775183Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:3bc91e64d629b1ac094d7ebcfa9a92ea19b69a354f84825b4cfad9b36da8eaa7","observation_id":"bdcbd77c-fc68-4b4e-a573-adfa71abcc53","resolution":{"observed_at":"2026-08-08T22:49:52.284402Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1809.02789","last_updated":"2018-09-08T11:47:16Z","snapshot_observed_at":"2026-08-14T07:00:02.529732Z","submitted_at":"2018-09-08T11:47:16Z","title":"Can a Suit of Armor Conduct Electricity? A New Dataset for Open Book Question Answering","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1809.02789","snapshot_observed_at":"2026-08-08T22:49:51.777778Z","title":"Can a suit of armor conduct electricity? a new dataset for open book question answering","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.777778Z"},"links":{"cited_paper":"/paper/1809.02789","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:522c4433fcbee80ceb3d2fb6cfd0d26670abba107ea83dda61c57dbc92df39ac","observation_id":"e00f3d42-fc83-4f39-b1ab-76de36c74747","resolution":{"observed_at":"2026-08-08T22:49:51.777778Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.02060","last_updated":"2025-03-03T01:25:46Z","snapshot_observed_at":"2026-08-15T00:07:51.079800Z","submitted_at":"2024-09-03T17:08:20Z","title":"OLMoE: Open Mixture-of-Experts Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.02060","snapshot_observed_at":"2026-08-08T22:49:51.780442Z","title":"Ong, I., Almahairi, A., Wu, V ., Chiang, W.-L., Wu, T., Gon- zalez, J","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.780442Z"},"links":{"cited_paper":"/paper/2409.02060","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:5aff247596709c37424f0db558e53a4703d91f0d5090da26b755c08fa7de3f56","observation_id":"37b624ca-6a68-4237-8132-cfcafa82d03f","resolution":{"observed_at":"2026-08-08T22:49:51.780442Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.269015Z","title":null,"venue":null,"work_id":"2f82d375-13a2-4fbf-8dbb-1b0d7c2311fe","year":2021},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.783139Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:27b0d5d207c1bfb89621902d2f52396e890918cda588473bce275dd965c3613e","observation_id":"df72b521-bb48-4f07-8299-4d723edcfa79","resolution":{"observed_at":"2026-08-08T22:49:52.272584Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13048","last_updated":"2023-12-11T03:58:56Z","snapshot_observed_at":"2026-08-17T17:58:38.402665Z","submitted_at":"2023-05-22T13:57:41Z","title":"RWKV: Reinventing RNNs for the Transformer Era","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13048","snapshot_observed_at":"2026-08-08T22:49:51.785535Z","title":"Rwkv: Reinventing rnns for the transformer era","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.785535Z"},"links":{"cited_paper":"/paper/2305.13048","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:0ea5376a5cc57bb808d966ee5414835cfb09c4fe7dd63883ec4c97ab6a7c0af0","observation_id":"1e0da4b9-a75f-4a6d-b8f9-6c4467fb7e6f","resolution":{"observed_at":"2026-08-08T22:49:51.785535Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.09276","last_updated":"2024-07-12T14:09:40Z","snapshot_observed_at":"2026-08-16T13:34:45.121495Z","submitted_at":"2024-07-12T14:09:40Z","title":"H2O-Danube3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.09276","snapshot_observed_at":"2026-08-08T22:49:51.788735Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.788735Z"},"links":{"cited_paper":"/paper/2407.09276","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:4b0c128d7e527320cab3f119a627ed9f0558fe0b1a0ef6504cb61319de5ab9cb","observation_id":"471afd14-88eb-4719-b4ba-73a98c730532","resolution":{"observed_at":"2026-08-08T22:49:51.788735Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1608.01413","last_updated":"2016-08-20T11:50:41Z","snapshot_observed_at":"2026-08-18T09:21:00.040496Z","submitted_at":"2016-08-04T01:47:23Z","title":"Solving General Arithmetic Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1608.01413","snapshot_observed_at":"2026-08-08T22:49:51.791988Z","title":"and Roth, D","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.791988Z"},"links":{"cited_paper":"/paper/1608.01413","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:041e9cc48af78b85db34547cbe4c8fc1ca93ec8ec6611b6d878ca98c52e7b995","observation_id":"d482f99d-af9e-486a-bf39-18db2a1daa15","resolution":{"observed_at":"2026-08-08T22:49:51.791988Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.257626Z","title":"Social iqa: Commonsense reasoning about social interactions","venue":null,"work_id":"5982d6de-9ecc-4a00-b1a0-fc8700f5357e","year":2019},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.795164Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:bbf76367ee463770794d555ecfb0f5085dcd651f6517e7508e47e5c45e78e268","observation_id":"01ceda28-21f7-4eaf-985e-c1443065f06b","resolution":{"observed_at":"2026-08-08T22:49:52.261153Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.12320","last_updated":"2024-10-23T18:11:42Z","snapshot_observed_at":"2026-08-17T20:35:06.121552Z","submitted_at":"2024-08-22T11:57:07Z","title":"TensorOpera Router: A Multi-Model Router for Efficient LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.12320","snapshot_observed_at":"2026-08-08T22:49:51.798116Z","title":"Polyrouter: A multi- llm querying system","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.798116Z"},"links":{"cited_paper":"/paper/2408.12320","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:044aaf4ac279a366427ba40c4b07228743becbf4f68888424915eadd824d6046","observation_id":"6bd6de6c-f446-4a20-876a-8e2dfb717e85","resolution":{"observed_at":"2026-08-08T22:49:51.798116Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.245490Z","title":"Com- monsenseQA: A question answering challenge targeting commonsense knowledge","venue":null,"work_id":"2405c7cf-4668-425f-8ded-29414aeab127","year":2019},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.801572Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:42439e11c072405bfb761b1c77eab91049984aa5c0c346c0228c527376cda0c3","observation_id":"0a3ad4a9-1538-454c-b9a3-99d3dd588fc2","resolution":{"observed_at":"2026-08-08T22:49:52.249179Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.233192Z","title":null,"venue":null,"work_id":"33fd10f0-6e37-47a9-94e7-3c626d9ceb4d","year":2023},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.805079Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:84a5ad081e6a3ed0a7a72612281b9d8119e9b74722237630536abd0202ec21f1","observation_id":"b98b9a32-93a5-4424-9d08-3263171c10c5","resolution":{"observed_at":"2026-08-08T22:49:52.237635Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-08T22:49:51.808067Z","title":"Llama: Open and efficient foundation lan- guage models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.808067Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:b25a7d41211675d0f3a27c6c5938156afa9a8deda9a930420d9ca1d5653e457d","observation_id":"c23fc0c9-9d5d-4eba-9085-257a2cf1b101","resolution":{"observed_at":"2026-08-08T22:49:51.808067Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.221163Z","title":"Efficient out-of-domain detection for sequence to sequence models","venue":null,"work_id":"58871e25-d09a-45db-a46b-b2868afffbf7","year":2023},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.811338Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:9c3a56ed2b5448c96e45167dc2ef3c0d92a4e0f207ff05ec916ac3f780ce07ff","observation_id":"f15208aa-59f6-4af0-a11b-effed09ec339","resolution":{"observed_at":"2026-08-08T22:49:52.225335Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.211007Z","title":"P., Delucia, A., and Dredze, M","venue":null,"work_id":"0b64d78b-b414-40fe-bd01-3529ae1e8d60","year":2023},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.815312Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:1665a88b916f99392591e4f8aa1bdc1b2ccc5048d1b21e9953686e7ba8313c7a","observation_id":"617f0116-9939-4e88-89f1-d0a55cad236b","resolution":{"observed_at":"2026-08-08T22:49:52.214900Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.13063","last_updated":"2024-03-17T04:38:48Z","snapshot_observed_at":"2026-08-15T10:58:51.256503Z","submitted_at":"2023-06-22T17:31:44Z","title":"Can LLMs Express Their Uncertainty? An Empirical Evaluation of Confidence Elicitation in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.13063","snapshot_observed_at":"2026-08-08T22:49:51.818627Z","title":"Can llms express their uncertainty? an empirical evaluation of confidence elicitation in llms.arXiv preprint arXiv:2306.13063,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.818627Z"},"links":{"cited_paper":"/paper/2306.13063","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:5a3f1ebbb0ef47263138471cdb9c6ccde9a71d327af7e891d1eede476a250d4d","observation_id":"53cf45e0-4a02-4e0b-b466-579d79a89f42","resolution":{"observed_at":"2026-08-08T22:49:51.818627Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10671","last_updated":"2024-09-10T13:25:53Z","snapshot_observed_at":"2026-08-17T11:08:48.802438Z","submitted_at":"2024-07-15T12:35:42Z","title":"Qwen2 Technical Report","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10671","snapshot_observed_at":"2026-08-08T22:49:51.821773Z","title":"Qwen2 technical report","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.821773Z"},"links":{"cited_paper":"/paper/2407.10671","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:1dc99d1c8add40b76344e925f3012f8be65ff865ee2d2c2513f15763706bdc49","observation_id":"c56e0578-04c2-460e-99dd-9f52340e97b7","resolution":{"observed_at":"2026-08-08T22:49:51.821773Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16908","last_updated":"2024-09-26T08:53:01Z","snapshot_observed_at":"2026-08-18T23:48:41.069839Z","submitted_at":"2024-05-27T07:56:23Z","title":"Can Large Language Models Faithfully Express Their Intrinsic Uncertainty in Words?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16908","snapshot_observed_at":"2026-08-08T22:49:51.825552Z","title":"Can large language models faithfully express their intrinsic uncertainty in words? arXiv preprint arXiv:2405.16908,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.825552Z"},"links":{"cited_paper":"/paper/2405.16908","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:32425d174eb3b2cc05f3a517616f9e08dfa8a2752d441b1e9643dc06fb40be39","observation_id":"cb0d7eee-d9d4-4b3c-9117-27a1c2e6778f","resolution":{"observed_at":"2026-08-08T22:49:51.825552Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.02385","last_updated":"2024-06-04T02:05:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-04T17:54:59Z","title":"TinyLlama: An Open-Source Small Language Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.02385","snapshot_observed_at":"2026-08-08T22:49:51.829271Z","title":"Tinyllama: An open-source small language model","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.829271Z"},"links":{"cited_paper":"/paper/2401.02385","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:8cf2c8c2e4667aa842fe53baeecf0c932ed8c342fee801e21ddbb5c8d80a4fd3","observation_id":"b4ed9843-160a-41b4-9e8e-37e63a187db6","resolution":{"observed_at":"2026-08-08T22:49:51.829271Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.01068","last_updated":"2022-06-21T17:04:40Z","snapshot_observed_at":"2026-08-06T03:13:37.403059Z","submitted_at":"2022-05-02T17:49:50Z","title":"OPT: Open Pre-trained Transformer Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.01068","snapshot_observed_at":"2026-08-08T22:49:51.832594Z","title":"V ., et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.832594Z"},"links":{"cited_paper":"/paper/2205.01068","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:c7d00c2c4779b0b262cf10b453b1463ce6144d45df7e0cf7864637f59d756671","observation_id":"c62a21fb-2f1b-4edb-875b-eedfaaa0e91d","resolution":{"observed_at":"2026-08-08T22:49:51.832594Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.18223","last_updated":"2026-03-18T05:34:39Z","snapshot_observed_at":"2026-08-14T10:40:26.323157Z","submitted_at":"2023-03-31T17:28:46Z","title":"A Survey of Large Language Models","version":19},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.18223","snapshot_observed_at":"2026-08-08T22:49:51.835823Z","title":"X., Zhou, K., Li, J., Tang, T., Wang, X., Hou, Y ., Min, Y ., Zhang, B., Zhang, J., Dong, Z., et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.835823Z"},"links":{"cited_paper":"/paper/2303.18223","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:20d7e4bfa36765ad02237b59f13b6d32b5dd32e4cbf7b3a0d8d88c6d5d1130bd","observation_id":"180ffce6-3923-48b9-9311-9142254dacad","resolution":{"observed_at":"2026-08-08T22:49:51.835823Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.15518","last_updated":"2024-10-29T00:20:45Z","snapshot_observed_at":"2026-08-16T13:16:06.274919Z","submitted_at":"2024-09-23T20:10:10Z","title":"Eagle: Efficient Training-Free Router for Multi-LLM Inference","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.15518","snapshot_observed_at":"2026-08-08T22:49:51.838616Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.838616Z"},"links":{"cited_paper":"/paper/2409.15518","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:3bf4dadd77e87eab483cc3178219b63f38f55373a846c877976b30deec025970","observation_id":"bbd77d84-474a-4913-8b03-b6b035e119d7","resolution":{"observed_at":"2026-08-08T22:49:51.838616Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.08189","last_updated":"2024-07-07T06:42:31Z","snapshot_observed_at":"2026-08-16T23:01:59.512240Z","submitted_at":"2023-07-17T01:35:56Z","title":"Mini-Giants: \"Small\" Language Models and Open Source Win-Win","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.08189","snapshot_observed_at":"2026-08-08T22:49:51.841392Z","title":"Mini-giants:\" small\" language models and open source win-win.arXiv preprint arXiv:2307.08189,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.841392Z"},"links":{"cited_paper":"/paper/2307.08189","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:a6135e0d7f7bf7c1f4d22e69a37582cb119c171674f280ee4bc049eabced11dc","observation_id":"be1e8f29-6557-41b4-90c0-4bf57f1e86cb","resolution":{"observed_at":"2026-08-08T22:49:51.841392Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-08T22:49:52.301808Z","title":"A survey of confidence estimation and calibration in large language models","venue":null,"work_id":"57d7164e-fddf-40d0-ae68-c9fae6462512","year":2024},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":2016,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.734943Z"},"links":{"citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:fb9a30a11e51d44428e0d84d32e223814fb95d4af95506886e4ce910e1022cdf","observation_id":"9bc13ef9-9f95-487c-934c-7b5a71f985d4","resolution":{"observed_at":"2026-08-08T22:49:52.305076Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-18T06:34:40.430872+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.15993","last_updated":"2024-10-23T08:33:54Z","snapshot_observed_at":"2026-08-18T07:54:59.904276Z","submitted_at":"2024-04-24T17:10:35Z","title":"Uncertainty Estimation and Quantification for LLMs: A Simple Supervised Approach","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.15993","snapshot_observed_at":"2026-08-08T22:49:51.766316Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.766316Z"},"links":{"cited_paper":"/paper/2404.15993","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:4f37be0bb85020c805df0a56aee429a6f4435290a4c2f18d4d55fd98b00388d2","observation_id":"5b821494-1088-44df-8176-a17f2a9289fd","resolution":{"observed_at":"2026-08-08T22:49:51.766316Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-08-14T02:43:01.480086Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-08-08T22:49:51.715003Z","title":"Training verifiers to solve math word problems","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.715003Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:66ab1e148cd072927e847e607a00c10aa4adc07d25fc400d174e64da8981b6b2","observation_id":"c3aa7bb1-8702-49cd-81fe-a2bd8370c28b","resolution":{"observed_at":"2026-08-08T22:49:51.715003Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.03827","last_updated":"2024-03-02T21:33:53Z","snapshot_observed_at":"2026-08-17T12:54:32.587457Z","submitted_at":"2022-12-07T18:17:56Z","title":"Discovering Latent Knowledge in Language Models Without Supervision","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.03827","snapshot_observed_at":"2026-08-08T22:49:51.698774Z","title":"Discovering latent knowledge in language models without supervision","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.698774Z"},"links":{"cited_paper":"/paper/2212.03827","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:16ddaac8fec9ec7fe9ba0c19ca1a61d4bf4fd7688864d044ebe0613048d1c36d","observation_id":"92e0b825-6ad6-4248-8d42-71a72784f47f","resolution":{"observed_at":"2026-08-08T22:49:51.698774Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.21060","last_updated":"2024-05-31T17:50:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-05-31T17:50:01Z","title":"Transformers are SSMs: Generalized Models and Efficient Algorithms Through Structured State Space Duality","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.21060","snapshot_observed_at":"2026-08-08T22:49:51.718501Z","title":"and Gu, A","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.718501Z"},"links":{"cited_paper":"/paper/2405.21060","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:3f029b866c314be4bdcc004818c502199abb481372fa4eba6f4f7af9ac1cce95","observation_id":"f942d742-5198-4b63-a132-003de8d12982","resolution":{"observed_at":"2026-08-08T22:49:51.718501Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.06857","last_updated":"2026-08-03T10:27:28Z","snapshot_observed_at":"2026-08-18T02:45:08.653847Z","submitted_at":"2024-09-10T20:45:43Z","title":"What is the Role of Small Models in the LLM Era: A Survey","version":8},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.06857","snapshot_observed_at":"2026-08-08T22:49:51.701759Z","title":"and Varoquaux, G","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":2022,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.701759Z"},"links":{"cited_paper":"/paper/2409.06857","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:e3db9fc89ac0cc0168a2351c0edebf1c92a97ba79eb1cff0a2ea74793fadd906","observation_id":"c44fd253-9c0f-434d-a3a7-eb485b754587","resolution":{"observed_at":"2026-08-08T22:49:51.701759Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.16609","last_updated":"2023-09-28T17:07:49Z","snapshot_observed_at":"2026-08-09T21:25:20.369782Z","submitted_at":"2023-09-28T17:07:49Z","title":"Qwen Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.16609","snapshot_observed_at":"2026-08-08T22:49:51.695700Z","title":"Qwen technical report","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.695700Z"},"links":{"cited_paper":"/paper/2309.16609","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:9fb6dfd418db9731957f8197eb16f91e4cd54f7693be947ee9aa5374847885f5","observation_id":"f4473d77-4f6f-4d7b-9fe7-3353c1be7c2b","resolution":{"observed_at":"2026-08-08T22:49:51.695700Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-17T09:58:46.058102Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-08T22:49:51.689459Z","title":"L., Almeida, D., Altenschmidt, J., Altman, S., Anadkat, S., et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-08T22:49:51.689459Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2502.04428"},"observation_digest":"sha256:0255546423fcdeae662d9db32885895c6695b644974b01a58ebe5dccaa46727b","observation_id":"82580ddc-e3e9-409d-848b-517a52985a89","resolution":{"observed_at":"2026-08-08T22:49:51.689459Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2502.04428","last_updated":"2025-02-06T18:59:11Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-18T08:36:55.882093Z","submitted_at":"2025-02-06T18:59:11Z","title":"Confident or Seek Stronger: Exploring Uncertainty-Based On-device LLM Routing From Benchmarking to Generalization"},"reference_resolution":{"displayed":51,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":39,"verified_exact":0,"verified_fuzzy":12},"total_outbound_references":51},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-18T06:34:40.430872+00:00","source":"crossref"},{"observed_at":"2026-08-18T06:34:34.496301+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 51 of 51 outbound references and 9 inbound Pith citation observations for arXiv:2502.04428."}