{"as_of":"2026-08-12T10:55:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:9657d02574a41f89a7ac54639bb6b9e4006f7a4815eef4f5044b20f060f5d044","coverage":[{"denominator":38,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":38,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T21:45:55.440332Z","state":"measured"},{"denominator":40,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":40,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-12T03:10:22.314719Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"cited_work":{"arxiv_id":"2501.03991","doi":"10.48550/arxiv.2501.03991","metadata_source":"arxiv_reference","pith_arxiv_id":"2501.03991","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Influences on","venue":"arXiv (Cornell University)","work_id":"f34cef3f-b293-4923-b684-685569772108","year":null},"citing_paper":{"arxiv_id":"2605.00195","last_updated":"2026-05-11T12:48:16Z","snapshot_observed_at":"2026-07-31T18:51:54.041588Z","submitted_at":"2026-04-30T20:20:59Z","title":"Diversity in Large Language Models under Supervised Fine-Tuning","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-05-09T20:32:37.788283Z"},"links":{"cited_paper":"/paper/2501.03991","citing_paper":"/paper/2605.00195"},"observation_digest":"sha256:3c5ad66a023dc9ce901b2a3c2573e14c4a02d6b71777b99f3147d5268a86c1a9","observation_id":"e23a4143-c4fd-4ba1-9487-e5d342c18bc6","resolution":{"observed_at":"2026-05-09T20:37:32.184304Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"cited_work":{"arxiv_id":"2501.03991","doi":"10.48550/arxiv.2501.03991","metadata_source":"arxiv_reference","pith_arxiv_id":"2501.03991","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Influences on","venue":"arXiv (Cornell University)","work_id":"f34cef3f-b293-4923-b684-685569772108","year":null},"citing_paper":{"arxiv_id":"2605.00195","last_updated":"2026-05-11T12:48:16Z","snapshot_observed_at":"2026-07-31T18:51:54.041588Z","submitted_at":"2026-04-30T20:20:59Z","title":"Diversity in Large Language Models under Supervised Fine-Tuning","version":2},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-05-12T03:10:22.314719Z"},"links":{"cited_paper":"/paper/2501.03991","citing_paper":"/paper/2605.00195"},"observation_digest":"sha256:3ccbcc1bdc7b2dd0be2353e9ce3b920242d50ad3bb131dd7e807134e04e0bef5","observation_id":"1c83a074-fcde-4da3-b10c-d6ad65543441","resolution":{"observed_at":"2026-05-12T03:11:18.163273Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2501.03991/citation-record","integrity":"/paper/2501.03991/integrity","json":"/paper/2501.03991/citation-record.json","paper":"/paper/2501.03991"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2404.14219","last_updated":"2024-08-30T21:17:17Z","snapshot_observed_at":"2026-08-10T14:07:02.234322Z","submitted_at":"2024-04-22T14:32:33Z","title":"Phi-3 Technical Report: A Highly Capable Language Model Locally on Your Phone","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.14219","snapshot_observed_at":"2026-08-10T21:45:55.223260Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.223260Z"},"links":{"cited_paper":"/paper/2404.14219","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:3c9703ee005c6819e6d48aa681e2cdb7745e89759457df03e3d6ebac17f7a18b","observation_id":"2437d2a3-1f25-4f30-8035-dab3d49d55f2","resolution":{"observed_at":"2026-08-10T21:45:55.223260Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.231174Z","title":null,"venue":null,"work_id":null,"year":1950},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.231174Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:1b3503b619fbfb2330d7bd58985d8b46561ddddcb3e0a8084e2bdf3c5b73c6c4","observation_id":"d4accea8-d40d-4512-a807-6766d7f8e59a","resolution":{"observed_at":"2026-08-10T21:45:55.231174Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.14735","last_updated":"2025-05-11T09:23:41Z","snapshot_observed_at":"2026-08-10T12:40:24.930519Z","submitted_at":"2023-10-23T09:15:18Z","title":"Unleashing the potential of prompt engineering for large language models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.14735","snapshot_observed_at":"2026-08-10T21:45:55.237824Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.237824Z"},"links":{"cited_paper":"/paper/2310.14735","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:65dbb56bdf102d56e84341078d9e831062594097ff5563996835daf2bff251a7","observation_id":"6fe7c2af-1b1f-4142-8d55-05f0dbd5efbf","resolution":{"observed_at":"2026-08-10T21:45:55.237824Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1810.04805","last_updated":"2019-05-24T20:37:26Z","snapshot_observed_at":"2026-07-30T09:12:38.100527Z","submitted_at":"2018-10-11T00:50:01Z","title":"BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.04805","snapshot_observed_at":"2026-08-10T21:45:55.246319Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.246319Z"},"links":{"cited_paper":"/paper/1810.04805","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:2d4615717d67edaa7d3dac1ae9b5d3229dfa3cf0e023de12d1ab606ed5349559","observation_id":"9655f459-5a0d-4bc2-8e71-fcf3c3814996","resolution":{"observed_at":"2026-08-10T21:45:55.246319Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:56.265689Z","title":null,"venue":null,"work_id":"d3f015cd-5acf-4b98-9568-eae3eeb9a0bf","year":2017},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.254079Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:b6fea26b75b15631813777e5046d849c14ff4404eebabc1810d4008bde54df16","observation_id":"977b23d0-960d-4453-abd3-0df00bb25af4","resolution":{"observed_at":"2026-08-10T21:45:56.271323Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.260066Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.260066Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:0135275c86ec6110064c8dbd2b5b2dcf63a96e006f83e0bf8247f4c9d83a88eb","observation_id":"18b1d759-3c19-4c63-9ea9-c89897dd98bf","resolution":{"observed_at":"2026-08-10T21:45:55.260066Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-10T16:40:37.411115Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-10T21:45:55.266290Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.266290Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:336e5b88a77ee0cb7d008b7c5f8736076bdb63192dd1f07610a2d85f1a7f4160","observation_id":"363c05d4-b30c-44fe-90be-f9457651271b","resolution":{"observed_at":"2026-08-10T21:45:55.266290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:56.224040Z","title":"Weinberger","venue":null,"work_id":"214afc15-3820-43c8-ae1b-aa2dec82adf8","year":2017},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.273080Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:e48bddbd53de4f9c46fd18cac5e8ceaf2a722dad40957a292df8fa714ce80e32","observation_id":"93975104-9ac5-4b2b-af92-4c63caedd451","resolution":{"observed_at":"2026-08-10T21:45:56.229309Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.279578Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.279578Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:88ba06a46d0eb9dab016c67ffdc2b4d4dc1b9b0200fef5f75f5d4ae7ff19a52d","observation_id":"da6b8bd1-9adf-44ba-9e19-0a974c85cf97","resolution":{"observed_at":"2026-08-10T21:45:55.279578Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04088","last_updated":"2024-01-08T18:47:34Z","snapshot_observed_at":"2026-08-08T06:16:25.839566Z","submitted_at":"2024-01-08T18:47:34Z","title":"Mixtral of Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04088","snapshot_observed_at":"2026-08-10T21:45:55.285607Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.285607Z"},"links":{"cited_paper":"/paper/2401.04088","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:8b93e601befff8e00c37db519317d5cef08ca515902c98c3cbe068bae3148795","observation_id":"82d3a353-8b65-4b43-8054-ee8a490139cc","resolution":{"observed_at":"2026-08-10T21:45:55.285607Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.291623Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.291623Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:3b6094756ba6df8fc04ecb2f847e332b653343cfba3dee56a7313ac828c77405","observation_id":"b4be6e2f-c567-4069-a37c-74949b1e0b61","resolution":{"observed_at":"2026-08-10T21:45:55.291623Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06209","last_updated":"2017-07-19T17:28:46Z","snapshot_observed_at":"2026-08-06T06:05:27.671303Z","submitted_at":"2017-07-19T17:28:46Z","title":"Crowdsourcing Multiple Choice Science Questions","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06209","snapshot_observed_at":"2026-08-10T21:45:55.296875Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.296875Z"},"links":{"cited_paper":"/paper/1707.06209","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:c17f3945a2a66805580752c1697f298b487cb5f15956336a3df5f35c5c2402b2","observation_id":"b6fd4399-7601-495f-87fb-6ccacd27f204","resolution":{"observed_at":"2026-08-10T21:45:55.296875Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.302037Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.302037Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:bfbf84fbfe6e41db83e1f9d76a232ba21aeed1bf49cd512f53ba98055f81c48f","observation_id":"134f8437-793e-4540-a239-9cdf0ac97f70","resolution":{"observed_at":"2026-08-10T21:45:55.302037Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2023.findings-eacl.40","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.554793Z","title":null,"venue":null,"work_id":"124063aa-c2fe-42b1-9e9d-92a5a50c418d","year":2023},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.307831Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:51d80ac17d7b58d372d72fce1b7e5aaa4e887d2bc2ece3bf705e21c3c69cbed0","observation_id":"ebaccc24-93fa-4d61-99c7-0d4886b80ac2","resolution":{"observed_at":"2026-08-10T21:45:55.564236Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.01535","last_updated":"2024-12-04T19:23:17Z","snapshot_observed_at":"2026-08-06T22:44:08.607902Z","submitted_at":"2024-05-02T17:59:35Z","title":"Prometheus 2: An Open Source Language Model Specialized in Evaluating Other Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.01535","snapshot_observed_at":"2026-08-10T21:45:55.313121Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.313121Z"},"links":{"cited_paper":"/paper/2405.01535","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:d049b6679349a5e85590207150a580ba55fc45dd46955f9d455b116d8f6e2790","observation_id":"c065bc01-aca6-499e-9fdd-dfc70ec4dbf6","resolution":{"observed_at":"2026-08-10T21:45:55.313121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.318696Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.318696Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:b7ab9a88558c8f2a01939bb70a3e6efae56d5541d23f3d2587ae44979d70ec71","observation_id":"e7fd3ae0-bc93-4cdb-ae68-e9540c7f249f","resolution":{"observed_at":"2026-08-10T21:45:55.318696Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.324814Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.324814Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:f172a4d855eabf1bce98cadd19b234fcaea62ae500b12f4a1113c44981f45644","observation_id":"3dd76152-0712-44d3-ad7d-3194c91e4685","resolution":{"observed_at":"2026-08-10T21:45:55.324814Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.330193Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.330193Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:e2394415c874639caa0e3457dade21dd8e0bb675ad4c65fccd1911ff9f9367ce","observation_id":"28d3090a-ccab-4bab-853d-278f878b076f","resolution":{"observed_at":"2026-08-10T21:45:55.330193Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.19208","last_updated":"2024-03-13T05:11:57Z","snapshot_observed_at":"2026-08-09T12:55:13.317965Z","submitted_at":"2023-10-30T00:30:34Z","title":"LitCab: Lightweight Language Model Calibration over Short- and Long-form Responses","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.19208","snapshot_observed_at":"2026-08-10T21:45:55.335235Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.335235Z"},"links":{"cited_paper":"/paper/2310.19208","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:c0b55aff01f4f5dea619a5c1885da1db9809713a77c4928578f1e9d2e8d09daa","observation_id":"9ef77ced-f288-44a3-80b1-730665737ce7","resolution":{"observed_at":"2026-08-10T21:45:55.335235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.340506Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.340506Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:8562e05492fd1c32f4a377d92305b03c5b8721ce31f5869060c154c1798558b1","observation_id":"372df8b3-dd6e-4c24-ad8d-dcd4bcf7af54","resolution":{"observed_at":"2026-08-10T21:45:55.340506Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:56.125646Z","title":null,"venue":null,"work_id":"901ec109-d981-40b1-b619-12751bf0f2e1","year":2020},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.345545Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:349aac9f4793aec4df87f56f80f6ff6883d4752c13a55504576b00d9f98e7520","observation_id":"07db8076-e51c-4e73-8c41-08e2a516f43d","resolution":{"observed_at":"2026-08-10T21:45:56.130568Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.349813Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.349813Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:e435019361067ccbb209561eb224982dce72d1da850ea4fbcf0bde23100cdd1e","observation_id":"dad6e6fa-fff3-4fb9-859d-998ddbae3248","resolution":{"observed_at":"2026-08-10T21:45:55.349813Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.09773","last_updated":"2024-08-19T08:01:11Z","snapshot_observed_at":"2026-08-07T02:37:14.627726Z","submitted_at":"2024-08-19T08:01:11Z","title":"Are Large Language Models More Honest in Their Probabilistic or Verbalized Confidence?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.09773","snapshot_observed_at":"2026-08-10T21:45:55.353869Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.353869Z"},"links":{"cited_paper":"/paper/2408.09773","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:e97018a43d72c4263d918946279f64ab1be56a2a30c91e9a6ad1964ad752b914","observation_id":"75e18df9-c47d-42a1-bab0-bab40123489a","resolution":{"observed_at":"2026-08-10T21:45:55.353869Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:56.099113Z","title":null,"venue":null,"work_id":"f37e37da-ccc3-49e3-b134-19f145c7f4b4","year":1999},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.358561Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:649de94501e14239e5359c7a66ba0320cac82032db82da0f22b0f21c739e8d6d","observation_id":"80b879ff-c89a-44b4-b2e4-181cbc11883e","resolution":{"observed_at":"2026-08-10T21:45:56.104511Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:56.084134Z","title":null,"venue":null,"work_id":"57c5c24a-b1d9-4a12-85f9-eb440488a254","year":2020},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.362780Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:15ff26add7e2734d038f232df8e96b09a65772aff8fdd2cd71eb0f9bd824c217","observation_id":"861caf24-c6bd-4b84-9189-ba350bda04df","resolution":{"observed_at":"2026-08-10T21:45:56.088795Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00118","last_updated":"2024-10-02T15:22:49Z","snapshot_observed_at":"2026-08-02T16:20:09.773989Z","submitted_at":"2024-07-31T19:13:07Z","title":"Gemma 2: Improving Open Language Models at a Practical Size","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.00118","snapshot_observed_at":"2026-08-10T21:45:55.367568Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.367568Z"},"links":{"cited_paper":"/paper/2408.00118","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:5f88a0e7aea8d81d028645cfc9a7b1e3382c35b5f96a6464a119b23e75d00a4c","observation_id":"70da037e-6b04-43f2-8887-60240f1d9332","resolution":{"observed_at":"2026-08-10T21:45:55.367568Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.372940Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.372940Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:ef6aa56b142152e4bbaefd5bacad3bb4e5dffc7be63e74b943249c2ac0f2f8a5","observation_id":"69d81529-fee0-4dc7-acd6-c91bafa35aa4","resolution":{"observed_at":"2026-08-10T21:45:55.372940Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-07T12:56:43.323460Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-10T21:45:55.378290Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.378290Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:b0217cff5fe5250d6001bd63061b39e5071c273023ebad8174a24f668b64811a","observation_id":"366aef7f-7944-4f1b-b868-6b5d828a89e9","resolution":{"observed_at":"2026-08-10T21:45:55.378290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.387807Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.387807Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:9860b76d8f3a35d0f85598ac549b00b5c138d1c84983a76a0fc3b03b14b2ad1e","observation_id":"6dd64b09-9785-4ce2-871a-99066fc79a19","resolution":{"observed_at":"2026-08-10T21:45:55.387807Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:56.065594Z","title":null,"venue":null,"work_id":"81823051-1058-42f6-bd83-7ba30e5d9b49","year":2023},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.393570Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:2ebe01016c46fe597b3a5983d90a38ea6c4331c282c3696a9fa46bcedbaf1f12","observation_id":"5e0abe4a-1789-4269-8252-b65a78c52366","resolution":{"observed_at":"2026-08-10T21:45:56.072357Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.399566Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.399566Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:3846311fac2d5df54cb2d48334ab4fc356069da5ff28ba731974763e2ce2fbf1","observation_id":"c2e6744a-f9d1-4f49-b42c-d932526ee4ca","resolution":{"observed_at":"2026-08-10T21:45:55.399566Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.13063","last_updated":"2024-03-17T04:38:48Z","snapshot_observed_at":"2026-07-06T15:45:39.649725Z","submitted_at":"2023-06-22T17:31:44Z","title":"Can LLMs Express Their Uncertainty? An Empirical Evaluation of Confidence Elicitation in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.13063","snapshot_observed_at":"2026-08-10T21:45:55.405417Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.405417Z"},"links":{"cited_paper":"/paper/2306.13063","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:78e27842a46d208e73de7c9d3b2772306c9741e12e74b483db9375f42ed350d7","observation_id":"9eb3183a-f463-4224-a7f9-498e8e1e4d67","resolution":{"observed_at":"2026-08-10T21:45:55.405417Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10671","last_updated":"2024-09-10T13:25:53Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-07-15T12:35:42Z","title":"Qwen2 Technical Report","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10671","snapshot_observed_at":"2026-08-10T21:45:55.411727Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.411727Z"},"links":{"cited_paper":"/paper/2407.10671","citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:acb9191a69a82b477cd97eeb818d578bfb6f7094a5dbd9517735742fcbe5a60b","observation_id":"7f1e3389-1cde-4c35-9e7f-39ce8b4296ed","resolution":{"observed_at":"2026-08-10T21:45:55.411727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.418252Z","title":null,"venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.418252Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:f690cf3e6347ac636526826c04df301d2523eb0d9003be46f506d9b1eae289c3","observation_id":"57e8619a-f0e9-43e9-a684-5692d682f9b7","resolution":{"observed_at":"2026-08-10T21:45:55.418252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.423444Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.423444Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:e63b75d48ab245570bf3128606ea71364599a883b1bab1eb320cbe4d7f8a22d9","observation_id":"85412576-37ee-46f4-952a-f5af7f6e1603","resolution":{"observed_at":"2026-08-10T21:45:55.423444Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.428399Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.428399Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:d8e091a507d8c7ac8b42b8ee9bac103bf6fcda9a0b4829714c5707fa9d607c36","observation_id":"e6299e48-8543-4707-871c-c69a9962124b","resolution":{"observed_at":"2026-08-10T21:45:55.428399Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.434586Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.434586Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:fee16aace23b3d943c1a73999cc5c003d1ab1b117bcea9c328a7acbd12aac09c","observation_id":"e9d2127f-85c6-4a37-92af-0b03ad31f6c6","resolution":{"observed_at":"2026-08-10T21:45:55.434586Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T21:45:55.440332Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-10T21:45:55.440332Z"},"links":{"citing_paper":"/paper/2501.03991"},"observation_digest":"sha256:33ce37029ec155bfda0c5ac8e269459c536791cf01fbfe804d9a0104abe86d29","observation_id":"064d0475-4941-4b4c-bc54-af5ad6faeda5","resolution":{"observed_at":"2026-08-10T21:45:55.440332Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2501.03991","last_updated":"2025-01-07T18:48:42Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-11T18:45:26.121373Z","submitted_at":"2025-01-07T18:48:42Z","title":"Influences on LLM Calibration: A Study of Response Agreement, Loss Functions, and Prompt Styles"},"reference_resolution":{"displayed":38,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":36,"verified_exact":1,"verified_fuzzy":1},"total_outbound_references":38},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 38 of 38 outbound references and 2 inbound Pith citation observations for arXiv:2501.03991."}