{"as_of":"2026-08-11T00:05:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3cc6d55915010b8c38c35a08b28239d8aa7d75bea3bee7570fa10588d0906ca1","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":26,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":26,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":26,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":26,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T14:06:14.562327Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":72,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2201.11903","last_updated":"2023-01-10T23:07:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-01-28T02:33:07Z","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","version":6},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-05-10T12:54:44.636760Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2201.11903"},"observation_digest":"sha256:f18ced9baeff448f4b279b440f3389b4cedc65f376408ac9a9c8a8ad3fa364d9","observation_id":"43c2c1c1-123e-40e2-81a8-a558ebf1caf8","resolution":{"observed_at":"2026-05-10T12:54:44.815033Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2206.07682","last_updated":"2022-10-26T05:06:24Z","snapshot_observed_at":"2026-08-02T15:56:35.249569Z","submitted_at":"2022-06-15T17:32:01Z","title":"Emergent Abilities of Large Language Models","version":2},"reference_index":101,"source":"arxiv_source","source_observed_at":"2026-05-11T07:38:37.734402Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2206.07682"},"observation_digest":"sha256:bdd1e428498f8575c113816f3b23bd6e1b9c8687d0e54e30da5078198ce062f6","observation_id":"9575c1b5-daab-40de-863c-fa385ffb3314","resolution":{"observed_at":"2026-05-11T07:38:38.283201Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2212.03827","last_updated":"2024-03-02T21:33:53Z","snapshot_observed_at":"2026-08-10T11:37:23.229129Z","submitted_at":"2022-12-07T18:17:56Z","title":"Discovering Latent Knowledge in Language Models Without Supervision","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-15T20:34:08.207848Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2212.03827"},"observation_digest":"sha256:2541496e6d60dcebdb042d6332c524d531246b5a41e2d9adc6990756b70fe26a","observation_id":"44ec5d43-9a94-4655-9930-48e61a3a461d","resolution":{"observed_at":"2026-05-15T20:34:08.272099Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2308.03958","last_updated":"2024-02-15T01:03:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-07T23:48:36Z","title":"Simple synthetic data reduces sycophancy in large language models","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-16T14:48:08.508109Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2308.03958"},"observation_digest":"sha256:02d3591964fbbbacbd606cabf49b3c83ab84e9595e29386c2d2ef8ce25dc1747","observation_id":"f115fa0c-df30-4ac7-a3c3-83b8846c5a1e","resolution":{"observed_at":"2026-05-16T14:48:08.751738Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-10T14:06:14.562327Z","title":"Calibrate before use: Improving few-shot performance of language models","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.15708","last_updated":"2025-04-18T08:09:26Z","snapshot_observed_at":"2026-08-10T14:22:34.514445Z","submitted_at":"2025-01-27T00:05:12Z","title":"StaICC: Standardized Evaluation for Classification Task in In-context Learning","version":3},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-08-10T14:06:14.562327Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2501.15708"},"observation_digest":"sha256:fdd4965f7db8bc24e6c61eb6e1477953c5425128e6c923d33c4ce4801937f5cd","observation_id":"22e0aee5-c82d-4e50-a060-0a527e33de59","resolution":{"observed_at":"2026-08-10T14:06:14.562327Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-08T22:01:38.783090Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2502.05233","last_updated":"2025-02-07T04:24:07Z","snapshot_observed_at":"2026-08-09T03:26:00.376129Z","submitted_at":"2025-02-07T04:24:07Z","title":"Efficient Knowledge Feeding to Language Models: A Novel Integrated Encoder-Decoder Architecture","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-08T22:01:38.783090Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2502.05233"},"observation_digest":"sha256:521c1335db99105ede6f0cec06c6f254b72125f23f1d99a7e0b391ef984a53ea","observation_id":"0336a6c1-be87-4060-a2d0-22664eed745c","resolution":{"observed_at":"2026-08-08T22:01:38.783090Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-07T14:30:36.967490Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models, 2021, [arXiv:cs.CL/2102.09690]","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.18754","last_updated":"2025-05-24T15:43:25Z","snapshot_observed_at":"2026-08-09T19:53:58.077035Z","submitted_at":"2025-05-24T15:43:25Z","title":"Few-Shot Optimization for Sensor Data Using Large Language Models: A Case Study on Fatigue Detection","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T14:30:36.967490Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2505.18754"},"observation_digest":"sha256:fe4b75fc90247ff029de7ba1d1cc406058cd6a7470b727b3f799afc52cb6359a","observation_id":"1c03bcd1-f052-4941-b54f-48c7515882c8","resolution":{"observed_at":"2026-08-07T14:30:36.967490Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-07T14:17:35.841319Z","title":"Zhao, Eric Wallace, Shi Feng, Dan Klein, and Same er Singh","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.19402","last_updated":"2025-05-26T01:38:02Z","snapshot_observed_at":"2026-08-08T15:10:22.162332Z","submitted_at":"2025-05-26T01:38:02Z","title":"Recalibrating the Compass: Integrating Large Language Models into Classical Research Methods","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T14:17:35.841319Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2505.19402"},"observation_digest":"sha256:aca4f045fb3eb434e2d72090f26f15acd053be99dde8acb1bcbe5ae43d4a9dcb","observation_id":"292d72e3-ab6c-4730-bd59-2d16bf3002d5","resolution":{"observed_at":"2026-08-07T14:17:35.841319Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-07T05:41:36.464492Z","title":"ArXiv:2102.09690 [cs]","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07448","last_updated":"2025-06-09T05:52:03Z","snapshot_observed_at":"2026-08-10T20:04:42.302588Z","submitted_at":"2025-06-09T05:52:03Z","title":"Extending Epistemic Uncertainty Beyond Parameters Would Assist in Designing Reliable LLMs","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T05:41:36.464492Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2506.07448"},"observation_digest":"sha256:8632da22b575415765d4ae5ea82d407f89e3b827617bc28b55a171d06962b0dd","observation_id":"a1fcb81f-1314-4bf3-8c8f-daffe68d9d19","resolution":{"observed_at":"2026-08-07T05:41:36.464492Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-06T22:44:52.166529Z","title":"Zhao, Eric Wallace, Shi Feng, Dan Klein, and Sameer Singh","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.20815","last_updated":"2025-07-08T17:25:34Z","snapshot_observed_at":"2026-08-10T20:03:59.531089Z","submitted_at":"2025-06-25T20:29:46Z","title":"Dynamic Context-Aware Prompt Recommendation for Domain-Specific AI Applications","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T22:44:52.166529Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2506.20815"},"observation_digest":"sha256:dbe839d7a8f98f1680a301fe23710f539966dab8cb3b61c6efcff2fc4b3635f7","observation_id":"2f006f19-1cd3-4b3f-a637-6a3f3cc29bec","resolution":{"observed_at":"2026-08-06T22:44:52.166529Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-06T19:42:02.778812Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.04889","last_updated":"2025-07-07T11:23:20Z","snapshot_observed_at":"2026-08-07T08:41:40.870535Z","submitted_at":"2025-07-07T11:23:20Z","title":"Fine-tuning on simulated data outperforms prompting for agent tone of voice","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T19:42:02.778812Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2507.04889"},"observation_digest":"sha256:2929b2f13e36311356e4efe4ec3bcc3899ad82b2e32e2de447766dd7a61ce360","observation_id":"d87a4944-d290-4699-bcfd-913853402e8d","resolution":{"observed_at":"2026-08-06T19:42:02.778812Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-06T17:20:24.364838Z","title":"Zhao, Eric Wallace, Shi Feng, Dan Klein, and Sameer Singh","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2507.11216","last_updated":"2025-07-15T11:37:30Z","snapshot_observed_at":"2026-08-10T10:24:53.691246Z","submitted_at":"2025-07-15T11:37:30Z","title":"EsBBQ and CaBBQ: The Spanish and Catalan Bias Benchmarks for Question Answering","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-06T17:20:24.364838Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2507.11216"},"observation_digest":"sha256:1108f3b6f4f6c32acd650f561dbed23349896ab4ab327e00ccb8f6e31cc0f145","observation_id":"546cbb8c-00c8-40de-8048-358cbc578a2d","resolution":{"observed_at":"2026-08-06T17:20:24.364838Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T22:58:50.285582Z","title":"cuttable,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.06124","last_updated":"2025-08-08T08:43:24Z","snapshot_observed_at":"2026-08-08T14:46:37.801934Z","submitted_at":"2025-08-08T08:43:24Z","title":"AURA: Affordance-Understanding and Risk-aware Alignment Technique for Large Language Models","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-05T22:58:50.285582Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2508.06124"},"observation_digest":"sha256:443638a203abf50f315044083dfba2c5fd3287d90e0480692415da6416d48adb","observation_id":"283fd66d-35df-44a8-adf0-b197b90a45af","resolution":{"observed_at":"2026-08-05T22:58:50.285582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T10:27:28.012925Z","title":"Calibrate before use: Improving few-shot performance of language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.05375","last_updated":"2025-09-04T11:52:19Z","snapshot_observed_at":"2026-08-09T16:09:15.396679Z","submitted_at":"2025-09-04T11:52:19Z","title":"Characterizing Fitness Landscape Structures in Prompt Engineering","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-05T10:27:28.012925Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2509.05375"},"observation_digest":"sha256:f7bf6fa82925a8a9e9e583100f07cc4236f313e1f91f7eacaec17f91504d0598","observation_id":"553866de-7244-428a-9558-ca78112819e6","resolution":{"observed_at":"2026-08-05T10:27:28.012925Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-04T12:45:27.038741Z","title":"Zhao, Eric Wallace, Shi Feng, Dan Klein, and Sameer Singh","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2510.02480","last_updated":"2026-05-27T22:55:59Z","snapshot_observed_at":"2026-08-09T10:13:00.774448Z","submitted_at":"2025-10-02T18:36:10Z","title":"Controlling the Risk of Corrupted Contexts for Language Models via Early-Exiting","version":3},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-04T12:45:27.038741Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2510.02480"},"observation_digest":"sha256:748e937bbd8f36b06557b259c56c4266e576c1e9a5b2d07ca7b4590afada6ee1","observation_id":"754327b0-c56c-4ab9-ae12-09ffe3fe8fdb","resolution":{"observed_at":"2026-08-04T12:45:27.038741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2603.11689","last_updated":"2026-07-01T08:11:59Z","snapshot_observed_at":"2026-08-09T19:37:08.421363Z","submitted_at":"2026-03-12T08:56:14Z","title":"Explicit Logic Channel for Validation and Enhancement of MLLMs on Zero-Shot Tasks","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-05-21T11:09:51.816554Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2603.11689"},"observation_digest":"sha256:84f3a3248e0de16d9d938374483d81e9c3906fb712cf6df514223474820551ea","observation_id":"f3caae90-762f-40e2-a475-caa43ca87d81","resolution":{"observed_at":"2026-05-21T11:10:02.029576Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2604.15547","last_updated":"2026-04-16T21:52:11Z","snapshot_observed_at":"2026-08-02T19:53:23.399369Z","submitted_at":"2026-04-16T21:52:11Z","title":"Consistency Analysis of Sentiment Predictions using Syntactic & Semantic Context Assessment Summarization (SSAS)","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T10:55:20.435471Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2604.15547"},"observation_digest":"sha256:cb5eeb4597fb22e8ad84e62736895053583a84ef99f6b2328b98719b28071489","observation_id":"6719b4f3-b09e-4215-ad88-0d2cf0016e0e","resolution":{"observed_at":"2026-05-10T11:00:03.692322Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2604.23371","last_updated":"2026-04-25T16:35:25Z","snapshot_observed_at":"2026-07-06T23:09:38.799345Z","submitted_at":"2026-04-25T16:35:25Z","title":"When Context Sticks: Studying Interference in In-Context Learning","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-08T08:31:14.231710Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2604.23371"},"observation_digest":"sha256:c202f6ffbc71815f4c4861463cbf48eb722c455d72997994d8b4a88eef37f23e","observation_id":"cbe87f13-4d7c-4db7-8727-a861c6df8d48","resolution":{"observed_at":"2026-05-11T20:36:09.464038Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2604.24678","last_updated":"2026-04-27T16:38:01Z","snapshot_observed_at":"2026-07-06T23:10:38.548416Z","submitted_at":"2026-04-27T16:38:01Z","title":"Leveraging LLMs for Multi-File DSL Code Generation: An Industrial Case Study","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-05-08T02:49:14.533263Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2604.24678"},"observation_digest":"sha256:538fa0c98c721a2ae0109be13eca3085eb0d1204b72fa1083202316ff294b154","observation_id":"08df5648-85f0-49bd-ba16-739ea8688977","resolution":{"observed_at":"2026-05-11T22:26:13.413182Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2605.22714","last_updated":"2026-06-09T17:40:55Z","snapshot_observed_at":"2026-08-03T02:17:57.312538Z","submitted_at":"2026-05-21T16:51:04Z","title":"AMEL: Accumulated Message Effects on LLM Judgments","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-22T05:08:30.607268Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2605.22714"},"observation_digest":"sha256:17d6c7c9a565f53ce5f256d1af92e54c0f38eb7d03071abb6235f3930a59a974","observation_id":"3518060e-cf38-468b-af03-e914ad9f9640","resolution":{"observed_at":"2026-05-22T05:11:06.762175Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2605.22714","last_updated":"2026-06-09T17:40:55Z","snapshot_observed_at":"2026-08-03T02:17:57.312538Z","submitted_at":"2026-05-21T16:51:04Z","title":"AMEL: Accumulated Message Effects on LLM Judgments","version":3},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-30T17:04:22.688250Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2605.22714"},"observation_digest":"sha256:670462921572048310ad0cf39e96804b11aa55c17dd98f0cf91c68e61662c963","observation_id":"7e90e6ff-17fc-448e-90be-f5211dfe758b","resolution":{"observed_at":"2026-06-30T17:04:56.429057Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2605.24718","last_updated":"2026-05-23T20:13:44Z","snapshot_observed_at":"2026-08-07T15:18:00.875763Z","submitted_at":"2026-05-23T20:13:44Z","title":"The Tokenizer Tax Across 25 European Languages: Domain Invariance, Cross-Lingual Few-Shot Effects, and the Ukrainian Penalty","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-30T13:08:33.172423Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2605.24718"},"observation_digest":"sha256:f7f8c1615d54d0760eb5781cc1add092d43c99d9c282b8a2534d64ed868a6948","observation_id":"0394aa37-285e-468d-8f6e-d7c57a26396a","resolution":{"observed_at":"2026-06-30T13:14:40.898485Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2605.29170","last_updated":"2026-05-27T23:12:20Z","snapshot_observed_at":"2026-07-06T23:38:37.178378Z","submitted_at":"2026-05-27T23:12:20Z","title":"UA-Legal-Bench: A Benchmark for Evaluating Large Language Models on Ukrainian Legal Reasoning","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-29T12:15:12.570661Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2605.29170"},"observation_digest":"sha256:09a6afbe24bb127cd61f4e4186cd26f5e7f4b89be57c06628fd45acff9b0d84c","observation_id":"3ffd7d1b-0f14-4730-bfe1-8b04604c36e4","resolution":{"observed_at":"2026-06-29T12:23:24.536134Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2606.07479","last_updated":"2026-06-05T17:34:07Z","snapshot_observed_at":"2026-08-05T14:47:33.060849Z","submitted_at":"2026-06-05T17:34:07Z","title":"Supervision versus Demonstration-Based In-Context Learning for Multiword Expression Classification","version":1},"reference_index":219,"source":"arxiv_source","source_observed_at":"2026-06-27T22:03:52.672777Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2606.07479"},"observation_digest":"sha256:d074f327a086bc8822d8deaed1270dd3bcb78b061dd197f8523d37e5146826fd","observation_id":"d93c0237-f6f9-4dba-9e95-256509827e67","resolution":{"observed_at":"2026-07-02T17:27:14.909483Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":"2102.09690","doi":"10.48550/arxiv.2102.09690","metadata_source":"arxiv_reference","pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Calibrate before use: Improving few- shot performance of language models","venue":"arXiv (Cornell University)","work_id":"cd67a9c7-e56a-43c3-8688-0cfddf32f697","year":2021},"citing_paper":{"arxiv_id":"2606.22470","last_updated":"2026-06-21T12:30:23Z","snapshot_observed_at":"2026-08-05T21:38:55.898342Z","submitted_at":"2026-06-21T12:30:23Z","title":"PRIME: Evaluating Prompt Resolution Under Incompatible Instructions in LLMs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-26T11:03:19.265228Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2606.22470"},"observation_digest":"sha256:220e102f96344020df51c725db88f093bcc341190568b185e57c072a71480afd","observation_id":"2cdc9b74-bbd5-45d9-9142-1971a6c52cc0","resolution":{"observed_at":"2026-07-04T08:49:41.704426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-08T21:08:07.955868+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.09690","snapshot_observed_at":"2026-07-31T12:20:06.584062Z","title":"and Wallace, Eric and Feng, Shi and others , year = 2021, month = jun, publisher =","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28282","last_updated":"2026-07-30T14:31:07Z","snapshot_observed_at":"2026-08-06T16:34:19.905250Z","submitted_at":"2026-07-30T14:31:07Z","title":"(Towards) Scalable Reliable Automated Evaluation with Large Language Models","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-07-31T12:20:06.584062Z"},"links":{"cited_paper":"/paper/2102.09690","citing_paper":"/paper/2607.28282"},"observation_digest":"sha256:593c64b9ac50c37e4144bb9ecd6c3af3f15f0ed29a5f21fb2bedb52cc3967fc9","observation_id":"b49f1bd0-4cc1-4abf-a062-4043cebe7476","resolution":{"observed_at":"2026-07-31T12:20:06.584062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2102.09690/citation-record","integrity":"/paper/2102.09690/integrity","json":"/paper/2102.09690/citation-record.json","paper":"/paper/2102.09690"},"outbound":[],"paper":{"arxiv_id":"2102.09690","last_updated":"2021-06-10T18:20:59Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-10T20:03:40.241346Z","submitted_at":"2021-02-19T00:23:59Z","title":"Calibrate Before Use: Improving Few-Shot Performance of Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 26 inbound Pith citation observations for arXiv:2102.09690."}