{"as_of":"2026-08-20T08:35:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:1786b5a5f025376d305a99140c47427ceabc9c4e02c65e5bd29a68496bd9db84","coverage":[{"denominator":9,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":9,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T14:56:08.354492Z","state":"measured"},{"denominator":9,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":9,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.03437/citation-record","integrity":"/paper/2608.03437/integrity","json":"/paper/2608.03437/citation-record.json","paper":"/paper/2608.03437"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T14:56:08.612745Z","title":"ambiguity","venue":null,"work_id":"b7c962e1-1161-44e9-b078-a9a1a32d8cba","year":2010},"citing_paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T14:56:08.354492Z"},"links":{"citing_paper":"/paper/2608.03437"},"observation_digest":"sha256:61f92127867e23fecec88dce43524cea20db150e34e7ef0f48933f683734e83d","observation_id":"1641ae9b-d2e4-4207-9c65-71450032d46b","resolution":{"observed_at":"2026-08-15T14:56:08.617830Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.20879","last_updated":"2025-05-12T16:33:58Z","snapshot_observed_at":"2026-08-17T00:43:37.683523Z","submitted_at":"2025-04-29T15:48:49Z","title":"The Leaderboard Illusion","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.20879","snapshot_observed_at":"2026-08-15T14:56:08.343727Z","title":"In ICLR 2024 Workshop on Mathematical and Empirical Understanding of Foundation Models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T14:56:08.343727Z"},"links":{"cited_paper":"/paper/2504.20879","citing_paper":"/paper/2608.03437"},"observation_digest":"sha256:914275b3684622995c325774970ee8d6132e7e9b243b59babc5826538428bf31","observation_id":"d084149c-080f-49d0-9537-740a245079bc","resolution":{"observed_at":"2026-08-15T14:56:08.343727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.06172","last_updated":"2025-02-26T21:53:59Z","snapshot_observed_at":"2026-08-19T21:01:54.758174Z","submitted_at":"2024-07-08T17:48:42Z","title":"On Speeding Up Language Model Evaluation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.06172","snapshot_observed_at":"2026-08-15T14:56:08.350983Z","title":"Advances in neural information pro ­ cessing systems 36:46595–46623","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T14:56:08.350983Z"},"links":{"cited_paper":"/paper/2407.06172","citing_paper":"/paper/2608.03437"},"observation_digest":"sha256:c508bc1c37cc996cdf8806bd610f01720065de08929c211e58126905b4b324b7","observation_id":"e694aa7b-e7d9-4cac-84b6-701bafa4fab4","resolution":{"observed_at":"2026-08-15T14:56:08.350983Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T14:56:08.322213Z","title":"Machine learning 47:235–256","venue":null,"work_id":null,"year":1985},"citing_paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking","version":1},"reference_index":2002,"source":"pdf_text","source_observed_at":"2026-08-15T14:56:08.322213Z"},"links":{"citing_paper":"/paper/2608.03437"},"observation_digest":"sha256:8a3f5a50e9ae5438ff4065da52d72920480b7e202d89e0cc75ad1c6a0f6affeb","observation_id":"8a04d905-15d1-4ec7-8fbd-4320f4df45ad","resolution":{"observed_at":"2026-08-15T14:56:08.322213Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T14:56:08.635402Z","title":"Journal of Machine Learning Research 7(39):1079–1105","venue":null,"work_id":"776729d9-e1fb-4503-b899-08f57613b430","year":2021},"citing_paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking","version":1},"reference_index":2006,"source":"pdf_text","source_observed_at":"2026-08-15T14:56:08.326626Z"},"links":{"citing_paper":"/paper/2608.03437"},"observation_digest":"sha256:c282d320d94f6e197a8379d84ed96a68195ebfd878484fd1e30f4c8f7f79cbbb","observation_id":"26a9566e-b7a9-40dd-8ef3-da7cad4ebe2f","resolution":{"observed_at":"2026-08-15T14:56:08.638714Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2601.07651","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T14:56:08.546147Z","title":"In Proceedings of the 30th International Conference on Machine Learning, pages 1238–1246, Atlanta, Georgia, USA","venue":null,"work_id":"9e97f37f-e359-476b-a828-cc0024054842","year":2025},"citing_paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking","version":1},"reference_index":2013,"source":"pdf_text","source_observed_at":"2026-08-15T14:56:08.333952Z"},"links":{"citing_paper":"/paper/2608.03437"},"observation_digest":"sha256:935e10c5f8bb680a8bfd0db67c8d29920600bfdb746c292db530690e7890a599","observation_id":"64b4135e-2fff-4683-b450-6d41b972b6f7","resolution":{"observed_at":"2026-08-15T14:56:08.552563Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T14:56:08.624787Z","title":"In Proceedings of the Eighth Conference on Machine Translation , pages 756–767, Singapore","venue":null,"work_id":"84da851d-7572-4936-8f73-2610e4c9b7a1","year":2023},"citing_paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-15T14:56:08.330369Z"},"links":{"citing_paper":"/paper/2608.03437"},"observation_digest":"sha256:6c3af8acc099d164567883ab528c50baf621acc5e4cd1e8c7c0225f600931a3b","observation_id":"b9f86210-17d7-4941-92c7-52afd3eec674","resolution":{"observed_at":"2026-08-15T14:56:08.628488Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.24664","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T14:56:08.484506Z","title":"In Proceedings of the 41st International Conference on Machine Learning, Vienna, Austria","venue":null,"work_id":"a289b700-09f1-4241-8b8c-b2c48d45a4f6","year":2025},"citing_paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-15T14:56:08.339704Z"},"links":{"citing_paper":"/paper/2608.03437"},"observation_digest":"sha256:cd46ef775df2a449981cdea224cad7671a31224a6ddea00039eb3102594d7cf3","observation_id":"594d7e53-d6e9-41f6-9845-5f45f8add87c","resolution":{"observed_at":"2026-08-15T14:56:08.492175Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.03814","last_updated":"2025-05-02T17:05:01Z","snapshot_observed_at":"2026-08-17T15:11:12.016051Z","submitted_at":"2025-05-02T17:05:01Z","title":"Cer-Eval: Certifiable and Cost-Efficient Evaluation Framework for LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.03814","snapshot_observed_at":"2026-08-15T14:56:08.347771Z","title":"ArXiv: 2505.03814 [stat.ML]","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-15T14:56:08.347771Z"},"links":{"cited_paper":"/paper/2505.03814","citing_paper":"/paper/2608.03437"},"observation_digest":"sha256:6e40ec86f5cfb439980c90c68d9789e18422479f57fd9b46c5e606daeb3d8d28","observation_id":"60bbfc65-7c72-4eae-aa35-660be58f2965","resolution":{"observed_at":"2026-08-15T14:56:08.347771Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2608.03437","last_updated":"2026-08-04T10:33:04Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-20T06:03:14.325775Z","submitted_at":"2026-08-04T10:33:04Z","title":"Dynamically Allocating Evaluation Effort for Model Ranking"},"reference_resolution":{"displayed":9,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":4,"verified_exact":2,"verified_fuzzy":3},"total_outbound_references":9},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 9 of 9 outbound references and 0 inbound Pith citation observations for arXiv:2608.03437."}