{"as_of":"2026-08-09T17:01:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7001a7270d4769fe8effba8ab6b5624588964a44b4546a6c3948fa6a85c688a7","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":17,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":17,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":17,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":17,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:12:24.340787Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T17:48:45.447795Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":"2407.10457","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-07-03T17:48:45.447795Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism","venue":null,"work_id":"7b257233-7d88-48df-8208-6dea20dbedea","year":2024},"citing_paper":{"arxiv_id":"2407.21787","last_updated":"2024-12-30T19:03:24Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:57:25Z","title":"Large Language Monkeys: Scaling Inference Compute with Repeated Sampling","version":3},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-05-12T04:42:23.297389Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2407.21787"},"observation_digest":"sha256:38d71effc6611a954529a67cd14e62c6770f946fcdeb290158db6e1fb33708f5","observation_id":"77c12347-f8b1-4204-aee5-3b194378bfd0","resolution":{"observed_at":"2026-05-12T04:42:23.449917Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-07T15:12:24.340787Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17131","last_updated":"2025-05-22T01:59:54Z","snapshot_observed_at":"2026-08-07T23:47:01.265288Z","submitted_at":"2025-05-22T01:59:54Z","title":"Relative Bias: A Comparative Framework for Quantifying Bias in LLMs","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-07T15:12:24.340787Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2505.17131"},"observation_digest":"sha256:1f363af00f3860fd5baabc6a7f3eb7e772d3de21a7b918d744195291cfd3788a","observation_id":"3ea5021e-aa06-4c4f-a1cb-bce3c7a11016","resolution":{"observed_at":"2026-08-07T15:12:24.340787Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-07T12:35:33.261795Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24613","last_updated":"2025-05-30T14:04:30Z","snapshot_observed_at":"2026-08-08T07:31:58.850783Z","submitted_at":"2025-05-30T14:04:30Z","title":"When Harry Meets Superman: The Role of The Interlocutor in Persona-Based Dialogue Generation","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T12:35:33.261795Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2505.24613"},"observation_digest":"sha256:2b3931f82be69af2ee73e172484d09f5c812a9926020af7dc717afe8adec4375","observation_id":"cf93851c-6879-442c-8940-07f820961e16","resolution":{"observed_at":"2026-08-07T12:35:33.261795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-07T05:54:34.908848Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.06832","last_updated":"2025-06-22T12:08:32Z","snapshot_observed_at":"2026-08-09T17:01:08.454683Z","submitted_at":"2025-06-07T15:25:10Z","title":"Cross-Entropy Games for Language Models: From Implicit Knowledge to General Capability Measures","version":2},"reference_index":2018,"source":"pdf_text","source_observed_at":"2026-08-07T05:54:34.908848Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2506.06832"},"observation_digest":"sha256:0180eadaf56d96fdaa014467aeedf4c9479be05b67955834d4de92857a9bf368","observation_id":"814010f8-3418-43b5-a70c-ea56cc92e6ae","resolution":{"observed_at":"2026-08-07T05:54:34.908848Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-06T23:57:21.346054Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism.arXiv preprint arXiv:2407.10457, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15649","last_updated":"2025-06-18T17:23:36Z","snapshot_observed_at":"2026-08-07T09:16:49.443338Z","submitted_at":"2025-06-18T17:23:36Z","title":"Dual-Stage Value-Guided Inference with Margin-Based Reward Adjustment for Fast and Faithful VLM Captioning","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T23:57:21.346054Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2506.15649"},"observation_digest":"sha256:85ec18d17a33f25d554d0f761f186a5d39f118561a7f4bcea799a4ff096aedb6","observation_id":"2083ec24-9664-4998-aa1a-381b6a7f2f86","resolution":{"observed_at":"2026-08-06T23:57:21.346054Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-06T20:40:18.647502Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.02226","last_updated":"2025-07-03T01:17:44Z","snapshot_observed_at":"2026-08-06T20:31:57.968196Z","submitted_at":"2025-07-03T01:17:44Z","title":"DecoRTL: A Run-time Decoding Framework for RTL Code Generation with LLMs","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T20:40:18.647502Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2507.02226"},"observation_digest":"sha256:20c9a5a25ce56045d82c0820ac66edbb83c558dcc039fc7bc4c061975d74a4ca","observation_id":"6a2e6521-c894-4df1-a800-a5eac1bfc4ca","resolution":{"observed_at":"2026-08-06T20:40:18.647502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-06T19:38:19.498886Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.05056","last_updated":"2025-07-22T07:33:11Z","snapshot_observed_at":"2026-08-09T11:47:01.555437Z","submitted_at":"2025-07-07T14:38:53Z","title":"INTER: Mitigating Hallucination in Large Vision-Language Models by Interaction Guidance Sampling","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T19:38:19.498886Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2507.05056"},"observation_digest":"sha256:20cb5238b715a8426545d70fa221f273f627dfbafdad1d4f02797be663b2e830","observation_id":"4b4ff39a-1a9f-4ecc-9496-2c30b67cc78b","resolution":{"observed_at":"2026-08-06T19:38:19.498886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-06T18:24:52.743707Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.08432","last_updated":"2025-07-11T09:18:41Z","snapshot_observed_at":"2026-08-09T14:15:09.598647Z","submitted_at":"2025-07-11T09:18:41Z","title":"xpSHACL: Explainable SHACL Validation using Retrieval-Augmented Generation and Large Language Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T18:24:52.743707Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2507.08432"},"observation_digest":"sha256:adea137a087b401982d2baf71ca3abc678dec37344142103d170151bd6bed746","observation_id":"37f20783-4d39-41ca-8a24-b059a7821d1f","resolution":{"observed_at":"2026-08-06T18:24:52.743707Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-06T15:29:41.138123Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.15715","last_updated":"2025-08-05T20:19:37Z","snapshot_observed_at":"2026-08-09T11:28:54.946827Z","submitted_at":"2025-07-21T15:26:58Z","title":"From Queries to Criteria: Understanding How Astronomers Evaluate LLMs","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-06T15:29:41.138123Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2507.15715"},"observation_digest":"sha256:95996e0ce870de26209226c082dc80314a5e73909194c74d73345d45d55cab18","observation_id":"71985320-c1e8-4ecf-a5dc-8468c119c5d6","resolution":{"observed_at":"2026-08-06T15:29:41.138123Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-05T23:02:03.070200Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.06017","last_updated":"2025-08-08T05:04:47Z","snapshot_observed_at":"2026-08-08T10:44:52.352035Z","submitted_at":"2025-08-08T05:04:47Z","title":"Position: Intelligent Coding Systems Should Write Programs with Justifications","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-05T23:02:03.070200Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2508.06017"},"observation_digest":"sha256:8ae8fa0f86c582cc38654ca9eb0f46eb43bdadb559abfaa9788a193146f4d6d2","observation_id":"a54752ae-d3fb-46d3-b33a-5853707f61c2","resolution":{"observed_at":"2026-08-05T23:02:03.070200Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-05T05:30:12.058244Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.09705","last_updated":"2025-09-05T17:31:14Z","snapshot_observed_at":"2026-08-07T10:14:14.572744Z","submitted_at":"2025-09-05T17:31:14Z","title":"The Non-Determinism of Small LLMs: Evidence of Low Answer Consistency in Repetition Trials of Standard Multiple-Choice Benchmarks","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-05T05:30:12.058244Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2509.09705"},"observation_digest":"sha256:0a7bfd379496fad2082c0a30ac43151dc859e84dd4e22520fa5b8c8237a2a555","observation_id":"c64a07ba-1aca-4d10-bbe2-35b3646355a2","resolution":{"observed_at":"2026-08-05T05:30:12.058244Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-04T11:20:12.536154Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism.arXiv preprint arXiv:2407.10457,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2510.05681","last_updated":"2026-07-04T15:45:15Z","snapshot_observed_at":"2026-08-04T11:20:10.610981Z","submitted_at":"2025-10-07T08:38:08Z","title":"Verifier-free Test-Time Sampling for Vision-Language-Action Models","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-04T11:20:12.536154Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2510.05681"},"observation_digest":"sha256:067607f752df438a55ff32b4c2d83739dce9b58bc5a8dbc2f24eacb10985b541","observation_id":"cc762b08-94c5-4256-ba40-cb43a3bb11b8","resolution":{"observed_at":"2026-08-04T11:20:12.536154Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":"2407.10457","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-07-03T17:48:45.447795Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism","venue":null,"work_id":"7b257233-7d88-48df-8208-6dea20dbedea","year":2024},"citing_paper":{"arxiv_id":"2604.10350","last_updated":"2026-04-11T21:20:57Z","snapshot_observed_at":"2026-08-04T00:28:43.791967Z","submitted_at":"2026-04-11T21:20:57Z","title":"LLM-based Generation of Semantically Diverse and Realistic Domain Model Instances","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-10T15:25:15.447353Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2604.10350"},"observation_digest":"sha256:b56533b8e66f6fde831389018909a1ab18586fb1c02b82488d31aa0629596650","observation_id":"f5f89fe1-a490-4ca1-b1cc-e5f40e26081a","resolution":{"observed_at":"2026-05-11T10:36:04.165624Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":"2407.10457","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-07-03T17:48:45.447795Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism","venue":null,"work_id":"7b257233-7d88-48df-8208-6dea20dbedea","year":2024},"citing_paper":{"arxiv_id":"2604.19598","last_updated":"2026-04-23T01:33:33Z","snapshot_observed_at":"2026-08-02T21:34:48.936053Z","submitted_at":"2026-04-21T15:51:46Z","title":"Cross-Model Consistency of AI-Generated Exercise Prescriptions: A Repeated Generation Study Across Three Large Language Models","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-10T03:21:17.833685Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2604.19598"},"observation_digest":"sha256:8f63803734541e2efcd1eb0345d7dd3ce3d6873fb66444050cd602fbf9745f20","observation_id":"a0ef5d5a-2568-4f80-8059-4d0eddeba328","resolution":{"observed_at":"2026-05-11T12:41:01.290765Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":"2407.10457","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-07-03T17:48:45.447795Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism","venue":null,"work_id":"7b257233-7d88-48df-8208-6dea20dbedea","year":2024},"citing_paper":{"arxiv_id":"2604.22411","last_updated":"2026-04-24T10:11:06Z","snapshot_observed_at":"2026-07-06T23:08:47.078992Z","submitted_at":"2026-04-24T10:11:06Z","title":"Introducing Background Temperature to Characterise Hidden Randomness in Large Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-08T12:14:47.518940Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2604.22411"},"observation_digest":"sha256:8c6ba8546c03b0a1ead943d71bd543cb2c02af7a08ff73745f4e6139a8be912b","observation_id":"3e924719-3349-4d20-8122-d8a15fbf44a5","resolution":{"observed_at":"2026-05-11T19:21:07.498219Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":"2407.10457","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-07-03T17:48:45.447795Z","title":"The good, the bad, and the greedy: Evaluation of llms should not ignore non-determinism","venue":null,"work_id":"7b257233-7d88-48df-8208-6dea20dbedea","year":2024},"citing_paper":{"arxiv_id":"2606.17182","last_updated":"2026-06-15T18:19:34Z","snapshot_observed_at":"2026-08-04T04:35:27.820830Z","submitted_at":"2026-06-15T18:19:34Z","title":"Verified Detection and Prevention of Concurrency Anomalies in Multi-Agent Large Language Model Systems","version":1},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-06-27T03:53:14.728371Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2606.17182"},"observation_digest":"sha256:a8f164389874baebb629596370ebb59c076688aa18e95d6fd7e0c80e6f9bcbb1","observation_id":"7f4d4cb6-4d5e-4617-83e6-0704a679fda3","resolution":{"observed_at":"2026-07-03T17:48:45.449869Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.10457","snapshot_observed_at":"2026-08-02T06:23:48.374729Z","title":"The good, the bad, and the greedy: Evaluation of LLMs should not ignore non-determinism.arXiv preprint arXiv:2407.10457, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.12796","last_updated":"2026-07-25T18:18:32Z","snapshot_observed_at":"2026-08-07T16:30:53.506949Z","submitted_at":"2026-07-14T14:12:05Z","title":"The One-Word Census: Answer-Choice Conformity Across 44 Language Models","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-02T06:23:48.374729Z"},"links":{"cited_paper":"/paper/2407.10457","citing_paper":"/paper/2607.12796"},"observation_digest":"sha256:3a13df1d5fcf37b14e4819506bef57c73ae5de7e31b1b5091cb497b1e54982b7","observation_id":"4855884f-df11-46f7-928d-54c99b0614d4","resolution":{"observed_at":"2026-08-02T06:23:48.374729Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2407.10457/citation-record","integrity":"/paper/2407.10457/integrity","json":"/paper/2407.10457/citation-record.json","paper":"/paper/2407.10457"},"outbound":[],"paper":{"arxiv_id":"2407.10457","last_updated":"2024-07-15T06:12:17Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T18:46:18.540571Z","submitted_at":"2024-07-15T06:12:17Z","title":"The Good, The Bad, and The Greedy: Evaluation of LLMs Should Not Ignore Non-Determinism"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 17 inbound Pith citation observations for arXiv:2407.10457."}