{"as_of":"2026-08-16T19:15:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:a0b90415f982084531014643ba4bff9becc1627c523328275f3c2aa5877cac8b","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":13,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":13,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":13,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":13,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T19:48:47.164830Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T08:09:40.653154Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-15T19:48:47.164830Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.15068","last_updated":"2025-06-18T02:16:53Z","snapshot_observed_at":"2026-08-16T18:06:28.142966Z","submitted_at":"2025-06-18T02:16:53Z","title":"Semantically-Aware Rewards for Open-Ended R1 Training in Free-Form Generation","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-15T19:48:47.164830Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2506.15068"},"observation_digest":"sha256:099d710b2cd1b5e15140b2fe9539e144beb41de08ee71f7f82f61e23266e0d9f","observation_id":"07198add-f6ef-4dbe-a58c-f159f398f067","resolution":{"observed_at":"2026-08-15T19:48:47.164830Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2509.23542","last_updated":"2026-04-19T04:59:30Z","snapshot_observed_at":"2026-08-14T16:14:11.817391Z","submitted_at":"2025-09-28T00:43:52Z","title":"On the Shelf Life of Fine-Tuned LLM-Judges: Future-Proofing, Backward-Compatibility, and Question Generalization","version":2},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-18T12:53:45.767341Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2509.23542"},"observation_digest":"sha256:458656f9513edcc592a3ae4750beb359777012b422692401992543f96543ebd9","observation_id":"e7697d9a-a7b0-4a90-af30-9147f7f8c8f3","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2510.02837","last_updated":"2026-05-25T02:56:44Z","snapshot_observed_at":"2026-08-12T10:11:19.597552Z","submitted_at":"2025-10-03T09:19:15Z","title":"Beyond the Final Answer: Evaluating the Reasoning Trajectories of Tool-Augmented Agents","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-18T11:02:55.529271Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2510.02837"},"observation_digest":"sha256:7d9e1e8c75cc4b86abbc1a102231736117a88b61a07fee85023599f87d50c70e","observation_id":"e90a2fbc-d89e-4962-a6a7-3d9496f2359f","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2512.19728","last_updated":"2026-04-14T04:23:52Z","snapshot_observed_at":"2026-08-13T20:30:19.092758Z","submitted_at":"2025-12-17T06:15:52Z","title":"Hard Negative Sample-Augmented DPO Post-Training for Small Language Models","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-16T21:57:16.279587Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2512.19728"},"observation_digest":"sha256:ed2a43adca81ae2f850c5c060cd9de4b5f78bfe5cbf5ac152b18b85e885ecb31","observation_id":"d3be9a4e-a668-452a-acfc-b1b922cab7d6","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2604.06820","last_updated":"2026-07-18T03:33:00Z","snapshot_observed_at":"2026-08-15T04:07:32.928990Z","submitted_at":"2026-04-08T08:37:37Z","title":"When Direct Prediction Fails: Evidence from LLM-Based Misinformation Risk Evaluation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T17:57:21.981236Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2604.06820"},"observation_digest":"sha256:b1da00472f125998a8527958c44b67789ea15190b6f1e84505bb73f550e51009","observation_id":"fcfa7d4a-a0b7-4b81-92a0-0671ba477b79","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2604.08559","last_updated":"2026-03-17T09:03:09Z","snapshot_observed_at":"2026-08-07T20:04:32.737155Z","submitted_at":"2026-03-17T09:03:09Z","title":"Medical Reasoning with Large Language Models: A Survey and MR-Bench","version":1},"reference_index":126,"source":"pdf_text","source_observed_at":"2026-05-15T10:21:39.892271Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2604.08559"},"observation_digest":"sha256:78bdf331f784a6a1f475e4bffeb675b8dbf2c7b5ec96fb2a371868eebc00e065","observation_id":"141e6aa8-1ef3-4c27-b759-9604657c2c0c","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2604.17650","last_updated":"2026-04-19T22:45:10Z","snapshot_observed_at":"2026-08-16T17:00:15.780259Z","submitted_at":"2026-04-19T22:45:10Z","title":"Measuring Distribution Shift in User Prompts and Its Effects on LLM Performance","version":1},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-05-10T05:27:50.049798Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2604.17650"},"observation_digest":"sha256:5fd5f1179228b40805d5471b4fee737540104a9e45f7ce13a0078ca5220fcade","observation_id":"66cff144-3f95-4d56-8869-a36f86835261","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2604.18487","last_updated":"2026-04-20T16:37:27Z","snapshot_observed_at":"2026-08-16T05:38:16.741563Z","submitted_at":"2026-04-20T16:37:27Z","title":"Adversarial Humanities Benchmark: Results on Stylistic Robustness in Frontier Model Safety","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-10T04:17:43.880661Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2604.18487"},"observation_digest":"sha256:61cc385ebecc836ffd27466e590303d96f220c0b9c3db4bcb108dd64daf9a0c3","observation_id":"e1ced17e-5ab4-446b-9e9f-30a7ac1159fa","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2605.18805","last_updated":"2026-05-11T18:55:32Z","snapshot_observed_at":"2026-08-13T01:04:11.029880Z","submitted_at":"2026-05-11T18:55:32Z","title":"RecoAtlas: From Semantic Plausibility to Set-Level Utility in LLM Recommendation Agents","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-20T22:27:16.974169Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2605.18805"},"observation_digest":"sha256:fd18913bc01edf06d56a1bd59fbe17a609164da234c3cd6c95e4a7db426ed427","observation_id":"4a2f157e-b045-4c09-ae2c-2b4d601ff26b","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2605.20745","last_updated":"2026-05-20T05:48:16Z","snapshot_observed_at":"2026-08-16T06:41:37.484341Z","submitted_at":"2026-05-20T05:48:16Z","title":"The Hidden Signal of Verifier Strictness: Controlling and Improving Step-Wise Verification via Selective Latent Steering","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-21T06:51:56.556213Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2605.20745"},"observation_digest":"sha256:2eaf51d02915dd131720d62de66185f3dd531b49270d217a193ca7590de1872e","observation_id":"8ef0b962-b5b7-4453-b862-3e0aecdd86e5","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2606.06871","last_updated":"2026-06-05T03:39:58Z","snapshot_observed_at":"2026-08-13T17:41:57.581863Z","submitted_at":"2026-06-05T03:39:58Z","title":"Evidence-Grounded Ensemble Diagnosis of 802.11 Packet Captures: A Multi-Stage Pipeline with Deterministic Reliability Scoring","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-27T22:29:01.722453Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2606.06871"},"observation_digest":"sha256:0f72a55443e165e6bb54131262b246d43017e42e88d4857d6e4aa98df7581506","observation_id":"a84fdf79-e3e6-48f3-aede-20d1233afc3a","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2606.18648","last_updated":"2026-06-23T03:11:08Z","snapshot_observed_at":"2026-08-12T14:16:02.818023Z","submitted_at":"2026-06-17T03:32:06Z","title":"Deep Research in Physical Sciences: A Multi-Agent Framework and Comprehensive Benchmark","version":3},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-26T19:17:17.373463Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2606.18648"},"observation_digest":"sha256:d49ad35549c6c5508144f8cbc40efdccd3a51a9618fa17c91e654714b3bf325a","observation_id":"9cef6ea9-6279-44ea-9322-f04cdf25958e","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding","version":3},"cited_work":{"arxiv_id":"2503.05061","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.05061","snapshot_observed_at":"2026-08-12T02:26:03.833426Z","title":"No free labels: Limitations of llm-as-a-judge without human grounding.arXiv preprint arXiv:2503.05061","venue":null,"work_id":"066d19e0-d5c1-45d7-9959-86ebc413d347","year":2025},"citing_paper":{"arxiv_id":"2606.21943","last_updated":"2026-06-20T08:20:41Z","snapshot_observed_at":"2026-08-15T21:55:25.625601Z","submitted_at":"2026-06-20T08:20:41Z","title":"Modularized Reinforcement Learning on LLMs: From MDP Creation to Exploration and Learning","version":1},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-06-26T12:15:08.304150Z"},"links":{"cited_paper":"/paper/2503.05061","citing_paper":"/paper/2606.21943"},"observation_digest":"sha256:580b0f3a726a542a0c3228257798e6719537537779541c83c6eab390aae197f9","observation_id":"935e1686-125a-4747-8f98-7cda50e740e7","resolution":{"observed_at":"2026-08-12T02:26:03.833426Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2503.05061/citation-record","integrity":"/paper/2503.05061/integrity","json":"/paper/2503.05061/citation-record.json","paper":"/paper/2503.05061"},"outbound":[],"paper":{"arxiv_id":"2503.05061","last_updated":"2026-08-11T15:28:16Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-14T23:11:22.741737Z","submitted_at":"2025-03-07T00:42:08Z","title":"No Free Labels: Limitations of LLM-as-a-Judge Without Human Grounding"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 16 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 13 inbound Pith citation observations for arXiv:2503.05061."}