{"as_of":"2026-08-10T03:42:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:54cf66b865e72bf71e049600b9787f9f7cb0b183d0125e082d0f0b5f05a1e81d","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":33,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":33,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":33,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":33,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T20:34:01.258829Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":54,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2305.17926","last_updated":"2023-08-30T13:22:35Z","snapshot_observed_at":"2026-08-08T03:32:23.667881Z","submitted_at":"2023-05-29T07:41:03Z","title":"Large Language Models are not Fair Evaluators","version":2},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-05-17T12:10:42.248005Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2305.17926"},"observation_digest":"sha256:7885d8a43d0929905c3511ec021bf84461b5c17f930f76456cf044cc2c3e1793","observation_id":"c633b906-40bb-48c1-b2ad-7cb858c7dcb6","resolution":{"observed_at":"2026-05-17T12:10:42.405853Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2306.05685","last_updated":"2023-12-24T02:01:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-09T05:55:52Z","title":"Judging LLM-as-a-Judge with MT-Bench and Chatbot Arena","version":4},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T18:52:59.033645Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2306.05685"},"observation_digest":"sha256:2dee0be5ea8ccaf30d0988d12f8bfde9449811466b418d9148b6686253b275a6","observation_id":"41c776e8-d9d2-4e35-a3a2-e329d193b2fa","resolution":{"observed_at":"2026-05-10T18:52:59.178783Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2306.11644","last_updated":"2023-10-02T06:12:30Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-20T16:14:25Z","title":"Textbooks Are All You Need","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-13T04:44:03.148223Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2306.11644"},"observation_digest":"sha256:84ef355a456631569661c358b79e08723f04005d54c46c2d3cf634967dc20ff6","observation_id":"d04362a3-2474-402b-83e2-bd6d5380fe24","resolution":{"observed_at":"2026-05-13T04:44:03.204657Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2307.06435","last_updated":"2024-10-17T01:10:40Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-07-12T20:01:52Z","title":"A Comprehensive Overview of Large Language Models","version":10},"reference_index":174,"source":"pdf_text","source_observed_at":"2026-05-19T20:28:38.900026Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2307.06435"},"observation_digest":"sha256:aeb7d5d46f92f807dedd1d7d01c7eab9385370c565b2be86080d527baf7499c0","observation_id":"0404f66a-62af-479f-a6d8-65ffa17db7b1","resolution":{"observed_at":"2026-05-19T20:28:39.557301Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2309.00614","last_updated":"2023-09-04T17:47:36Z","snapshot_observed_at":"2026-07-06T16:13:23.343694Z","submitted_at":"2023-09-01T17:59:44Z","title":"Baseline Defenses for Adversarial Attacks Against Aligned Language Models","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-13T23:24:39.835347Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2309.00614"},"observation_digest":"sha256:d9e9476faa632743cd3b718626d2aa5a2e344f79bef8a76d3b1ee495a11090e3","observation_id":"cd184d77-cc36-4485-9d2c-7d16d87aa969","resolution":{"observed_at":"2026-05-13T23:24:40.111516Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2309.11495","last_updated":"2023-09-25T15:25:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-20T17:50:55Z","title":"Chain-of-Verification Reduces Hallucination in Large Language Models","version":2},"reference_index":138,"source":"arxiv_source","source_observed_at":"2026-05-18T01:06:49.811982Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2309.11495"},"observation_digest":"sha256:b7e3212eced34b81f17f46ecdf199dd8670c06aa0b8035278aa24b38263f5cbc","observation_id":"fbf82ab5-4dae-4e5f-ae5a-b0b1bec77a67","resolution":{"observed_at":"2026-05-18T01:06:50.331746Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2309.14525","last_updated":"2023-09-25T20:59:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-25T20:59:33Z","title":"Aligning Large Multimodal Models with Factually Augmented RLHF","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-15T17:58:17.699042Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2309.14525"},"observation_digest":"sha256:1f03e2598bdaf58893c946baff1cc0d97793cf121b11f0cf1fa7bce689e248f2","observation_id":"59a8c904-4c23-47ce-8f37-ff774f2cd88e","resolution":{"observed_at":"2026-05-15T17:58:17.843448Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2310.11511","last_updated":"2023-10-17T18:18:32Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-17T18:18:32Z","title":"Self-RAG: Learning to Retrieve, Generate, and Critique through Self-Reflection","version":1},"reference_index":137,"source":"arxiv_source","source_observed_at":"2026-05-12T14:15:10.907921Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2310.11511"},"observation_digest":"sha256:a173f5c98d95efc2008497ade381af939475ebfa05c83ebf9c9ca5e8b2c93d15","observation_id":"d214e8c3-495a-4794-b92a-6fdfe98f9429","resolution":{"observed_at":"2026-05-12T14:15:11.113391Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2311.16867","last_updated":"2023-11-29T19:45:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-11-28T15:12:47Z","title":"The Falcon Series of Open Language Models","version":2},"reference_index":276,"source":"arxiv_source","source_observed_at":"2026-05-16T09:46:09.701440Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2311.16867"},"observation_digest":"sha256:f8c5b3491e5462e13bdb0df0bbb5b4ced24a05ca646f801c75c443647287990f","observation_id":"6182808e-2263-4bc4-ae30-791c538dcd7c","resolution":{"observed_at":"2026-05-16T09:46:10.083261Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2401.10020","last_updated":"2025-03-28T00:06:51Z","snapshot_observed_at":"2026-08-07T08:02:34.857823Z","submitted_at":"2024-01-18T14:43:47Z","title":"Self-Rewarding Language Models","version":3},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-05-13T12:01:42.290502Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2401.10020"},"observation_digest":"sha256:97da28332aee5d83cb1b9a2b694b25929a9260a435539ebfdefc5546406041f5","observation_id":"650f1700-58a3-4bc1-8196-94d98ab8091e","resolution":{"observed_at":"2026-05-13T12:01:42.411055Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2401.15884","last_updated":"2024-10-07T02:19:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-29T04:36:39Z","title":"Corrective Retrieval Augmented Generation","version":3},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-12T11:19:17.120464Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2401.15884"},"observation_digest":"sha256:31c036ee434838640f8c23719dbfd9c207f68e216a150832783c4da6dfe8b78c","observation_id":"c0008a37-ac01-4e72-8eb4-452843efdb90","resolution":{"observed_at":"2026-05-12T11:19:17.301138Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2406.04244","last_updated":"2024-06-06T16:41:39Z","snapshot_observed_at":"2026-07-30T15:43:06.151242Z","submitted_at":"2024-06-06T16:41:39Z","title":"Benchmark Data Contamination of Large Language Models: A Survey","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-22T23:10:40.420241Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2406.04244"},"observation_digest":"sha256:9230abd06a31fe4dea492c31647affa503e67a5aa699c948b88850bdf6ab4faa","observation_id":"98eb55fd-94b7-446f-b2ff-518e5b4fab1b","resolution":{"observed_at":"2026-05-22T23:10:40.926678Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2408.00724","last_updated":"2025-03-03T07:53:32Z","snapshot_observed_at":"2026-08-09T23:53:42.648697Z","submitted_at":"2024-08-01T17:16:04Z","title":"Inference Scaling Laws: An Empirical Analysis of Compute-Optimal Inference for Problem-Solving with Language Models","version":3},"reference_index":288,"source":"arxiv_source","source_observed_at":"2026-05-18T06:38:36.517935Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2408.00724"},"observation_digest":"sha256:3f95ec2d5e9615ee1bbfed17bbf90e1bd3a14395e8881d95d51c72c6e21d10a2","observation_id":"a4007580-dd99-4d7b-99b3-570889290493","resolution":{"observed_at":"2026-05-18T06:38:37.165372Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-09T20:34:01.258829Z","title":"X., Taori, R., Zhang, T., Gulrajani, I., Ba, J., Guestrin, C., Liang, P","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.19358","last_updated":"2025-06-02T16:30:23Z","snapshot_observed_at":"2026-08-09T20:35:40.499093Z","submitted_at":"2025-01-31T18:10:53Z","title":"The Energy Loss Phenomenon in RLHF: A New Perspective on Mitigating Reward Hacking","version":3},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-09T20:34:01.258829Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2501.19358"},"observation_digest":"sha256:a68fb739cb55aecc34d8ea167016764030b9254afae43516c88b8212ed48e04b","observation_id":"d359e29c-cdbe-41d1-a64e-0fa9d0f514e9","resolution":{"observed_at":"2026-08-09T20:34:01.258829Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-07T11:59:22.616939Z","title":"Alpacafarm: A simulation framework for methods that learn from human feedback","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.01055","last_updated":"2025-06-01T15:48:06Z","snapshot_observed_at":"2026-08-09T00:34:14.908419Z","submitted_at":"2025-06-01T15:48:06Z","title":"Simple Prompt Injection Attacks Can Leak Personal Data Observed by LLM Agents During Task Execution","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:59:22.616939Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2506.01055"},"observation_digest":"sha256:aa97035040ad126af6b0c9d5598af738f5b5780ca22f0558868fb6c95c12392d","observation_id":"662f2b23-c9d8-4462-b276-496aa0e30648","resolution":{"observed_at":"2026-08-07T11:59:22.616939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-07T05:21:58.639987Z","title":"Hashimoto","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08266","last_updated":"2025-06-09T22:03:56Z","snapshot_observed_at":"2026-08-10T01:49:26.719459Z","submitted_at":"2025-06-09T22:03:56Z","title":"Reinforcement Learning from Human Feedback with High-Confidence Safety Constraints","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T05:21:58.639987Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2506.08266"},"observation_digest":"sha256:2b3307f38d28d6a9dba1d8199a496b83fa4b8525513b08ce668790d237456be4","observation_id":"e8c042f1-7573-493f-aac2-c9a443c337f4","resolution":{"observed_at":"2026-08-07T05:21:58.639987Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-06T14:51:53.887056Z","title":"arXiv:2305.14387","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.17477","last_updated":"2025-07-23T13:00:00Z","snapshot_observed_at":"2026-08-07T11:10:51.753612Z","submitted_at":"2025-07-23T13:00:00Z","title":"An Uncertainty-Driven Adaptive Self-Alignment Framework for Large Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T14:51:53.887056Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2507.17477"},"observation_digest":"sha256:9bae0f1119c3116199c94c39a4b8715802c5b395641248fff05a91fa403e1bf6","observation_id":"7ec1677e-173c-4452-ab99-ce10bfeee767","resolution":{"observed_at":"2026-08-06T14:51:53.887056Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-04T17:56:38.763239Z","title":"Hashimoto","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.10397","last_updated":"2026-07-25T20:09:07Z","snapshot_observed_at":"2026-08-07T01:01:23.391389Z","submitted_at":"2025-09-12T16:44:34Z","title":"RecoWorld: Building Simulated Environments for Agentic Recommender Systems","version":3},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-04T17:56:38.763239Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2509.10397"},"observation_digest":"sha256:060849eca78ef7fe4d24ecf4ce63d150b4118fcbdbb5a9f1dab4d3a1050087d5","observation_id":"c0e8ab4e-2af1-4680-84be-b13782bead09","resolution":{"observed_at":"2026-08-04T17:56:38.763239Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2509.23542","last_updated":"2026-04-19T04:59:30Z","snapshot_observed_at":"2026-08-03T18:10:00.237919Z","submitted_at":"2025-09-28T00:43:52Z","title":"On the Shelf Life of Fine-Tuned LLM-Judges: Future-Proofing, Backward-Compatibility, and Question Generalization","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-18T12:53:45.767341Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2509.23542"},"observation_digest":"sha256:81ae869741adb13456c2d8c8189b4983b241305eda15b7fceb44df0229243dfb","observation_id":"0954badc-258c-4482-983e-fd21e59143d4","resolution":{"observed_at":"2026-05-18T12:56:24.568395Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2601.14053","last_updated":"2026-04-16T15:26:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-01-20T15:06:19Z","title":"LLMOrbit: A Circular Taxonomy of Large Language Models -From Scaling Walls to Agentic AI Systems","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-16T12:47:28.248540Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2601.14053"},"observation_digest":"sha256:6a5f2e7faea749417ef45c1abaec4dd4b121a50c25cf7c915fa2b6f5ef87f0c0","observation_id":"603a3ffd-b380-48b2-b709-f295c5a3a440","resolution":{"observed_at":"2026-05-16T12:47:53.751787Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2605.00195","last_updated":"2026-05-11T12:48:16Z","snapshot_observed_at":"2026-07-31T18:51:54.041588Z","submitted_at":"2026-04-30T20:20:59Z","title":"Diversity in Large Language Models under Supervised Fine-Tuning","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-09T20:32:37.788283Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2605.00195"},"observation_digest":"sha256:6096f7788b666afea325ddfaf4124901b765a737d58bd21a3889cf22c107fa09","observation_id":"9fb779b0-903a-4a0c-898c-e1821bbcfcf5","resolution":{"observed_at":"2026-05-09T20:37:32.112040Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2605.00195","last_updated":"2026-05-11T12:48:16Z","snapshot_observed_at":"2026-07-31T18:51:54.041588Z","submitted_at":"2026-04-30T20:20:59Z","title":"Diversity in Large Language Models under Supervised Fine-Tuning","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-12T03:10:22.314719Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2605.00195"},"observation_digest":"sha256:b6292bcc4ce6ba2e8083d1fbf67dde43367b470a13d60b0c7b93ec44bb055796","observation_id":"42f50f90-736f-45cc-a9cd-a95063fe62fc","resolution":{"observed_at":"2026-05-12T03:11:18.198968Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2605.08904","last_updated":"2026-05-09T11:51:34Z","snapshot_observed_at":"2026-07-06T23:21:02.177557Z","submitted_at":"2026-05-09T11:51:34Z","title":"OPT-BENCH: Evaluating the Iterative Self-Optimization of LLM Agents in Large-Scale Search Spaces","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-05-12T02:57:15.521594Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2605.08904"},"observation_digest":"sha256:20d532a4fdd2f7333d59a1be66aa1a4268628cf7a4173ae89d928e13c37e51dd","observation_id":"642a627a-3497-4036-8e18-857282d8e03d","resolution":{"observed_at":"2026-05-12T03:01:18.800129Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2605.23171","last_updated":"2026-05-22T02:43:19Z","snapshot_observed_at":"2026-08-01T19:48:27.330086Z","submitted_at":"2026-05-22T02:43:19Z","title":"Understanding and Improving Noisy Embedding Techniques in Instruction Finetuning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-25T04:51:05.359636Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2605.23171"},"observation_digest":"sha256:898822f58a845f3b1d855a3fd7ca6265efd6bb58ad021fb72174199604b1b978","observation_id":"24387372-a44b-410c-8017-87796548bf21","resolution":{"observed_at":"2026-05-25T04:55:24.092587Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.01811","last_updated":"2026-06-01T07:27:43Z","snapshot_observed_at":"2026-08-01T23:50:35.081089Z","submitted_at":"2026-06-01T07:27:43Z","title":"\"I've Seen How This Goes\": Characterizing Diversity via Progressive Conditional Surprise","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-06-28T14:51:06.448678Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.01811"},"observation_digest":"sha256:a1984a64f98f7f5a5d84169f70bdf71c643e13b4526016b7c9d93f8669b3a958","observation_id":"d819dc1b-c8ce-47a8-913b-7bdd4b79b5c3","resolution":{"observed_at":"2026-07-01T22:56:20.797161Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.03137","last_updated":"2026-07-01T08:33:01Z","snapshot_observed_at":"2026-08-05T05:03:21.120024Z","submitted_at":"2026-06-02T04:26:01Z","title":"Think-Before-Speak: From Internal Evaluation to Public Expression in Multi-Agent Social Simulation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-28T10:24:30.660372Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.03137"},"observation_digest":"sha256:ad4d9895db56949b42e22237a399eed50d76bb7a54252a3e5325c9dbcfb21f1b","observation_id":"87967254-7898-4ef0-9b2e-114ae8fcd293","resolution":{"observed_at":"2026-07-02T03:06:29.378053Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.03137","last_updated":"2026-07-01T08:33:01Z","snapshot_observed_at":"2026-08-05T05:03:21.120024Z","submitted_at":"2026-06-02T04:26:01Z","title":"Think-Before-Speak: From Internal Evaluation to Public Expression in Multi-Agent Social Simulation","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-02T23:10:03.733636Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.03137"},"observation_digest":"sha256:efc36c82474a5893caa2cdc59211ccaf6ab092407559fd63c40a9c959338f3f3","observation_id":"1fa36e31-3765-47a6-896b-81e21a2f5453","resolution":{"observed_at":"2026-07-02T23:17:29.021269Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.19714","last_updated":"2026-06-18T02:26:05Z","snapshot_observed_at":"2026-08-01T18:27:22.876701Z","submitted_at":"2026-06-18T02:26:05Z","title":"AURA: Adaptive Uncertainty-aware Refinement for LLM-as-a-Judge Auditing","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-06-26T15:48:26.303462Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.19714"},"observation_digest":"sha256:63bf8303b3b341c9f79c33a23711df11a0cb3500b917a7508a433f288c59c133","observation_id":"5b97ba62-2938-4c33-adcf-a8c4fc23d29d","resolution":{"observed_at":"2026-07-04T05:39:40.156723Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.19993","last_updated":"2026-06-18T09:31:31Z","snapshot_observed_at":"2026-08-07T06:04:20.969024Z","submitted_at":"2026-06-18T09:31:31Z","title":"Activation- and Influence-Aware Ranks (AIR): Function-Preserving SVD Compression for LLMs","version":1},"reference_index":174,"source":"arxiv_source","source_observed_at":"2026-06-26T18:09:17.414031Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.19993"},"observation_digest":"sha256:71ff4121692a5cb975418578f91f7fbea30dce220cbee3a8ee8e12b75f77e0df","observation_id":"c68d1aae-9a53-4eb8-8a41-dcb39cc5f058","resolution":{"observed_at":"2026-07-04T03:19:31.548571Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.25445","last_updated":"2026-06-24T06:15:24Z","snapshot_observed_at":"2026-07-06T23:59:52.373272Z","submitted_at":"2026-06-24T06:15:24Z","title":"C3-Bench: A Context-Aware Change Captioning Benchmark","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-25T21:02:52.529391Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.25445"},"observation_digest":"sha256:2e9c49b0f99ca12e3d8c0e694f6000d9ef3de74a967c85b847161310d5dd6432","observation_id":"004c6577-3963-4e08-ba4f-a124e17a89a7","resolution":{"observed_at":"2026-07-04T19:50:10.108934Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":"2305.14387","doi":"10.48550/arxiv.2305.14387","metadata_source":"arxiv_reference","pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yann Dubois, Xuechen Li, Rohan Taori, Tianyi Zhang, Ishaan Gulrajani, Jimmy Ba, Carlos Guestrin, Percy Liang, and Tatsunori B Hashimoto","venue":"arXiv (Cornell University)","work_id":"a875adf2-8826-466d-bf52-896ee15632ea","year":2023},"citing_paper":{"arxiv_id":"2606.26346","last_updated":"2026-06-24T19:38:21Z","snapshot_observed_at":"2026-08-08T08:43:50.725073Z","submitted_at":"2026-06-24T19:38:21Z","title":"How Do Tool-Augmented LLM Agents Perform on Real-World Energy Analytics Tasks?","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-06-26T01:34:10.103638Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2606.26346"},"observation_digest":"sha256:fe744ec24da6dca2be7e689d1d619efada34fe92f43b82ec799cab9f5512d673","observation_id":"4fbfc001-c736-41d4-ac46-390a0ee1ca24","resolution":{"observed_at":"2026-07-04T15:29:56.664257Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-07-11T13:53:36.775836Z","title":"arXiv preprint arXiv:2305.14387 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04763","last_updated":"2026-07-26T14:17:18Z","snapshot_observed_at":"2026-08-02T10:24:43.977557Z","submitted_at":"2026-07-06T07:56:53Z","title":"Multi-Turn On-Policy Distillation with Prefix Replay","version":1},"reference_index":268,"source":"arxiv_source","source_observed_at":"2026-07-11T13:53:36.775836Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2607.04763"},"observation_digest":"sha256:5ca5ddcbc1792b0bddb5b20fa3befed0d85e90e66227eb752b865bf540e3df57","observation_id":"08d8ff0c-5ae4-41e2-b896-c58ddbfc6f99","resolution":{"observed_at":"2026-07-11T13:53:36.775836Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14387","snapshot_observed_at":"2026-08-02T08:41:03.696680Z","title":"arXiv preprint arXiv:2305.14387 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.04763","last_updated":"2026-07-26T14:17:18Z","snapshot_observed_at":"2026-08-02T10:24:43.977557Z","submitted_at":"2026-07-06T07:56:53Z","title":"Multi-Turn On-Policy Distillation with Prefix Replay","version":3},"reference_index":269,"source":"arxiv_source","source_observed_at":"2026-08-02T08:41:03.696680Z"},"links":{"cited_paper":"/paper/2305.14387","citing_paper":"/paper/2607.04763"},"observation_digest":"sha256:366aba32dd41cbc195e17fcbe559c6babf153af45ef96b9a5e8bd44ecb76b971","observation_id":"7cc20e6e-6e09-4de3-ace2-c50a812b577f","resolution":{"observed_at":"2026-08-02T08:41:03.696680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2305.14387/citation-record","integrity":"/paper/2305.14387/integrity","json":"/paper/2305.14387/citation-record.json","paper":"/paper/2305.14387"},"outbound":[],"paper":{"arxiv_id":"2305.14387","last_updated":"2024-01-08T04:46:56Z","latest_version":4,"primary_category":"cs.LG","snapshot_observed_at":"2026-07-06T15:31:51.249701Z","submitted_at":"2023-05-22T17:55:50Z","title":"AlpacaFarm: A Simulation Framework for Methods that Learn from Human Feedback"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 33 inbound Pith citation observations for arXiv:2305.14387."}