{"as_of":"2026-08-13T07:10:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0a7f2a504023619b931cb2515149d44a25ca09cf2fd6e4d5e022c401818fa4e7","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":60,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":60,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":60,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":60,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T18:15:15.970037Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":2,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-12T18:15:15.970037Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.11761","last_updated":"2025-02-20T16:22:19Z","snapshot_observed_at":"2026-08-13T01:12:11.693860Z","submitted_at":"2024-11-18T17:40:42Z","title":"Mapping out the Space of Human Feedback for Reinforcement Learning: A Conceptual Framework","version":2},"reference_index":176,"source":"pdf_text","source_observed_at":"2026-08-12T18:15:15.970037Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2411.11761"},"observation_digest":"sha256:58a04da8c3ff6d6025e0981592d6d9503f9c3bc9677efe7b1c0fa127b8636a2f","observation_id":"fd3428c7-f9a8-4efb-8e9a-a8e07503d536","resolution":{"observed_at":"2026-08-12T18:15:15.970037Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-12T13:06:27.901889Z","title":"A long way to go: Investigating length correlations in rlhf","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.16502","last_updated":"2025-02-26T16:46:25Z","snapshot_observed_at":"2026-08-12T16:06:09.957131Z","submitted_at":"2024-11-25T15:37:27Z","title":"Interpreting Language Reward Models via Contrastive Explanations","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T13:06:27.901889Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2411.16502"},"observation_digest":"sha256:155af525ded7833a74900285198532c288c853454f7d978c5be0796304c54576","observation_id":"bda6611a-462d-4a5e-a4f9-bfdf8590a683","resolution":{"observed_at":"2026-08-12T13:06:27.901889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-12T12:58:30.292523Z","title":"A long way to go: Investigating length correlations in rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.16646","last_updated":"2025-02-09T07:53:38Z","snapshot_observed_at":"2026-08-13T05:48:57.895519Z","submitted_at":"2024-11-25T18:28:26Z","title":"Self-Generated Critiques Boost Reward Modeling for Language Models","version":3},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-12T12:58:30.292523Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2411.16646"},"observation_digest":"sha256:0252e1c8519c519ee9848039f716294dbf467c64a6d6c17f796ee3d3dcc9cf70","observation_id":"34899596-ca26-4c55-96fb-253fc2dd3352","resolution":{"observed_at":"2026-08-12T12:58:30.292523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-11T14:40:29.387244Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.11803","last_updated":"2025-05-23T03:05:52Z","snapshot_observed_at":"2026-08-12T00:18:17.505630Z","submitted_at":"2024-12-16T14:14:27Z","title":"UAlign: Leveraging Uncertainty Estimations for Factuality Alignment on Large Language Models","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-11T14:40:29.387244Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2412.11803"},"observation_digest":"sha256:f35124e192ca00d7e77867a9681ce5a6994cf40e8e5705b0473d1e85e6599b5b","observation_id":"022c3f76-94d8-4d24-baa9-1edd7347be15","resolution":{"observed_at":"2026-08-11T14:40:29.387244Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-11T10:41:24.069209Z","title":"A long way to go: Investigating length correlations in rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.16475","last_updated":"2024-12-21T04:07:17Z","snapshot_observed_at":"2026-08-13T04:40:46.902476Z","submitted_at":"2024-12-21T04:07:17Z","title":"When Can Proxies Improve the Sample Complexity of Preference Learning?","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-11T10:41:24.069209Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2412.16475"},"observation_digest":"sha256:4148f91bf25ecffcba0faa9846aba0c48d42410807fca91956fb93ed6e55f5cf","observation_id":"699e0a8f-d3cd-489a-b6e3-dfbdf36a16e7","resolution":{"observed_at":"2026-08-11T10:41:24.069209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-10T22:26:46.616014Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.01668","last_updated":"2025-06-14T09:58:51Z","snapshot_observed_at":"2026-08-11T15:59:54.184603Z","submitted_at":"2025-01-03T06:50:06Z","title":"CoT-based Synthesizer: Enhancing LLM Performance through Answer Synthesis","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-10T22:26:46.616014Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2501.01668"},"observation_digest":"sha256:f7f0e5eff0fc29a13c95ace11500303b35e2fd72ca1a34f221453378d2c48e5c","observation_id":"fb2ff56d-3b00-4d5a-b7bf-0229e9336553","resolution":{"observed_at":"2026-08-10T22:26:46.616014Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-10T21:24:10.614544Z","title":"A long way to go: Investigating length correlations in rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.04952","last_updated":"2025-01-09T03:59:10Z","snapshot_observed_at":"2026-08-10T21:19:04.243505Z","submitted_at":"2025-01-09T03:59:10Z","title":"Open Problems in Machine Unlearning for AI Safety","version":1},"reference_index":120,"source":"arxiv_source","source_observed_at":"2026-08-10T21:24:10.614544Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2501.04952"},"observation_digest":"sha256:c0d75aa753d950b9c855e8bb44eb769b611631e95d05683bcf76a326327cd686","observation_id":"7baf5ce8-b921-4106-9d79-cf2c0f44a7b0","resolution":{"observed_at":"2026-08-10T21:24:10.614544Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-10T21:31:45.685491Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.06248","last_updated":"2025-02-25T18:04:50Z","snapshot_observed_at":"2026-08-10T23:14:12.533178Z","submitted_at":"2025-01-08T19:03:17Z","title":"Utility-inspired Reward Transformations Improve Reinforcement Learning Training of Language Models","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-10T21:31:45.685491Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2501.06248"},"observation_digest":"sha256:9bd9433b4d6e5ac2065b42f3312a47004bcc1d378794db1e90b0e210898fa153","observation_id":"571b9802-8a55-474e-81c0-c2e5a4ec0f20","resolution":{"observed_at":"2026-08-10T21:31:45.685491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-10T19:55:50.747644Z","title":"A long way to go: Investigating length correlations in rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09620","last_updated":"2025-05-29T02:21:03Z","snapshot_observed_at":"2026-08-13T04:40:45.853560Z","submitted_at":"2025-01-16T16:00:37Z","title":"Beyond Reward Hacking: Causal Rewards for Large Language Model Alignment","version":2},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-10T19:55:50.747644Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2501.09620"},"observation_digest":"sha256:3575c00e8966d325f1016713de08ef1559f2ba1c1bc93075aa14837db5a8e45e","observation_id":"9d594349-e429-4e62-a5f7-ff968947da6e","resolution":{"observed_at":"2026-08-10T19:55:50.747644Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-10T16:58:02.583705Z","title":"A long way to go: Investigating length correlations in rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.12735","last_updated":"2025-02-07T02:13:27Z","snapshot_observed_at":"2026-08-10T16:48:54.578052Z","submitted_at":"2025-01-22T09:12:09Z","title":"Online Preference Alignment for Language Models via Count-based Exploration","version":3},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-10T16:58:02.583705Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2501.12735"},"observation_digest":"sha256:b807ed71c1cee0aff46db22347c69a8637f3e0b0f2b00c5cb6aff41ad18d77ae","observation_id":"4974abf7-e451-4b7c-b3a2-52f85b3ce142","resolution":{"observed_at":"2026-08-10T16:58:02.583705Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-09T17:46:29.272056Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.00814","last_updated":"2025-05-19T08:24:51Z","snapshot_observed_at":"2026-08-12T21:19:56.699829Z","submitted_at":"2025-02-02T14:50:25Z","title":"Disentangling Length Bias In Preference Learning Via Response-Conditioned Modeling","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-09T17:46:29.272056Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2502.00814"},"observation_digest":"sha256:96bd8a04d4342609efb882ddbd302afcf56a1699e4bf9bd7aa99f6f559f90eca","observation_id":"5e01346b-9df2-4ca3-b2cd-0c293fc45593","resolution":{"observed_at":"2026-08-09T17:46:29.272056Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-07T23:18:40.210360Z","title":"A long way to go: Investigating length correlations in rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.08922","last_updated":"2025-02-13T03:15:31Z","snapshot_observed_at":"2026-08-11T00:52:33.304474Z","submitted_at":"2025-02-13T03:15:31Z","title":"Self-Consistency of the Internal Reward Models Improves Self-Rewarding Language Models","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T23:18:40.210360Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2502.08922"},"observation_digest":"sha256:41a67727454aa1b2b2139161ef372d8254e495469a027739b3f8faa3779aa878","observation_id":"7f7c52d3-b3d6-435b-9a49-90fde801f122","resolution":{"observed_at":"2026-08-07T23:18:40.210360Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-07T21:56:20.465030Z","title":"A long way to go: Investi- gating length correlations in rlhf","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.09328","last_updated":"2025-02-13T13:40:52Z","snapshot_observed_at":"2026-08-07T21:49:22.645033Z","submitted_at":"2025-02-13T13:40:52Z","title":"Copilot Arena: A Platform for Code LLM Evaluation in the Wild","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T21:56:20.465030Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2502.09328"},"observation_digest":"sha256:3c8bcfa6aed45f36fe85c28b596f6be85a673399a8cf4eef0a13d5284ac00cb8","observation_id":"b5e176e0-f468-4289-ba89-d9ffca8f0654","resolution":{"observed_at":"2026-08-07T21:56:20.465030Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-07T18:32:21.215175Z","title":"A long way to go: Investigating length correlations in rlhf, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.10505","last_updated":"2025-07-26T17:40:08Z","snapshot_observed_at":"2026-08-11T21:30:34.811222Z","submitted_at":"2025-02-14T19:01:34Z","title":"Preference learning made easy: Everything should be understood through win rate","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T18:32:21.215175Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2502.10505"},"observation_digest":"sha256:365137ab1585c2d2234be67608bc72e9bb0f3598fca6bbb4876a02331dc1f7aa","observation_id":"823f03aa-ab4a-4e29-9d59-69a72139f1aa","resolution":{"observed_at":"2026-08-07T18:32:21.215175Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2503.04697","last_updated":"2025-10-03T01:55:58Z","snapshot_observed_at":"2026-08-06T08:53:09.095000Z","submitted_at":"2025-03-06T18:43:29Z","title":"L1: Controlling How Long A Reasoning Model Thinks With Reinforcement Learning","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-18T00:19:22.140009Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2503.04697"},"observation_digest":"sha256:5b642088c3f3ef5acf912e82cd4832e223e1809007d928e58d007a580f788edd","observation_id":"b57e7a6a-64d3-4f8a-b014-b0291c7ecc5c","resolution":{"observed_at":"2026-05-18T00:19:22.253420Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2504.12501","last_updated":"2026-08-03T01:47:58Z","snapshot_observed_at":"2026-08-07T16:01:39.307745Z","submitted_at":"2025-04-16T21:36:46Z","title":"Reinforcement Learning from Human Feedback","version":9},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-22T19:27:40.991325Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2504.12501"},"observation_digest":"sha256:c383413b4893f5ba0317ab72a549dc78d3c0d088cb0ff96789ededb34114c595","observation_id":"bba7c405-e8ce-4887-a676-306ec8f27e85","resolution":{"observed_at":"2026-05-22T19:32:01.260438Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-07T14:57:31.186800Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.16869","last_updated":"2025-05-22T16:24:51Z","snapshot_observed_at":"2026-08-09T02:24:14.170168Z","submitted_at":"2025-05-22T16:24:51Z","title":"MPO: Multilingual Safety Alignment via Reward Gap Optimization","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-07T14:57:31.186800Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2505.16869"},"observation_digest":"sha256:d1536d4c38a2867faa2e70b0f9557005b42255553d74d9c20ab30fde9985db18","observation_id":"6b1f4f8d-a4ec-484e-8542-d04f3331f5f3","resolution":{"observed_at":"2026-08-07T14:57:31.186800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-07T14:05:12.431299Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.20088","last_updated":"2025-05-29T15:47:53Z","snapshot_observed_at":"2026-08-10T20:48:13.954759Z","submitted_at":"2025-05-26T15:01:56Z","title":"Multi-Domain Explainability of Preferences","version":2},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-07T14:05:12.431299Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2505.20088"},"observation_digest":"sha256:b5150f156db373d35200dd60694aceefb16f49dd045218da6f0cd44ea0081e35","observation_id":"c90c595a-7ad6-4f05-85fa-4eaec3bd05a3","resolution":{"observed_at":"2026-08-07T14:05:12.431299Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-07T06:05:32.676825Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.11098","last_updated":"2025-06-06T13:19:07Z","snapshot_observed_at":"2026-08-08T09:29:26.176523Z","submitted_at":"2025-06-06T13:19:07Z","title":"Debiasing Online Preference Learning via Preference Feature Preservation","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T06:05:32.676825Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2506.11098"},"observation_digest":"sha256:27eca54691567bdfd46c483c961f06dbc94a8ad9c85b80c1168b5412a4fe30ce","observation_id":"a3844b5c-fc83-4018-b6a9-5f94054c4f26","resolution":{"observed_at":"2026-08-07T06:05:32.676825Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-07T00:59:13.509910Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.12307","last_updated":"2025-06-20T01:43:46Z","snapshot_observed_at":"2026-08-07T20:54:59.398635Z","submitted_at":"2025-06-14T02:00:36Z","title":"Med-U1: Incentivizing Unified Medical Reasoning in LLMs via Large-scale Reinforcement Learning","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T00:59:13.509910Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2506.12307"},"observation_digest":"sha256:4f7344e22df2d2b6f49df877dac9fffa53d2af38b45ee89b3d6f5157f0610aaa","observation_id":"5b442089-496d-4b5a-98ed-c8938c82a44c","resolution":{"observed_at":"2026-08-07T00:59:13.509910Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-06T23:00:45.978822Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.20160","last_updated":"2025-08-08T09:23:42Z","snapshot_observed_at":"2026-08-09T21:06:17.179483Z","submitted_at":"2025-06-25T06:29:18Z","title":"AALC: Large Language Model Efficient Reasoning via Adaptive Accuracy-Length Control","version":2},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-06T23:00:45.978822Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2506.20160"},"observation_digest":"sha256:cfa12b5b61d5bea924311adbd55d5e6fad058043f5c7131eb575091b92195c20","observation_id":"73af5cdd-1b0c-49a1-ad4f-8d4ee91000e7","resolution":{"observed_at":"2026-08-06T23:00:45.978822Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2507.15698","last_updated":"2026-05-19T07:11:43Z","snapshot_observed_at":"2026-08-12T06:53:50.899254Z","submitted_at":"2025-07-21T15:07:59Z","title":"CoLD: Counterfactually-Guided Length Debiasing for Process Reward Models in Mathematical Reasoning","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-21T23:24:43.556606Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2507.15698"},"observation_digest":"sha256:eff7a318f230447ef656388d4792bec502e34378c22c1ff8d2a3ffb0068e82e3","observation_id":"24e63f3c-ad87-44a4-b51f-3942077965e3","resolution":{"observed_at":"2026-05-21T23:25:45.297132Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2507.17746","last_updated":"2025-10-03T01:55:55Z","snapshot_observed_at":"2026-08-12T14:23:22.826603Z","submitted_at":"2025-07-23T17:57:55Z","title":"Rubrics as Rewards: Reinforcement Learning Beyond Verifiable Domains","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-13T06:07:56.678339Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2507.17746"},"observation_digest":"sha256:9e6f84364069d4ad208e869f24fcd997834cb31a974c144fca7fd09f2d12f6f8","observation_id":"5536e894-0fde-41a3-a091-46ab96deea3a","resolution":{"observed_at":"2026-05-13T06:07:56.839208Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2601.21350","last_updated":"2026-05-16T02:45:10Z","snapshot_observed_at":"2026-07-06T22:43:26.250071Z","submitted_at":"2026-01-29T07:18:45Z","title":"Factored Causal Representation Learning for Robust Reward Modeling in RLHF","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-21T14:18:33.768962Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2601.21350"},"observation_digest":"sha256:680b45a8f631d8670afd99dac8ed6229bc057efea8ae7585facae342763a8377","observation_id":"44618efe-8da3-4200-a188-dd2b0be30815","resolution":{"observed_at":"2026-05-21T14:20:13.606475Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-03T03:04:44.349733Z","title":"A long way to go: Investigating length correlations in rlhf.arXiv preprint arXiv:2310.03716,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.09305","last_updated":"2026-06-28T18:13:37Z","snapshot_observed_at":"2026-08-12T17:14:29.445906Z","submitted_at":"2026-02-10T00:45:24Z","title":"Reward Modeling for Reinforcement Learning-Based LLM Reasoning: Design, Challenges, and Evaluation","version":2},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-03T03:04:44.349733Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2602.09305"},"observation_digest":"sha256:bc1ec630915e5fa900f85b15e24d847f65b35cd8a99f534a7c6b8fea10703b6c","observation_id":"06669edc-9c7d-46b5-a87f-e90a72d29701","resolution":{"observed_at":"2026-08-03T03:04:44.349733Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-03T01:06:24.896678Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2602.10623","last_updated":"2026-06-01T09:37:25Z","snapshot_observed_at":"2026-08-03T01:06:22.800179Z","submitted_at":"2026-02-11T08:14:11Z","title":"Mitigating Reward Hacking in RLHF via Bayesian Non-negative Reward Modeling","version":2},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-03T01:06:24.896678Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2602.10623"},"observation_digest":"sha256:062ee64fd8b15f7db9582f205099a81bfaebfe36cc9a3c80b11c39abde3f2b26","observation_id":"4182d5ab-7695-4523-805c-a384433e185d","resolution":{"observed_at":"2026-08-03T01:06:24.896678Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2602.22710","last_updated":"2026-05-06T21:24:15Z","snapshot_observed_at":"2026-08-11T13:34:15.111919Z","submitted_at":"2026-02-26T07:34:15Z","title":"Same Words, Different Judgments: How Preferences Vary Across Modalities","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-15T19:31:51.996480Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2602.22710"},"observation_digest":"sha256:2d1b9f329b9831c2a73f2ae4f33f20b014da9139126d1e8c37345c5fda3eba92","observation_id":"7e485b73-f283-4343-adc9-cc9062d905b9","resolution":{"observed_at":"2026-05-15T19:36:33.025808Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-07-13T21:26:43.149095Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2603.20510","last_updated":"2026-06-23T04:29:10Z","snapshot_observed_at":"2026-08-11T12:46:39.330679Z","submitted_at":"2026-03-20T21:24:28Z","title":"Grounded Chess Reasoning in Language Models via Master Distillation","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-07-13T21:26:43.149095Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2603.20510"},"observation_digest":"sha256:6d7ed6fc965a865a4472f2aa977c6ee6138c475d88e9a0631773bda1c417b6c4","observation_id":"7d7e801c-58d9-4e73-828a-34b0499054ed","resolution":{"observed_at":"2026-07-13T21:26:43.149095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2604.02686","last_updated":"2026-04-03T03:30:34Z","snapshot_observed_at":"2026-08-02T18:52:24.780545Z","submitted_at":"2026-04-03T03:30:34Z","title":"Beyond Semantic Manipulation: Token-Space Attacks on Reward Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-13T20:29:31.354743Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2604.02686"},"observation_digest":"sha256:679ee5c52fdf577df1000665db4d92a7cfe54a192dfbe4a9418c32ab660ec76d","observation_id":"13e706f7-5350-42c3-acef-a4f37a21c5d3","resolution":{"observed_at":"2026-05-13T20:33:16.846395Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2604.13602","last_updated":"2026-04-15T08:11:34Z","snapshot_observed_at":"2026-08-11T03:30:51.869417Z","submitted_at":"2026-04-15T08:11:34Z","title":"Reward Hacking in the Era of Large Models: Mechanisms, Emergent Misalignment, Challenges","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-10T13:58:53.430492Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2604.13602"},"observation_digest":"sha256:58d20e71f6a3e131e45e19ae2c45e4729dc0a9d202b0e550706eb939e42ff6f0","observation_id":"fb6bad19-b6c0-470f-90d7-156a88a82d12","resolution":{"observed_at":"2026-05-10T14:00:28.520666Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2604.13833","last_updated":"2026-04-16T05:28:11Z","snapshot_observed_at":"2026-08-11T14:46:48.035861Z","submitted_at":"2026-04-15T13:07:11Z","title":"Robust Reward Modeling for Large Language Models via Causal Decomposition","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T13:50:47.902911Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2604.13833"},"observation_digest":"sha256:347ce56662be8914e98aaa5287979c0a63dbdab69855c6cc0c71e647cd0c42bc","observation_id":"f5703def-dd90-49d8-8885-be064697dd6a","resolution":{"observed_at":"2026-05-10T14:00:29.614432Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2604.17328","last_updated":"2026-05-23T14:02:35Z","snapshot_observed_at":"2026-08-11T12:32:52.781184Z","submitted_at":"2026-04-19T08:48:46Z","title":"Rethinking the Comparison Unit in Sequence-Level Reinforcement Learning: An Equal-Length Paired Training Framework from Loss Correction to Sample Construction","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-05-10T06:23:14.905706Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2604.17328"},"observation_digest":"sha256:9ce91bcf9a1dce75b4c929bd77b963616a0f1655739ab6509e2ffeb7ebef9cc8","observation_id":"247c5220-c7e8-4d73-8e78-925f1bbfb5f3","resolution":{"observed_at":"2026-05-10T06:26:27.556587Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2604.21911","last_updated":"2026-04-23T17:54:36Z","snapshot_observed_at":"2026-08-11T07:45:06.895919Z","submitted_at":"2026-04-23T17:54:36Z","title":"When Prompts Override Vision: Prompt-Induced Hallucinations in LVLMs","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-09T22:11:40.905814Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2604.21911"},"observation_digest":"sha256:2dfc31b3ee2e80932958605674a8acd88b4a6592bee5eb3738be58aae9c07fb8","observation_id":"df56be47-4138-4850-ae37-4b4363f25a25","resolution":{"observed_at":"2026-05-09T22:54:16.716341Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2605.09269","last_updated":"2026-05-10T02:32:19Z","snapshot_observed_at":"2026-08-10T23:31:48.440468Z","submitted_at":"2026-05-10T02:32:19Z","title":"DeltaRubric: Generative Multimodal Reward Modeling via Joint Planning and Verification","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-12T04:41:44.833354Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2605.09269"},"observation_digest":"sha256:9f59c77c937fa6a5e739d59a97c695d6ab60d16895ef41820d662fbc0ee9e919","observation_id":"da880bd9-6413-4078-b427-eb13109afe13","resolution":{"observed_at":"2026-05-12T06:01:24.374651Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2605.11134","last_updated":"2026-05-29T17:16:57Z","snapshot_observed_at":"2026-08-11T09:54:15.520173Z","submitted_at":"2026-05-11T18:41:12Z","title":"Spurious Correlation Learning in Preference Optimization: Mechanisms, Consequences, and Mitigation via Tie Training","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-13T06:30:51.812541Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2605.11134"},"observation_digest":"sha256:6ea48784d8e53dbaeadb3929757b22103d9a8950988ee8ced581fc5607f7a6b1","observation_id":"d4a8e0f9-5531-4d4e-b214-f3b43b419a14","resolution":{"observed_at":"2026-05-13T06:32:24.323187Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2605.15207","last_updated":"2026-07-08T14:36:15Z","snapshot_observed_at":"2026-07-12T17:52:34.996306Z","submitted_at":"2026-05-01T23:42:57Z","title":"TeamTR: Trust-Region Fine-Tuning for Multi-Agent LLM Coordination","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-05-19T18:01:06.649723Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2605.15207"},"observation_digest":"sha256:bec2749295f8c6992aa76511cdd1f19cfa509d0b98ce840415ca13f9f264e409","observation_id":"793a6de8-4a36-483c-8c96-0eca0a88efe2","resolution":{"observed_at":"2026-05-19T18:02:42.351806Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2605.16339","last_updated":"2026-05-07T16:48:48Z","snapshot_observed_at":"2026-08-12T12:06:34.294930Z","submitted_at":"2026-05-07T16:48:48Z","title":"Preference Instability in Reward Models: Detection and Mitigation via Sparse Autoencoders","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-20T22:48:54.238767Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2605.16339"},"observation_digest":"sha256:7842c9ccedbeafeeeafad463b477c85d4d69c7a084c9f2551f93640d5dc34525","observation_id":"aacc4150-6bc4-421e-8d49-e55e0eda8a76","resolution":{"observed_at":"2026-05-20T22:49:10.113914Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2605.18721","last_updated":"2026-05-21T04:23:01Z","snapshot_observed_at":"2026-07-06T23:29:33.702647Z","submitted_at":"2026-05-18T17:50:27Z","title":"General Preference Reinforcement Learning","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-20T12:43:56.522345Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2605.18721"},"observation_digest":"sha256:afa69672f379dca8c34a239f75e90d09fa9f2e68946e9542f926cbe79e278051","observation_id":"ba470ddb-b949-421d-9f72-ae3edd45db43","resolution":{"observed_at":"2026-05-20T12:48:17.700027Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2605.18721","last_updated":"2026-05-21T04:23:01Z","snapshot_observed_at":"2026-07-06T23:29:33.702647Z","submitted_at":"2026-05-18T17:50:27Z","title":"General Preference Reinforcement Learning","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-21T07:50:00.963837Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2605.18721"},"observation_digest":"sha256:8bdbba5c9daa4483f80ef0685d54c56995e80c86665996f594d4c53dc34832b6","observation_id":"76522607-8b4f-426f-a9a5-41d61fd44997","resolution":{"observed_at":"2026-05-21T07:54:02.903499Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2605.18721","last_updated":"2026-05-21T04:23:01Z","snapshot_observed_at":"2026-07-06T23:29:33.702647Z","submitted_at":"2026-05-18T17:50:27Z","title":"General Preference Reinforcement Learning","version":3},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-22T09:24:39.228616Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2605.18721"},"observation_digest":"sha256:c39f37b0eeeddceb467251c00ec9e83abb7c321b7c6504ba4302a3d542cd5302","observation_id":"860c78bb-120c-4aaa-ab83-a06d8f332584","resolution":{"observed_at":"2026-05-22T09:24:45.692068Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2605.30844","last_updated":"2026-05-29T05:05:01Z","snapshot_observed_at":"2026-07-06T23:40:03.151978Z","submitted_at":"2026-05-29T05:05:01Z","title":"Fine-Tuning Improves Information Conveyance in Language Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-28T22:58:27.094125Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2605.30844"},"observation_digest":"sha256:b843c73ea237e0b6885cbcdc4998f523e3c9399182a00cfa8cb4c5ae7e2aad00","observation_id":"9c15c34c-e713-4b2c-869e-825dca8795c8","resolution":{"observed_at":"2026-06-28T23:02:46.635308Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.03131","last_updated":"2026-06-02T04:18:08Z","snapshot_observed_at":"2026-08-02T20:05:30.693121Z","submitted_at":"2026-06-02T04:18:08Z","title":"HARVE: Hacking-Aware Reward-Head Vector Editing for Robust Reward Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-28T11:32:16.166724Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.03131"},"observation_digest":"sha256:a182feb6bfc3e32c6d81d3a1305d2cdccf1db9d545c907bb79925c4402b5da45","observation_id":"24dd6509-284b-4608-be64-273d1da92917","resolution":{"observed_at":"2026-07-02T01:46:26.825310Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.04075","last_updated":"2026-06-18T13:14:32Z","snapshot_observed_at":"2026-08-02T22:57:00.184818Z","submitted_at":"2026-06-02T16:29:48Z","title":"Large Language Models Hack Rewards, and Society","version":2},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-06-28T11:30:35.285902Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.04075"},"observation_digest":"sha256:bf237da06947db4fede0f697c589447426ebf4e97713c2bba54c0088a1b2733f","observation_id":"8624b820-36b4-4fcf-8606-1021e9b205ae","resolution":{"observed_at":"2026-07-02T01:46:26.991887Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.04273","last_updated":"2026-08-07T18:52:22Z","snapshot_observed_at":"2026-08-13T06:10:36.387429Z","submitted_at":"2026-06-02T22:58:19Z","title":"Human agency in initial human-AI proof formalization workflows","version":1},"reference_index":297,"source":"arxiv_source","source_observed_at":"2026-06-28T09:29:50.282874Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.04273"},"observation_digest":"sha256:6443c545c34ed80f61ee32a0ffad7683f71b173d3adee20063184af20f683d7f","observation_id":"fba79bae-8cef-4297-9bb8-9c4fe7ed7e51","resolution":{"observed_at":"2026-07-02T04:06:34.867868Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.04781","last_updated":"2026-06-03T12:02:49Z","snapshot_observed_at":"2026-08-12T14:43:39.047557Z","submitted_at":"2026-06-03T12:02:49Z","title":"AIP: A Graph Representation for Learning and Governing Agent Skills","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-28T05:58:37.538980Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.04781"},"observation_digest":"sha256:9e37e7dba6b1d46f428981379f5cb979a9b8de348f0a7dc3ac9b05b981dffc26","observation_id":"e8546286-d4ec-4628-9077-4670da119ae9","resolution":{"observed_at":"2026-07-02T08:36:47.772402Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.05054","last_updated":"2026-06-03T16:12:30Z","snapshot_observed_at":"2026-08-12T12:36:47.059970Z","submitted_at":"2026-06-03T16:12:30Z","title":"Boosting Self-Consistency with Ranking","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-06-28T06:49:58.051659Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.05054"},"observation_digest":"sha256:4773299a69d0425c9ccd4e865716eea043ec534634b445fe4d14944ff8cd265d","observation_id":"4322c478-5d0e-4084-9a60-68818d879b10","resolution":{"observed_at":"2026-06-28T06:51:44.326365Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.07988","last_updated":"2026-06-06T05:35:31Z","snapshot_observed_at":"2026-07-06T23:47:33.680573Z","submitted_at":"2026-06-06T05:35:31Z","title":"PAFO: Pareto Fairness Optimization for Personalized Reward Modeling","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-27T20:00:05.900814Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.07988"},"observation_digest":"sha256:12a2fa3549d3554c017f5ba9c632491a4f7013d9c0c315c70a71585b1dedf7dd","observation_id":"27d1cab7-fd2c-4458-88ae-476ae13a737a","resolution":{"observed_at":"2026-07-02T20:57:23.623629Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.09073","last_updated":"2026-06-10T20:16:58Z","snapshot_observed_at":"2026-08-13T02:39:24.411779Z","submitted_at":"2026-06-08T06:15:30Z","title":"A Unifying Lens on Reward Uncertainty in RLHF","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-27T17:11:46.549150Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.09073"},"observation_digest":"sha256:7c6062a1aac17cfd6d7f66559b4b2efe0445dff27e44c07cb2f4d807de595b36","observation_id":"0733ceff-fd0a-4cca-af74-700a880fc8c7","resolution":{"observed_at":"2026-07-03T00:27:29.828062Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.09711","last_updated":"2026-06-08T16:32:54Z","snapshot_observed_at":"2026-07-06T23:49:03.237958Z","submitted_at":"2026-06-08T16:32:54Z","title":"Proxy Reward Internalization and Mechanistic Exploitation: A Learned Precursor to Reward Hacking and Its Generalization","version":1},"reference_index":263,"source":"arxiv_source","source_observed_at":"2026-06-27T16:26:34.918099Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.09711"},"observation_digest":"sha256:05f1436e0b7b9c3e350c89ba5f125d25e8b1c96581d10f180c6b3e85dab96297","observation_id":"8eaf8cee-75bb-42df-bc50-f16fd6e34f46","resolution":{"observed_at":"2026-07-03T01:37:30.510079Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.11635","last_updated":"2026-06-10T03:56:07Z","snapshot_observed_at":"2026-08-08T17:34:26.394532Z","submitted_at":"2026-06-10T03:56:07Z","title":"Are LLMs Bad at Moral Reasoning?","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-27T08:20:24.251540Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.11635"},"observation_digest":"sha256:78a438f1fabc2ffb383d32e7cfcd2a2366e5d572095724e04253aa1270333a22","observation_id":"567ca9a4-00b3-4f25-bd9d-3914b2a0ac2b","resolution":{"observed_at":"2026-07-03T13:18:12.661096Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.19057","last_updated":"2026-06-17T13:26:04Z","snapshot_observed_at":"2026-08-12T00:53:29.964214Z","submitted_at":"2026-06-17T13:26:04Z","title":"Quantifying and Auditing LLM Evaluation via Positive--Unlabeled Learning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-26T19:04:45.062426Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.19057"},"observation_digest":"sha256:63a49d55c29bc7f533a828c22fd89f3736f0a7d398a63d2f102e9cfd16d0f966","observation_id":"879baf7e-c75c-44cb-ac42-e2b1b5283fb9","resolution":{"observed_at":"2026-07-04T02:49:24.830653Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.21732","last_updated":"2026-06-19T20:43:14Z","snapshot_observed_at":"2026-08-11T08:34:07.751766Z","submitted_at":"2026-06-19T20:43:14Z","title":"Safe to Check, Unsafe to Use: Relinking at the Compression Boundary of LLM Agents","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-06-26T13:23:54.784902Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.21732"},"observation_digest":"sha256:3e123159d174f4e76036ca8f3130f5cdf440f5d4979b3dbbcf6b373b3b58c678","observation_id":"018f614d-34df-478e-9f14-e90945783695","resolution":{"observed_at":"2026-07-04T07:29:39.123429Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.22329","last_updated":"2026-06-21T04:35:43Z","snapshot_observed_at":"2026-07-31T23:50:31.635694Z","submitted_at":"2026-06-21T04:35:43Z","title":"BabelJudge: Measuring LLM-as-a-Judge Reliability Across Languages and Agent Trajectories","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-26T11:06:07.303335Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.22329"},"observation_digest":"sha256:a948d9ffa78834a8d13617e0f4ee12904b7f84d75534102205fc0ed5e62ec585","observation_id":"9c2fe608-fda4-420c-983c-7e2db25faa88","resolution":{"observed_at":"2026-07-04T08:49:41.493565Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.25432","last_updated":"2026-06-29T19:08:38Z","snapshot_observed_at":"2026-08-13T06:46:05.768847Z","submitted_at":"2026-06-24T05:50:28Z","title":"Brevity is the Soul of Inference Efficiency: Inducing Concision in VLMs via Data Curation","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-06-25T21:05:36.836361Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.25432"},"observation_digest":"sha256:9020d0213b95ccb0804e60f0c8e8a5c737bf2d2e07050b15bf36e76f19fb60aa","observation_id":"83f3ef1f-c4fa-4943-8d0e-a742b4e7142b","resolution":{"observed_at":"2026-07-04T19:40:07.164699Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2606.25432","last_updated":"2026-06-29T19:08:38Z","snapshot_observed_at":"2026-08-13T06:46:05.768847Z","submitted_at":"2026-06-24T05:50:28Z","title":"Brevity is the Soul of Inference Efficiency: Inducing Concision in VLMs via Data Curation","version":2},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-07-01T06:30:27.178950Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2606.25432"},"observation_digest":"sha256:2574312574bd396e529b85e60c8bbeb84aeb5c7f1ffd941aadfa6791132ba613","observation_id":"b4ceae58-6152-4a00-b8e1-c7389e305a3a","resolution":{"observed_at":"2026-07-01T09:35:39.580781Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-07-11T16:47:52.768236Z","title":"A long way to go: Investigating length correlations in RLHF","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.04590","last_updated":"2026-07-06T01:33:12Z","snapshot_observed_at":"2026-08-05T13:00:19.375765Z","submitted_at":"2026-07-06T01:33:12Z","title":"Attention Limited Reward Learning","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-07-11T16:47:52.768236Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2607.04590"},"observation_digest":"sha256:da2d26d2a76eda4842a7b918b31d5ae1e0e8dc8c37ef0e73eff8b2d81b38e137","observation_id":"03d1d8a0-9e02-4791-b5f4-f5df460b8a94","resolution":{"observed_at":"2026-07-11T16:47:52.768236Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":"2310.03716","doi":"10.48550/arxiv.2310.03716","metadata_source":"pith","pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"A long way to go: Investigating length correlations in rlhf","venue":"cs.CL","work_id":"23d0dfe6-c919-4498-8696-1e298c7834e9","year":2023},"citing_paper":{"arxiv_id":"2607.05904","last_updated":"2026-07-07T06:59:30Z","snapshot_observed_at":"2026-08-13T05:58:01.677242Z","submitted_at":"2026-07-07T06:59:30Z","title":"More Convincing, Not More Correct: Self-Play Reward Hacking of Reference-Free LLM Judges","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-08T21:20:07.569587Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2607.05904"},"observation_digest":"sha256:0070ac2a24d6619fe2fe982696ffef01227a9eb07baffbdec7cbcc3f1cad0e8d","observation_id":"fad3a143-a4bd-4dc4-9ba6-2eac61ce336e","resolution":{"observed_at":"2026-07-08T21:25:38.638530Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:49:21.130019+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-01T15:17:58.814427Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.18508","last_updated":"2026-07-20T21:05:08Z","snapshot_observed_at":"2026-08-11T17:34:00.607706Z","submitted_at":"2026-07-20T21:05:08Z","title":"Style over Substance: A Shortcut Audit of Emotion-Description Preference Evaluation","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-01T15:17:58.814427Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2607.18508"},"observation_digest":"sha256:2bb2ff4c94b14ffec7ad67d1d68d984d271875aee383d6f64149ab1ff3f99f6b","observation_id":"3b333ee2-fdde-45d5-8f8d-ccbc3e91ba2d","resolution":{"observed_at":"2026-08-01T15:17:58.814427Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-06T16:45:20.503770Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.04783","last_updated":"2026-08-06T09:16:21Z","snapshot_observed_at":"2026-08-12T20:13:23.820545Z","submitted_at":"2026-08-05T12:49:36Z","title":"RepoProbe: Benchmarking Architecture-Aware Repository Comprehension with Checklists","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T16:45:20.503770Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2608.04783"},"observation_digest":"sha256:9bbc9c455b344223173c30a806d22d48edcf26220c8f4493477cc172e24f9520","observation_id":"c1e5f632-57cd-45c6-84be-0c6db6879aa0","resolution":{"observed_at":"2026-08-06T16:45:20.503770Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03716","snapshot_observed_at":"2026-08-08T17:43:23.532685Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.04783","last_updated":"2026-08-06T09:16:21Z","snapshot_observed_at":"2026-08-12T20:13:23.820545Z","submitted_at":"2026-08-05T12:49:36Z","title":"RepoProbe: Benchmarking Architecture-Aware Repository Comprehension with Checklists","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-08T17:43:23.532685Z"},"links":{"cited_paper":"/paper/2310.03716","citing_paper":"/paper/2608.04783"},"observation_digest":"sha256:e755478a595d3b3d121cf5a2806583b31bb8e7158397e6d781a83dee9689d472","observation_id":"b2cbe56d-c148-4c65-b989-9cd4ecd9fb20","resolution":{"observed_at":"2026-08-08T17:43:23.532685Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2310.03716/citation-record","integrity":"/paper/2310.03716/integrity","json":"/paper/2310.03716/citation-record.json","paper":"/paper/2310.03716"},"outbound":[],"paper":{"arxiv_id":"2310.03716","last_updated":"2024-07-10T23:15:49Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-13T05:56:00.586159Z","submitted_at":"2023-10-05T17:38:28Z","title":"A Long Way to Go: Investigating Length Correlations in RLHF"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 60 inbound Pith citation observations for arXiv:2310.03716."}