{"as_of":"2026-08-14T23:16:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:390b24e58a15e94f178148ca3c6b9ac32ebfc7d9beaad2ad395a60f233f9dca5","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":31,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":31,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":31,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":31,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T20:31:53.198146Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T13:49:52.398618Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-12T20:31:53.198146Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.09661","last_updated":"2024-11-14T18:31:39Z","snapshot_observed_at":"2026-08-14T09:19:14.315436Z","submitted_at":"2024-11-14T18:31:39Z","title":"Adaptive Decoding via Latent Preference Optimization","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-12T20:31:53.198146Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2411.09661"},"observation_digest":"sha256:ab97afccea89e3533e723a90be2082ea7a7ed802bd30e83b5dffef5335aa3768","observation_id":"65ddd036-d75f-47de-afe4-99952039e0d8","resolution":{"observed_at":"2026-08-12T20:31:53.198146Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-12T17:15:47.933909Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.12843","last_updated":"2024-11-19T20:17:04Z","snapshot_observed_at":"2026-08-12T19:56:50.710873Z","submitted_at":"2024-11-19T20:17:04Z","title":"Reward Modeling with Ordinal Feedback: Wisdom of the Crowd","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-12T17:15:47.933909Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2411.12843"},"observation_digest":"sha256:d38ded98731d4d1d7461b9c3b113b561839e5efc896e9875dcf112bf7cac6447","observation_id":"6586546e-2a46-48dc-8144-e1e528e25074","resolution":{"observed_at":"2026-08-12T17:15:47.933909Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-12T13:06:27.906346Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.16502","last_updated":"2025-02-26T16:46:25Z","snapshot_observed_at":"2026-08-12T16:06:09.957131Z","submitted_at":"2024-11-25T15:37:27Z","title":"Interpreting Language Reward Models via Contrastive Explanations","version":2},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-12T13:06:27.906346Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2411.16502"},"observation_digest":"sha256:5191c47c784e327d744150bf2ace7c83b08a6d2e26d74d6da7fa11a2903f391a","observation_id":"179382e3-09a0-4fee-93bb-aa3b7174946c","resolution":{"observed_at":"2026-08-12T13:06:27.906346Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-11T23:15:56.399577Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.02685","last_updated":"2024-12-03T18:56:07Z","snapshot_observed_at":"2026-08-12T01:39:13.359412Z","submitted_at":"2024-12-03T18:56:07Z","title":"T-REG: Preference Optimization with Token-Level Reward Regularization","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-11T23:15:56.399577Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2412.02685"},"observation_digest":"sha256:76b0a489571f3cf3910589e32e3e9e1c699d2690caece335bbac08cfc43b06a6","observation_id":"ae6ea1ec-a58e-4b27-bf42-0d54fce4fbca","resolution":{"observed_at":"2026-08-11T23:15:56.399577Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-11T05:38:36.946910Z","title":"Wang, H., Xiong, W., Xie, T., Zhao, H., and Zhang, T","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.17365","last_updated":"2024-12-23T08:01:24Z","snapshot_observed_at":"2026-08-14T18:15:33.924925Z","submitted_at":"2024-12-23T08:01:24Z","title":"Boosting LLM via Learning from Data Iteratively and Selectively","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T05:38:36.946910Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2412.17365"},"observation_digest":"sha256:fcdc63b6fa19b6774c504cb151def8e236dfe3cddf99c9851ea321364dbb2277","observation_id":"4bbd1de7-89bb-4245-845b-a3ce80ee060e","resolution":{"observed_at":"2026-08-11T05:38:36.946910Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-10T22:26:46.653799Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.01668","last_updated":"2025-06-14T09:58:51Z","snapshot_observed_at":"2026-08-13T20:27:58.783300Z","submitted_at":"2025-01-03T06:50:06Z","title":"CoT-based Synthesizer: Enhancing LLM Performance through Answer Synthesis","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-10T22:26:46.653799Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2501.01668"},"observation_digest":"sha256:980bb6800c8273f584f9db806890c92c1b054756de8bcdc98d80203d2c0d25dc","observation_id":"76dd9384-d843-481b-9c36-4ba97b5fbc21","resolution":{"observed_at":"2026-08-10T22:26:46.653799Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-10T21:51:09.084843Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts, 2024 b","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.03884","last_updated":"2025-05-30T04:56:02Z","snapshot_observed_at":"2026-08-14T02:18:54.465541Z","submitted_at":"2025-01-07T15:46:42Z","title":"AlphaPO: Reward Shape Matters for LLM Alignment","version":4},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-10T21:51:09.084843Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2501.03884"},"observation_digest":"sha256:fa398ee692b6b9b8304a80a8d1842fd9d2fd7808886ed44a19d8d0b632a1195c","observation_id":"ff6c558d-bff6-401c-8d5d-ec1eeff439c6","resolution":{"observed_at":"2026-08-10T21:51:09.084843Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-10T17:18:41.059884Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.12368","last_updated":"2025-05-20T11:36:34Z","snapshot_observed_at":"2026-08-14T20:27:26.429867Z","submitted_at":"2025-01-21T18:47:32Z","title":"InternLM-XComposer2.5-Reward: A Simple Yet Effective Multi-Modal Reward Model","version":2},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-08-10T17:18:41.059884Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2501.12368"},"observation_digest":"sha256:3b29e46355cc065dcc082c328d254a26da8d22c29fd3366a891074a04025c0fd","observation_id":"e7092d19-35bc-4643-861b-ceaa4b504d6c","resolution":{"observed_at":"2026-08-10T17:18:41.059884Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-10T14:25:13.202305Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.15453","last_updated":"2025-01-28T06:35:32Z","snapshot_observed_at":"2026-08-14T01:18:38.620764Z","submitted_at":"2025-01-26T08:49:46Z","title":"Data-adaptive Safety Rules for Training Reward Models","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T14:25:13.202305Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2501.15453"},"observation_digest":"sha256:e298280649265e71624e01f9e494d55b71e9088ef8b9cae35efe4e8889246f65","observation_id":"589cc8e1-0d44-499e-99c3-9300818f6f27","resolution":{"observed_at":"2026-08-10T14:25:13.202305Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-10T18:56:37.826261Z","title":"In- terpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.16352","last_updated":"2025-01-18T20:17:31Z","snapshot_observed_at":"2026-08-13T09:16:20.668359Z","submitted_at":"2025-01-18T20:17:31Z","title":"Mixture of Experts (MoE): A Big Data Perspective","version":1},"reference_index":164,"source":"pdf_text","source_observed_at":"2026-08-10T18:56:37.826261Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2501.16352"},"observation_digest":"sha256:9d972c482defbb477dc503c2355b794d3d27cc36d561df8c455bfc5b06b0c76f","observation_id":"59b02dbf-8291-4708-9a20-58cfbc4b8699","resolution":{"observed_at":"2026-08-10T18:56:37.826261Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-10T13:47:45.821883Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.17195","last_updated":"2025-01-27T15:09:08Z","snapshot_observed_at":"2026-08-14T14:26:31.081387Z","submitted_at":"2025-01-27T15:09:08Z","title":"Atla Selene Mini: A General Purpose Evaluation Model","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T13:47:45.821883Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2501.17195"},"observation_digest":"sha256:ffe41e49602c38ddffbb071c4a584a6479e45c5c365e397a685c9d847f7027f2","observation_id":"b1bdb19a-0c51-4eb1-9ac9-686627d8924a","resolution":{"observed_at":"2026-08-10T13:47:45.821883Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-09T23:02:23.407554Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.18578","last_updated":"2025-02-26T18:58:53Z","snapshot_observed_at":"2026-08-13T04:54:13.331187Z","submitted_at":"2025-01-30T18:50:25Z","title":"R.I.P.: Better Models by Survival of the Fittest Prompts","version":2},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-09T23:02:23.407554Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2501.18578"},"observation_digest":"sha256:e7ce3057a3f602f970fa07051042138be0be5d0d4f0e03a934e63295f4ef8eeb","observation_id":"6b6b7f60-0c5e-4956-a978-5c3389a3891c","resolution":{"observed_at":"2026-08-09T23:02:23.407554Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-07T15:09:59.967922Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.17147","last_updated":"2025-05-22T08:22:57Z","snapshot_observed_at":"2026-08-14T04:09:37.700554Z","submitted_at":"2025-05-22T08:22:57Z","title":"MTSA: Multi-turn Safety Alignment for LLMs through Multi-round Red-teaming","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T15:09:59.967922Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2505.17147"},"observation_digest":"sha256:315d886ed05f9c5a852d1cb9c078a72f36117c212d580fe9e1a22d0cd8f56e42","observation_id":"338590b3-d24c-4faf-bf56-9d251942cf9c","resolution":{"observed_at":"2026-08-07T15:09:59.967922Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-07T14:31:05.677821Z","title":"Haoxiang Wang, Wei Xiong, Tengyang Xie, Han Zhao, and Tong Zhang","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.20336","last_updated":"2025-05-24T12:22:21Z","snapshot_observed_at":"2026-08-13T19:02:07.935526Z","submitted_at":"2025-05-24T12:22:21Z","title":"MOSLIM:Align with diverse preferences in prompts through reward classification","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-07T14:31:05.677821Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2505.20336"},"observation_digest":"sha256:a59f5750172f1f7311795043bfe8983e19d740315349fd60ffea19f5b0882bca","observation_id":"7323d348-b0fc-4847-9203-2a850d839a04","resolution":{"observed_at":"2026-08-07T14:31:05.677821Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-07T12:59:47.775977Z","title":"W ang, W","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.23247","last_updated":"2025-06-17T06:41:40Z","snapshot_observed_at":"2026-08-10T02:52:48.102229Z","submitted_at":"2025-05-29T08:54:06Z","title":"Accelerating RLHF Training with Reward Variance Increase","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T12:59:47.775977Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2505.23247"},"observation_digest":"sha256:20205fdd1ad4f1868115d2c9d798d063d1b57aad55288a0b4048a24d5524d8fc","observation_id":"4be40304-6593-490b-9633-e738ab1320a2","resolution":{"observed_at":"2026-08-07T12:59:47.775977Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":"2406.12845","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-07-04T13:49:52.398618Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":"a0d3f792-49d2-4bfb-ba7e-b4cdf3bd3fd6","year":2024},"citing_paper":{"arxiv_id":"2506.01937","last_updated":"2026-04-23T14:42:19Z","snapshot_observed_at":"2026-07-06T21:35:12.872021Z","submitted_at":"2025-06-02T17:54:04Z","title":"RewardBench 2: Advancing Reward Model Evaluation","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-19T11:18:03.965711Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2506.01937"},"observation_digest":"sha256:8e2499f856f5dde5612444211a5d32af1d27b296755a575bbcbfb5b614a4cf7e","observation_id":"60e0b91e-3536-448c-ad4b-885071fab307","resolution":{"observed_at":"2026-05-19T11:22:16.873818Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-07T05:41:59.868635Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.07434","last_updated":"2025-06-09T05:21:22Z","snapshot_observed_at":"2026-08-11T21:49:15.530339Z","submitted_at":"2025-06-09T05:21:22Z","title":"Well Begun is Half Done: Low-resource Preference Alignment by Weak-to-Strong Decoding","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T05:41:59.868635Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2506.07434"},"observation_digest":"sha256:1aa2efd2b377c3e589bee2dc75ed614f5bc123fc6efc9e9994e030481b87c827","observation_id":"875c2e27-052c-4e7f-bbbd-d4b8646455d2","resolution":{"observed_at":"2026-08-07T05:41:59.868635Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-06T23:35:36.079485Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.17533","last_updated":"2025-06-21T01:11:01Z","snapshot_observed_at":"2026-08-14T05:37:59.022330Z","submitted_at":"2025-06-21T01:11:01Z","title":"DuaShepherd: Integrating Stepwise Correctness and Potential Rewards for Mathematical Reasoning","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-06T23:35:36.079485Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2506.17533"},"observation_digest":"sha256:87c39a098c4390d211ffa1177bacb03cd665c4b00774a01d99428be69d147bda","observation_id":"e82a1853-8285-4461-8166-fcdc7998f07b","resolution":{"observed_at":"2026-08-06T23:35:36.079485Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-06T22:05:28.911043Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.22716","last_updated":"2025-06-28T01:52:50Z","snapshot_observed_at":"2026-08-14T16:20:26.515800Z","submitted_at":"2025-06-28T01:52:50Z","title":"BEST-Route: Adaptive LLM Routing with Test-Time Optimal Compute","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-06T22:05:28.911043Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2506.22716"},"observation_digest":"sha256:055af4283d0730eec25da67839c63b3a129addf41d29e15f1d4b3366cfd3e08b","observation_id":"cceac6b1-0107-4437-8266-c51db615d615","resolution":{"observed_at":"2026-08-06T22:05:28.911043Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-06T20:15:52.760921Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.03483","last_updated":"2025-07-08T05:05:04Z","snapshot_observed_at":"2026-08-13T03:42:11.397329Z","submitted_at":"2025-07-04T11:20:09Z","title":"BMMR: A Large-Scale Bilingual Multimodal Multi-Discipline Reasoning Dataset","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T20:15:52.760921Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2507.03483"},"observation_digest":"sha256:5ff66da0b3036c3a69dbccb14c3ea8a121a8acd0d7e62960e4b48e946818d939","observation_id":"849056ab-8732-4605-8bed-aa8161486bed","resolution":{"observed_at":"2026-08-06T20:15:52.760921Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-04T13:52:12.590452Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts.arXiv preprint arXiv:2406.12845,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.24696","last_updated":"2026-05-30T14:42:04Z","snapshot_observed_at":"2026-08-13T16:21:09.152242Z","submitted_at":"2025-09-29T12:28:23Z","title":"T-POP: Test-Time Personalization with Online Preference Feedback","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-04T13:52:12.590452Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2509.24696"},"observation_digest":"sha256:7e6ecb72e358e900a1e2865c7c964dfc21c7350beebc25adbbd4373616582327","observation_id":"b06ce073-c8ad-4ce4-bb51-450be7245415","resolution":{"observed_at":"2026-08-04T13:52:12.590452Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-04T13:15:47.225695Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.01167","last_updated":"2026-05-30T06:40:26Z","snapshot_observed_at":"2026-08-14T18:05:24.054010Z","submitted_at":"2025-10-01T17:54:15Z","title":"Simultaneous Multi-objective Alignment Across Verifiable and Non-verifiable Rewards","version":2},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-04T13:15:47.225695Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2510.01167"},"observation_digest":"sha256:a95acd6fd4797419e20e633cc1b5e800d9a7ea04368db598d0b9c4f9cd59224b","observation_id":"5daa66ef-e109-4667-939c-3f89ab5c0f62","resolution":{"observed_at":"2026-08-04T13:15:47.225695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":"2406.12845","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-07-04T13:49:52.398618Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":"a0d3f792-49d2-4bfb-ba7e-b4cdf3bd3fd6","year":2024},"citing_paper":{"arxiv_id":"2601.21619","last_updated":"2026-05-09T01:57:30Z","snapshot_observed_at":"2026-08-11T01:08:48.247252Z","submitted_at":"2026-01-29T12:22:45Z","title":"On the Overscaling Curse of Parallel Thinking: System Efficacy Contradicts Sample Efficiency","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-16T10:38:09.875786Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2601.21619"},"observation_digest":"sha256:6353d98d9b9238bff64b48f9e1cd585ccb75adba7e1384ebba157f7a85999529","observation_id":"73b6f605-1e24-4871-bb6d-233259723df4","resolution":{"observed_at":"2026-05-16T10:40:51.396297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":"2406.12845","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-07-04T13:49:52.398618Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":"a0d3f792-49d2-4bfb-ba7e-b4cdf3bd3fd6","year":2024},"citing_paper":{"arxiv_id":"2604.25872","last_updated":"2026-04-28T17:10:15Z","snapshot_observed_at":"2026-08-14T07:50:02.401349Z","submitted_at":"2026-04-28T17:10:15Z","title":"When Errors Can Be Beneficial: A Categorization of Imperfect Rewards for Policy Gradient","version":1},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-05-07T16:24:43.688967Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2604.25872"},"observation_digest":"sha256:809f5e6797ffede9aaa73620e72bd742ec2cfd2e739f1409251ae7f2043d1ed3","observation_id":"78c93807-1306-4273-b3fb-328262e3caf1","resolution":{"observed_at":"2026-05-11T23:41:19.698135Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":"2406.12845","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-07-04T13:49:52.398618Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":"a0d3f792-49d2-4bfb-ba7e-b4cdf3bd3fd6","year":2024},"citing_paper":{"arxiv_id":"2605.02906","last_updated":"2026-07-31T03:14:41Z","snapshot_observed_at":"2026-08-11T14:45:25.895700Z","submitted_at":"2026-04-06T02:40:18Z","title":"OpsLLM: Construction of Large Language Model for Software Operations with Multi-stage Learning","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-10T20:07:07.548384Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2605.02906"},"observation_digest":"sha256:38a6da0268ff5bd9af2e9f8d022934539076ce01c1644256f2f9ed5a275d11f5","observation_id":"aa5406f5-3998-43dc-a48e-2dcc17ae23a3","resolution":{"observed_at":"2026-05-10T22:15:48.664158Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":"2406.12845","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-07-04T13:49:52.398618Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":"a0d3f792-49d2-4bfb-ba7e-b4cdf3bd3fd6","year":2024},"citing_paper":{"arxiv_id":"2605.02906","last_updated":"2026-07-31T03:14:41Z","snapshot_observed_at":"2026-08-11T14:45:25.895700Z","submitted_at":"2026-04-06T02:40:18Z","title":"OpsLLM: Construction of Large Language Model for Software Operations with Multi-stage Learning","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-13T06:25:16.650306Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2605.02906"},"observation_digest":"sha256:a4d09828bf20bdc73a81261cdb6aad699a6b63a842e4e54a635ea4dabaa1b955","observation_id":"ef04ec64-c2e6-464d-91e8-d7ff1db14216","resolution":{"observed_at":"2026-05-13T06:27:24.620457Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-03T02:30:35.617309Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2605.02906","last_updated":"2026-07-31T03:14:41Z","snapshot_observed_at":"2026-08-11T14:45:25.895700Z","submitted_at":"2026-04-06T02:40:18Z","title":"OpsLLM: Construction of Large Language Model for Software Operations with Multi-stage Learning","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-03T02:30:35.617309Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2605.02906"},"observation_digest":"sha256:683bebdcb711fcb7dc260c5e1111bd703929e3c47f9551e73393dd12e23193ba","observation_id":"684ed106-daed-4f8e-9b11-24a1aced04e7","resolution":{"observed_at":"2026-08-03T02:30:35.617309Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":"2406.12845","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-07-04T13:49:52.398618Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":"a0d3f792-49d2-4bfb-ba7e-b4cdf3bd3fd6","year":2024},"citing_paper":{"arxiv_id":"2605.11679","last_updated":"2026-05-13T09:28:34Z","snapshot_observed_at":"2026-08-11T08:28:39.520767Z","submitted_at":"2026-05-12T07:38:59Z","title":"Explaining and Breaking the Safety-Helpfulness Ceiling via Preference Dimensional Expansion","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-13T01:03:10.263663Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2605.11679"},"observation_digest":"sha256:1e2b6a92be91075317543b8ca26ca0288fac13fbcb266011d5f759d55daa0f7c","observation_id":"881263e6-a814-4432-ba11-50152e4a4d62","resolution":{"observed_at":"2026-05-13T01:07:00.595220Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":"2406.12845","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-07-04T13:49:52.398618Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":"a0d3f792-49d2-4bfb-ba7e-b4cdf3bd3fd6","year":2024},"citing_paper":{"arxiv_id":"2605.11679","last_updated":"2026-05-13T09:28:34Z","snapshot_observed_at":"2026-08-11T08:28:39.520767Z","submitted_at":"2026-05-12T07:38:59Z","title":"Explaining and Breaking the Safety-Helpfulness Ceiling via Preference Dimensional Expansion","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-14T21:12:06.989077Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2605.11679"},"observation_digest":"sha256:a2170f3db7a886fb4110250810df422deaf8b3412f27ce0d34b74f4ed81c2e03","observation_id":"afc53f3c-b12b-409c-94b1-edb69dece7a5","resolution":{"observed_at":"2026-05-14T21:12:58.918870Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":"2406.12845","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-07-04T13:49:52.398618Z","title":"Interpretable preferences via multi-objective reward modeling and mixture-of-experts","venue":null,"work_id":"a0d3f792-49d2-4bfb-ba7e-b4cdf3bd3fd6","year":2024},"citing_paper":{"arxiv_id":"2606.26917","last_updated":"2026-06-25T11:53:37Z","snapshot_observed_at":"2026-08-14T19:58:24.816006Z","submitted_at":"2026-06-25T11:53:37Z","title":"GEOALIGN: Geometric Rollout Curation for Robust LLM Reinforcement Learning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-26T04:49:28.598430Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2606.26917"},"observation_digest":"sha256:386863436adc0f7bf56c46a588e30704b129bf9efd62f97f7ef98636d3906e50","observation_id":"aae0c89a-f6f3-440e-8496-99ba88e4727b","resolution":{"observed_at":"2026-07-04T13:49:52.400308Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12845","snapshot_observed_at":"2026-08-03T00:55:25.496824Z","title":"arXiv preprint arXiv:2406.12845 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28636","last_updated":"2026-05-19T13:56:13Z","snapshot_observed_at":"2026-08-10T20:43:57.798906Z","submitted_at":"2026-05-19T13:56:13Z","title":"Chain-of-Models: Cross-Model Auditing for Bias-Robust LLM Judges","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-03T00:55:25.496824Z"},"links":{"cited_paper":"/paper/2406.12845","citing_paper":"/paper/2607.28636"},"observation_digest":"sha256:61b44638a85d8a09e7af3f28cd343d48c6886dae9e7c7d06185f6109a41c3e85","observation_id":"863bc362-198e-4315-9319-54c2924a9a3c","resolution":{"observed_at":"2026-08-03T00:55:25.496824Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2406.12845/citation-record","integrity":"/paper/2406.12845/integrity","json":"/paper/2406.12845/citation-record.json","paper":"/paper/2406.12845"},"outbound":[],"paper":{"arxiv_id":"2406.12845","last_updated":"2024-06-18T17:58:28Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-12T23:39:50.330661Z","submitted_at":"2024-06-18T17:58:28Z","title":"Interpretable Preferences via Multi-Objective Reward Modeling and Mixture-of-Experts"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 31 inbound Pith citation observations for arXiv:2406.12845."}