{"as_of":"2026-08-19T12:22:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d466b4f5459b49053891e34c5763577b88a840d147a1f4852363001e0be486e7","coverage":[{"denominator":27,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":27,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T23:02:32.171548Z","state":"measured"},{"denominator":43,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":43,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":16,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":16,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T17:47:22.122408Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T20:40:08.281402Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2510.07286","last_updated":"2026-04-11T09:47:15Z","snapshot_observed_at":"2026-08-15T02:51:05.476470Z","submitted_at":"2025-10-08T17:46:02Z","title":"Evolutionary Profiles for Protein Fitness Prediction","version":3},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-18T08:56:44.544404Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2510.07286"},"observation_digest":"sha256:f772093fbd81290394b480ba520f9e7bf497896ed9319d217f359b048d0f9c8f","observation_id":"a6e26b5f-d64a-4902-a3dc-db33497432dc","resolution":{"observed_at":"2026-05-18T09:01:09.671919Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-08-03T11:44:22.344517Z","title":"agent”, “I agree","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2601.19921","last_updated":"2026-06-03T15:14:01Z","snapshot_observed_at":"2026-08-17T06:45:22.012582Z","submitted_at":"2026-01-09T02:38:30Z","title":"Demystifying Multi-Agent Debate: The Role of Confidence and Diversity","version":3},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-03T11:44:22.344517Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2601.19921"},"observation_digest":"sha256:371bd276d78b11db17eb56d2ef1a42af5bdda584ca1b6643b85e533a94ed78e5","observation_id":"62e3d405-6efa-40e4-ab05-18fd40f6e5a3","resolution":{"observed_at":"2026-08-03T11:44:22.344517Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2601.21257","last_updated":"2026-04-19T21:04:27Z","snapshot_observed_at":"2026-08-11T02:55:47.386292Z","submitted_at":"2026-01-29T04:36:52Z","title":"MoCo: A One-Stop Shop for Model Collaboration Research","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-16T10:17:37.129753Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2601.21257"},"observation_digest":"sha256:c379b239bab385ed16ab59a2c25a2079b815a9db28c988a5d58a45893f633639","observation_id":"a1f4205f-c244-443a-9783-f5a56707a7bd","resolution":{"observed_at":"2026-05-16T10:17:43.763828Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-13T23:28:12.790404Z","title":"The ma- jority is not always right: Rl training for solution aggregation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.16867","last_updated":"2026-06-03T09:37:20Z","snapshot_observed_at":"2026-08-14T22:25:11.633208Z","submitted_at":"2026-03-17T17:59:51Z","title":"Efficient Reasoning on the Edge","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-07-13T23:28:12.790404Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2603.16867"},"observation_digest":"sha256:51dc45c3f6e25ccdefecfeded71af9de550556d420f8378fef4930180650f146","observation_id":"ddb9fb00-fec4-4199-aca8-14bcbad2a379","resolution":{"observed_at":"2026-07-13T23:28:12.790404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2604.05868","last_updated":"2026-04-07T13:28:09Z","snapshot_observed_at":"2026-08-13T03:00:47.388182Z","submitted_at":"2026-04-07T13:28:09Z","title":"Understanding Performance Gap Between Parallel and Sequential Sampling in Large Reasoning Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-10T18:36:01.200412Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2604.05868"},"observation_digest":"sha256:62da19a4469cc949ddff000091c618fb3422b3fe089304104f4d511800c3cf45","observation_id":"fb0f0e4e-ee17-45d2-b207-1a6183c422ba","resolution":{"observed_at":"2026-05-11T00:20:52.463558Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2605.08083","last_updated":"2026-05-12T03:46:05Z","snapshot_observed_at":"2026-08-14T03:41:26.776348Z","submitted_at":"2026-05-08T17:59:40Z","title":"LLMs Improving LLMs: Agentic Discovery for Test-Time Scaling","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-11T01:54:34.000216Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2605.08083"},"observation_digest":"sha256:8f584e861052d5dde91b7002799c644f6ed4676ca388faba61dbf71cad093cf9","observation_id":"8d1d1980-7186-4efb-96b2-00e724a171aa","resolution":{"observed_at":"2026-05-11T04:10:58.109816Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2605.08083","last_updated":"2026-05-12T03:46:05Z","snapshot_observed_at":"2026-08-14T03:41:26.776348Z","submitted_at":"2026-05-08T17:59:40Z","title":"LLMs Improving LLMs: Agentic Discovery for Test-Time Scaling","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-13T07:09:02.672233Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2605.08083"},"observation_digest":"sha256:09a0dbca5e7bed268704078efe3f31047bc17ffaff43691b8a2902d23f69a112","observation_id":"40c90944-5dd9-48cf-912b-3373a93a9598","resolution":{"observed_at":"2026-05-13T07:12:28.331875Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2605.15513","last_updated":"2026-05-15T01:16:12Z","snapshot_observed_at":"2026-08-01T17:14:08.866403Z","submitted_at":"2026-05-15T01:16:12Z","title":"CAPS: Cascaded Adaptive Pairwise Selection for Efficient Parallel Reasoning","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-19T15:39:56.255871Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2605.15513"},"observation_digest":"sha256:1debd4c735386f82470e6d11b7d92a8b47148994726f020a73d5c44276b1429a","observation_id":"7732f651-d6c9-4533-94a3-a5e0a393515b","resolution":{"observed_at":"2026-05-19T15:42:38.386015Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2605.24486","last_updated":"2026-05-23T09:21:51Z","snapshot_observed_at":"2026-08-14T02:55:27.894688Z","submitted_at":"2026-05-23T09:21:51Z","title":"AgentFugue: Agent Scaling for Long-Horizon Tasks through Collective Reasoning","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-06-30T13:29:36.710152Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2605.24486"},"observation_digest":"sha256:da942d036eba54552b6a738dcdacdccfb30404cb6d582172ecf83d1f0de42566","observation_id":"48fca032-9f85-4661-bca5-d45c888c86e6","resolution":{"observed_at":"2026-06-30T13:34:40.439558Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2606.00660","last_updated":"2026-05-30T10:21:20Z","snapshot_observed_at":"2026-07-06T23:41:19.955034Z","submitted_at":"2026-05-30T10:21:20Z","title":"FineVerify: Scaling Test-Time Compute with Fine-Grained Self-Verification for Agentic Search","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-28T19:05:28.264896Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2606.00660"},"observation_digest":"sha256:8e5b2983425e020870ee662b64b1e194666aaf7dc7d93053e2a3d44cf0ea5d94","observation_id":"895aa897-db85-468f-8112-f3c2c45045e1","resolution":{"observed_at":"2026-06-28T19:12:35.005229Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2606.01533","last_updated":"2026-06-01T01:29:36Z","snapshot_observed_at":"2026-08-13T23:51:51.905766Z","submitted_at":"2026-06-01T01:29:36Z","title":"Multi-Agent Computer Use","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-06-28T12:28:23.880150Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2606.01533"},"observation_digest":"sha256:7a4b8fe17717e61c900418e753993609b1fd1e79f2aa8ada608b5f4311781f1d","observation_id":"629e359f-7e5d-4ea1-9475-5c25d3f90d17","resolution":{"observed_at":"2026-07-02T01:06:25.097303Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2606.07812","last_updated":"2026-06-05T19:39:35Z","snapshot_observed_at":"2026-08-16T00:44:28.545690Z","submitted_at":"2026-06-05T19:39:35Z","title":"Scaling Participation in Modular AI Systems","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-27T21:49:27.042616Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2606.07812"},"observation_digest":"sha256:dbe4971fc45ce7f201dd47077a403e7fd51a660b652e42c0d8e0fddd19b333ec","observation_id":"4db2c2e3-fe49-4d3c-9b4c-c3679d3e74c8","resolution":{"observed_at":"2026-07-02T18:57:16.583268Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2606.25996","last_updated":"2026-07-04T15:07:51Z","snapshot_observed_at":"2026-08-07T11:35:11.679286Z","submitted_at":"2026-06-24T16:08:31Z","title":"Autodata: An agentic data scientist to create high quality synthetic data","version":1},"reference_index":124,"source":"arxiv_source","source_observed_at":"2026-06-25T19:50:35.574454Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2606.25996"},"observation_digest":"sha256:1ba12b5d1b2d47c9afa04e72540547e1507b080ac906f22940030e37530002ca","observation_id":"e8c9da9b-852a-4bd0-a110-ea27ddf38415","resolution":{"observed_at":"2026-07-04T20:40:08.283233Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":"2509.06870","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-04T20:40:08.281402Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":"fff415c1-d63b-4ecc-868b-e956a6b87ed0","year":2025},"citing_paper":{"arxiv_id":"2606.25996","last_updated":"2026-07-04T15:07:51Z","snapshot_observed_at":"2026-08-07T11:35:11.679286Z","submitted_at":"2026-06-24T16:08:31Z","title":"Autodata: An agentic data scientist to create high quality synthetic data","version":2},"reference_index":124,"source":"arxiv_source","source_observed_at":"2026-06-26T05:16:12.361470Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2606.25996"},"observation_digest":"sha256:3d1756fa7788ca0c1a970b4ced9ffe5edf94c06732cfa97f5fd55d5068052e81","observation_id":"3b2edf05-29c8-49d0-991c-ff18678d75e3","resolution":{"observed_at":"2026-07-04T13:19:51.228081Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-07-12T12:08:06.206832Z","title":"The majority is not always right: Rl training for solution aggregation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.25996","last_updated":"2026-07-04T15:07:51Z","snapshot_observed_at":"2026-08-07T11:35:11.679286Z","submitted_at":"2026-06-24T16:08:31Z","title":"Autodata: An agentic data scientist to create high quality synthetic data","version":3},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-07-12T12:08:06.206832Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2606.25996"},"observation_digest":"sha256:b03ecb19f255f393517be3e9b3940e6b53d60095584044c9d323242855575d25","observation_id":"cd10e78d-3eb5-4460-8ce0-13b9e574f31e","resolution":{"observed_at":"2026-07-12T12:08:06.206832Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.06870","snapshot_observed_at":"2026-08-05T17:47:22.122408Z","title":"arXiv preprint arXiv:2509.06870 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.03506","last_updated":"2026-08-04T11:45:46Z","snapshot_observed_at":"2026-08-15T14:06:43.024132Z","submitted_at":"2026-08-04T11:45:46Z","title":"When Many Answers Are Valid, Voting Fails: Symbolic Verification for Best-of-K Causal Reasoning in LLMs","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-05T17:47:22.122408Z"},"links":{"cited_paper":"/paper/2509.06870","citing_paper":"/paper/2608.03506"},"observation_digest":"sha256:dc24c270b6f40e164846ddf3a88e13e1607078e5f98efb1940c74a8d81ea0c37","observation_id":"c4e31add-bd7c-47c0-8c6f-9d01ed1c5ff7","resolution":{"observed_at":"2026-08-05T17:47:22.122408Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2509.06870/citation-record","integrity":"/paper/2509.06870/integrity","json":"/paper/2509.06870/citation-record.json","paper":"/paper/2509.06870"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.021495Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.021495Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:cd5b8e97a37b7baf349c1cd7cef4ccb21559c48547db5f88b62863fa87edfd69","observation_id":"65a5569d-d701-4b56-aa2f-c03be79d1f1b","resolution":{"observed_at":"2026-08-04T23:02:32.021495Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.031193Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.031193Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:7b216e23fd9030933abdba013a39989f3dc17dcabed342da82c95eccf2205dc0","observation_id":"c7955311-c588-4c39-a06f-703d1d87030a","resolution":{"observed_at":"2026-08-04T23:02:32.031193Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.036556Z","title":"Let ' s sample step by step: Adaptive-consistency for efficient reasoning and coding with LLM s","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.036556Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:8301899062e00e32a6241f677335c7cf042040dd40717d6d891352ebc40af5e7","observation_id":"115a7992-46b3-44cd-9365-575c74f66574","resolution":{"observed_at":"2026-08-04T23:02:32.036556Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.23281","last_updated":"2026-01-14T21:39:58Z","snapshot_observed_at":"2026-08-02T10:05:44.330694Z","submitted_at":"2025-05-29T09:28:06Z","title":"MathArena: Evaluating LLMs on Uncontaminated Math Competitions","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.23281","snapshot_observed_at":"2026-08-04T23:02:32.041489Z","title":"Matharena: Evaluating llms on uncontaminated math competitions","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.041489Z"},"links":{"cited_paper":"/paper/2505.23281","citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:82c809f5a17b978efb9cb76b6a4f36af856542e0955678dc50166bf762c2c137","observation_id":"033d5373-fa45-4400-b7a3-8c5f52de97b8","resolution":{"observed_at":"2026-08-04T23:02:32.041489Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21787","last_updated":"2024-12-30T19:03:24Z","snapshot_observed_at":"2026-08-14T10:00:52.343929Z","submitted_at":"2024-07-31T17:57:25Z","title":"Large Language Monkeys: Scaling Inference Compute with Repeated Sampling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21787","snapshot_observed_at":"2026-08-04T23:02:32.047668Z","title":"Large language monkeys: Scaling inference compute with repeated sampling","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.047668Z"},"links":{"cited_paper":"/paper/2407.21787","citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:585e25a78e92b32c7380c58044e3c709e198943da1e9e575173e3f0af8f4c11a","observation_id":"246d77d3-215c-4a63-8da4-d69ba22489f9","resolution":{"observed_at":"2026-08-04T23:02:32.047668Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.567714Z","title":"Universal self-consistency for large language models","venue":null,"work_id":"a74da777-370d-47fa-9913-1a3cd559b111","year":2024},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.053765Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:b089b2b78e9b3462efd925708722f4b1bb87039fc5a392b0867c3143b4d50319","observation_id":"2d04ee91-586a-4181-a3b0-90c75595cfe2","resolution":{"observed_at":"2026-08-04T23:02:32.574128Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.15260","last_updated":"2025-08-21T05:48:38Z","snapshot_observed_at":"2026-08-12T17:10:09.978771Z","submitted_at":"2025-08-21T05:48:38Z","title":"Deep Think with Confidence","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.15260","snapshot_observed_at":"2026-08-04T23:02:32.059071Z","title":"Deep think with confidence","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.059071Z"},"links":{"cited_paper":"/paper/2508.15260","citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:6cd51799231c3a46b6c75f28bccf83945343328595120672d0b830ac7b211aec","observation_id":"8fea90f3-c4ed-4c8a-9a24-4353b7a04898","resolution":{"observed_at":"2026-08-04T23:02:32.059071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-08-15T12:33:55.451951Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-04T23:02:32.067617Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.067617Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:12cb55570f039cb122a37dcb9ca4a288d02104947d372f27e92862e8b88439cc","observation_id":"75db96b4-81f9-439c-b4a3-2edc13b8a22b","resolution":{"observed_at":"2026-08-04T23:02:32.067617Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.073554Z","title":"Mirror-consistency: Harnessing inconsistency in majority voting","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.073554Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:e6a6d30e6a2aed004c5f6529a04736c099e7dc999c437c7abdac723e753229dd","observation_id":"120fbfd4-5356-4034-8bcc-5b941e5ce156","resolution":{"observed_at":"2026-08-04T23:02:32.073554Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16720","last_updated":"2026-04-30T02:46:40Z","snapshot_observed_at":"2026-08-19T11:46:55.171293Z","submitted_at":"2024-12-21T18:04:31Z","title":"OpenAI o1 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16720","snapshot_observed_at":"2026-08-04T23:02:32.078812Z","title":"Openai o1 system card","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.078812Z"},"links":{"cited_paper":"/paper/2412.16720","citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:abfa4565bbfbbe6ca5078666cac70fdb1dd1bf28bfca430eab6368940e3cd8e8","observation_id":"861251c2-8fe3-41e0-ab9b-8113a5cabc50","resolution":{"observed_at":"2026-08-04T23:02:32.078812Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.550722Z","title":"Enhancing language model reasoning via weighted reasoning in self-consistency","venue":null,"work_id":"f277b5be-b80d-4fde-86b5-8569f823bb48","year":2024},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.085460Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:9a79d2c34835a7846e5ef12e86d6abe16fa8964f138335a8a19956c1be7ec59a","observation_id":"d3f2bac6-d7b6-4738-bb30-1d3183d69cb4","resolution":{"observed_at":"2026-08-04T23:02:32.555884Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15084","last_updated":"2025-01-17T07:12:55Z","snapshot_observed_at":"2026-08-15T17:39:14.326195Z","submitted_at":"2024-12-19T17:29:44Z","title":"AceMath: Advancing Frontier Math Reasoning with Post-Training and Reward Modeling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15084","snapshot_observed_at":"2026-08-04T23:02:32.090666Z","title":"Acemath: Advancing frontier math reasoning with post-training and reward modeling","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.090666Z"},"links":{"cited_paper":"/paper/2412.15084","citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:a18c806c55ec6dc253391b4fbedb99738a038233ebfc56bdccd307236c70b34d","observation_id":"760542d6-a34c-4a94-a573-b9607593e220","resolution":{"observed_at":"2026-08-04T23:02:32.090666Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.096151Z","title":"Tang, Manan Roongta, Colin Cai, Jeffrey Luo, Li Erran Li, Raluca Ada Popa, and Ion Stoica","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.096151Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:0827b398299d8d619a9362415b9c88d29b4397a12ce089197d751bc02d30b39f","observation_id":"0d349a34-4ea0-468d-814b-dd3db792921a","resolution":{"observed_at":"2026-08-04T23:02:32.096151Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.100574Z","title":"Learning to reason across parallel samples for llm reasoning","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.100574Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:c05b9a4ee8564111753e3a9c1e0d53799639c8fe5a1eec5571ce8385d1abef04","observation_id":"ac312e43-5cd1-4dd9-8c50-6a717c53fd0d","resolution":{"observed_at":"2026-08-04T23:02:32.100574Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-04T23:02:32.105419Z","title":"Deepseekmath: Pushing the limits of mathematical reasoning in open language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.105419Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:e79931f53c983dc6fb01bc2137ff2dadcf6e44efb23fe26eb3649180895e5675","observation_id":"240ad768-35aa-40f3-b903-309304cabc36","resolution":{"observed_at":"2026-08-04T23:02:32.105419Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.110199Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.110199Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:d0c6f3f46a498b071953ea46e4f3e408c1f27f9ecb1fdd04e31c1a011527f485","observation_id":"428b07c8-10b0-4b9c-afc4-3a10362ec88d","resolution":{"observed_at":"2026-08-04T23:02:32.110199Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2022.acl-long.591","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Uncertainty determines the adequacy of the mode and the tractability of decoding in sequence-to-sequence models","venue":"Proceedings of the 60th Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers)","work_id":"afff9d90-9acf-44d4-9e50-d4b2f0ee7903","year":2022},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.115148Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:2331a8ee6e53b5983b671f8ed138baeccbb1d6ae273e14d522b8e9d2ebccf07f","observation_id":"73305e87-efcb-4de9-ba9f-302dd6640b2d","resolution":{"observed_at":"2026-08-04T23:02:32.230658Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.119563Z","title":"Chi, Sharan Narang, Aakanksha Chowdhery, and Denny Zhou","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.119563Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:eaffe232b88da675c40ab244c4a9707b8cb592ade5a104116e92b1327a39d974","observation_id":"37665e65-4eab-4280-89a5-b6eec33415cd","resolution":{"observed_at":"2026-08-04T23:02:32.119563Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.124791Z","title":"Chain-of-thought prompting elicits reasoning in large language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.124791Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:97c76ff97b14c66dbbce46e66f394e70a37c66fd2e8b151834f6ea492fc0eb13","observation_id":"9b759c66-b9b2-4138-a68a-93ecb2f37c6c","resolution":{"observed_at":"2026-08-04T23:02:32.124791Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.502267Z","title":"From decoding to meta-generation: Inference-time algorithms for large language models","venue":null,"work_id":"4a54c2dc-ead2-40b0-be24-8e6f3211c6e6","year":2024},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.130996Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:65d7a37a3122c5220c0adad078a67c272aa83e2e7df40a5e9a35cb6eef5b0ff7","observation_id":"434977b0-f388-491b-81fb-d49356359a1c","resolution":{"observed_at":"2026-08-04T23:02:32.507340Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.485630Z","title":"Inference scaling laws: An empirical analysis of compute-optimal inference for LLM problem-solving","venue":null,"work_id":"00f0817b-ebf2-49a9-a490-bc57c42b6b78","year":2025},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.135572Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:6c1a6c8999296529b2576903fec8ad418dfe09fbc22e6e002108d8413425bd4f","observation_id":"0f88c749-02b8-4fb2-a427-74b0d1fe1d00","resolution":{"observed_at":"2026-08-04T23:02:32.490870Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2023.findings-emnlp.203","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Dynamic voting for efficient reasoning in large language models","venue":null,"work_id":"920a5dc9-23e0-4124-babe-5d8782657a9d","year":2023},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.139814Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:05adb83c7ec8d5044ca99b09cf3716cb61c2fa7bdd697c014b5576682a02c88e","observation_id":"6161fd55-7761-4ca0-85c3-e764885a745b","resolution":{"observed_at":"2026-08-04T23:02:32.214285Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12122","last_updated":"2024-09-18T16:45:37Z","snapshot_observed_at":"2026-08-17T18:51:13.219936Z","submitted_at":"2024-09-18T16:45:37Z","title":"Qwen2.5-Math Technical Report: Toward Mathematical Expert Model via Self-Improvement","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12122","snapshot_observed_at":"2026-08-04T23:02:32.144273Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.144273Z"},"links":{"cited_paper":"/paper/2409.12122","citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:1693e90623ef59643433c5140f1b5027a7fd2fe3da34e5bad2426ca4ea9187c4","observation_id":"8cc1524e-e085-4503-ad3b-1c760259cebf","resolution":{"observed_at":"2026-08-04T23:02:32.144273Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-08-04T23:02:32.149619Z","title":"Qwen3 technical report","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.149619Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:542264511a71c4419d4cfc5cd4912d194a0f9edb9512f8af127e71365f6536fb","observation_id":"9c3a1fcc-83ac-4853-a31e-c9f6dc6d8a3e","resolution":{"observed_at":"2026-08-04T23:02:32.149619Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.154689Z","title":"@esa (Ref","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.154689Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:ed00f7a858c962feaee6985615fa2934ba42394d0590f6892ec97589329a8dd2","observation_id":"082a8ff6-e8f2-4d1c-90ac-4cac7451fe62","resolution":{"observed_at":"2026-08-04T23:02:32.154689Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.159579Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.159579Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:15766b2a157d6f726bf4e02b42c7e89cd122b549b213fea62c52e5e9a749004d","observation_id":"0d3dffb8-0d6b-4e46-b50b-23d60207ba7a","resolution":{"observed_at":"2026-08-04T23:02:32.159579Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-04T23:02:32.447275Z","title":"MIdaId: b5VȮBd)G̶ Rૉ,l","venue":null,"work_id":"3cd46bdd-c506-4358-9de5-624d4de0ad06","year":1976},"citing_paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-04T23:02:32.171548Z"},"links":{"citing_paper":"/paper/2509.06870"},"observation_digest":"sha256:bf5ddfc28cf3e256a9286ad14869efb52b4ccd6b18bd1aa75a55aac990b16fb7","observation_id":"82f22e92-311a-45d7-8f29-c7bfc1eb54bd","resolution":{"observed_at":"2026-08-04T23:02:32.452332Z","resolver_source":"raw_fallback","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2509.06870","last_updated":"2025-09-08T16:39:38Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-13T12:05:04.948661Z","submitted_at":"2025-09-08T16:39:38Z","title":"The Majority is not always right: RL training for solution aggregation"},"reference_resolution":{"displayed":27,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":20,"verified_exact":2,"verified_fuzzy":4},"total_outbound_references":27},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 27 of 27 outbound references and 16 inbound Pith citation observations for arXiv:2509.06870."}