{"as_of":"2026-08-09T23:55:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5e40ff84ec0278b0ade7ff0f90ac93dd514b97a6fee5d7040dbbf5e21740db52","coverage":[{"denominator":42,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":42,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T00:19:14.894582Z","state":"measured"},{"denominator":42,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":42,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.00816/citation-record","integrity":"/paper/2608.00816/integrity","json":"/paper/2608.00816/citation-record.json","paper":"/paper/2608.00816"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.06507","last_updated":"2025-07-14T07:46:11Z","snapshot_observed_at":"2026-08-08T04:13:41.618907Z","submitted_at":"2025-07-09T03:13:08Z","title":"GR-LLMs: Recent Advances in Generative Recommendation Based on Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.06507","doi":null,"metadata_source":"pith","pith_arxiv_id":"2507.06507","snapshot_observed_at":"2026-08-05T00:19:15.461289Z","title":"GR-LLMs: Recent Advances in Generative Recommendation Based on Large Language Models","venue":"cs.IR","work_id":"7d1d5c18-0e72-4711-b9b5-fef023919a7b","year":2025},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:10.824627Z"},"links":{"cited_paper":"/paper/2507.06507","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:7b6a98611d853e2af30fef4ff8cc165af40341cd665f693e88220114637a85c9","observation_id":"8b7ac46a-b531-47ff-b5fc-8ca7ce0599d3","resolution":{"observed_at":"2026-08-05T00:19:15.537312Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:17.228609Z","title":"Proceedings of the 30th ACM SIGKDD conference on Knowledge Discovery and Data Mining , pages=","venue":null,"work_id":"251107ba-4043-49fa-aeb8-43b802910cc1","year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:10.924810Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:ce941df4fa799f63d848b911f8f51594d03e2b7934761d723fe8e0f6f14fbf15","observation_id":"1b20a2cc-3392-49c1-8446-6fcc17730ddd","resolution":{"observed_at":"2026-08-05T00:19:17.287995Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:17.100837Z","title":"World Wide Web , volume=","venue":null,"work_id":"27ff870e-fd3f-4fe9-86fd-085cd70fb7c0","year":2024},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:11.058092Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:6363e059c672d1edb34d98eb374f1a1f037bebb2f2e5b73a5744b75bdd073e5c","observation_id":"8456ce5d-4208-4bc5-b7a5-eff25ce1bb68","resolution":{"observed_at":"2026-08-05T00:19:17.145515Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:11.200723Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:11.200723Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:7bb5c4e882de1eb6f9e73df5a67d7816b6a59affaea2cfdb12f5b0a6472f18e5","observation_id":"b70799e8-04ac-4b39-8e3b-fc3eda2a73d4","resolution":{"observed_at":"2026-08-05T00:19:11.200723Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:16.926761Z","title":"International Conference on Machine Learning (ICML) , year=","venue":null,"work_id":"390a6b66-230b-4c93-9e3b-61d00ec8a170","year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:11.304724Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:a6099c1b351497d9b2d98392e88436d7fd8659d99e33d44f0ba3b882cd12cf75","observation_id":"5cfc5d06-c236-4aaf-a350-5b0554f65af9","resolution":{"observed_at":"2026-08-05T00:19:17.018679Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:16.792003Z","title":"Proceedings of the 41st International Conference on Machine Learning , pages =","venue":null,"work_id":"5dc41235-74d1-4a56-a274-fbb3433ad936","year":2024},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:11.426050Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:dd3f7c710bc360e7366bd888de26adfcb0d60ebb0096f0dbd3ba0b7c68664938","observation_id":"084da148-7a8b-4d78-96e7-1832ff6a7c20","resolution":{"observed_at":"2026-08-05T00:19:16.832458Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:11.586804Z","title":"2018 IEEE international conference on data mining (ICDM) , pages=","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:11.586804Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:44cad639ca207b3005d0d0f9e08d863ebabe637a4e29bdeb2e10d6c5434f9e2c","observation_id":"b04b891e-61c5-4812-9284-43979c9be78a","resolution":{"observed_at":"2026-08-05T00:19:11.586804Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.18965","last_updated":"2025-02-26T09:25:10Z","snapshot_observed_at":"2026-07-06T20:42:55.911327Z","submitted_at":"2025-02-26T09:25:10Z","title":"OneRec: Unifying Retrieve and Rank with Generative Recommender and Iterative Preference Alignment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.18965","snapshot_observed_at":"2026-08-05T00:19:11.737576Z","title":"arXiv preprint arXiv:2502.18965 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:11.737576Z"},"links":{"cited_paper":"/paper/2502.18965","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:95822981313dc989008ad9f1609015959d2f1c81580cab5ead7aa293ed005641","observation_id":"9a9774e2-2e22-4421-862f-ae355f691df6","resolution":{"observed_at":"2026-08-05T00:19:11.737576Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:11.844487Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:11.844487Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:446e25d366f62dc44e4c2051aa02d1758650ef5a9055991d620339a62019cf58","observation_id":"5ec8ec22-014b-4036-8fd6-911d5be069c8","resolution":{"observed_at":"2026-08-05T00:19:11.844487Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:11.956367Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:11.956367Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:5399cce3043260ac6d5ab72de11bc3b4225e6a63ef2a0c10295464848456b72f","observation_id":"b0ca95dc-e35c-4856-98e3-c60acac5c8bf","resolution":{"observed_at":"2026-08-05T00:19:11.956367Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:12.070497Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.070497Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:ca6068d6d13dc4e5d9b9ecc53fc7710146fab07ca2cbfde3cbe36a91a3c42c6c","observation_id":"b20ad193-e0e7-4071-a486-71c6e07e7445","resolution":{"observed_at":"2026-08-05T00:19:12.070497Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2005.01643","last_updated":"2020-11-01T23:50:25Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-05-04T17:00:15Z","title":"Offline Reinforcement Learning: Tutorial, Review, and Perspectives on Open Problems","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.01643","snapshot_observed_at":"2026-08-05T00:19:12.211566Z","title":"arXiv preprint arXiv:2005.01643 , year=","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.211566Z"},"links":{"cited_paper":"/paper/2005.01643","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:4f42e3e7c66f48dc08baf018cbcea8cd56d9de87a9c3e8db5d8e0fc99a6bfc29","observation_id":"528aa86a-2d76-45d3-bf5a-4dad349ed9c8","resolution":{"observed_at":"2026-08-05T00:19:12.211566Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-05T00:19:12.292009Z","title":"arXiv preprint arXiv:1707.06347 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.292009Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:1bccd108ad13d440e76c2118ea2c69a87a29d54b3699881b8bbda82573117478","observation_id":"9e6dcee5-181b-44d7-88a8-7d6421722f5a","resolution":{"observed_at":"2026-08-05T00:19:12.292009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-05T00:19:12.295874Z","title":"arXiv preprint arXiv:2402.03300 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.295874Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:b1a22c91d68c9973691e8ed371e62b047e7c2a45ea6157e77b790a438efb066a","observation_id":"64151d39-3180-47d6-930c-a763b193f0c0","resolution":{"observed_at":"2026-08-05T00:19:12.295874Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.11431","last_updated":"2023-04-26T22:49:15Z","snapshot_observed_at":"2026-07-06T14:33:47.984474Z","submitted_at":"2022-12-22T00:47:40Z","title":"Local Policy Improvement for Recommender Systems","version":2},"cited_work":{"arxiv_id":"2212.11431","doi":null,"metadata_source":"pith","pith_arxiv_id":"2212.11431","snapshot_observed_at":"2026-08-05T00:19:15.260998Z","title":"Local Policy Improvement for Recommender Systems","venue":"cs.LG","work_id":"9a3660d1-b647-48d2-bad3-615d3465d58c","year":2022},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.368388Z"},"links":{"cited_paper":"/paper/2212.11431","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:763412271bb149ae9d8ff2718317d5af5e73faad3d4d3f76f913c68167e3136d","observation_id":"ba916abc-18dc-4c4f-b922-4a3fdd784c66","resolution":{"observed_at":"2026-08-05T00:19:15.338670Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:12.466083Z","title":"Proceedings of the 10th ACM conference on recommender systems , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.466083Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:630c51926de48f4a5fd1acc89bb53c52cd1611d372130e1058e46df4f3a6dac2","observation_id":"5f2c86f5-9b0a-4b96-b9af-bfaf21d7ebbd","resolution":{"observed_at":"2026-08-05T00:19:12.466083Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2206.02620","last_updated":"2023-06-16T08:38:52Z","snapshot_observed_at":"2026-07-06T13:17:51.474784Z","submitted_at":"2022-06-01T02:45:35Z","title":"ResAct: Reinforcing Long-term Engagement in Sequential Recommendation with Residual Actor","version":2},"cited_work":{"arxiv_id":"2206.02620","doi":null,"metadata_source":"pith","pith_arxiv_id":"2206.02620","snapshot_observed_at":"2026-08-05T00:19:15.109060Z","title":"ResAct: Reinforcing Long-term Engagement in Sequential Recommendation with Residual Actor","venue":"cs.IR","work_id":"132a95ac-8c84-456e-91b3-98a55904d8fa","year":2022},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.591937Z"},"links":{"cited_paper":"/paper/2206.02620","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:969179a14fc5f578f5ac1d6d1c1b62886874642317e64760219aca923d21556a","observation_id":"dc58d33e-6c97-4e4b-a486-51e9b9888aab","resolution":{"observed_at":"2026-08-05T00:19:15.168696Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:16.569701Z","title":"The Journal of Machine Learning Research , volume=","venue":null,"work_id":"b7844d25-fc59-41ec-8fb9-5f2222b057c3","year":2015},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.679862Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:e8ffbaa8132001f2a1e6957757daa8efab0176c6aade47bfba7f61051e4f5993","observation_id":"17e6eb35-c945-4ad1-89d1-579f4e12468c","resolution":{"observed_at":"2026-08-05T00:19:16.683494Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1103.4601","last_updated":"2011-05-06T02:38:18Z","snapshot_observed_at":"2026-07-06T02:25:17.083332Z","submitted_at":"2011-03-23T19:37:45Z","title":"Doubly Robust Policy Evaluation and Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1103.4601","snapshot_observed_at":"2026-08-05T00:19:12.758286Z","title":"arXiv preprint arXiv:1103.4601 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.758286Z"},"links":{"cited_paper":"/paper/1103.4601","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:547c0382ace00d1975936128a630bf7327967b84bba83546238339f8341cc6d7","observation_id":"6991502d-713a-42ba-b217-4212fcdcc6ed","resolution":{"observed_at":"2026-08-05T00:19:12.758286Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:12.848680Z","title":"1998 , publisher=","venue":null,"work_id":null,"year":1998},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.848680Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:30890877508bc05a61bb4c5091491ab629e5f0fa9e2bb655d43b2bcf50c21617","observation_id":"4f387d92-364e-41f3-a1e8-1470086c1582","resolution":{"observed_at":"2026-08-05T00:19:12.848680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1512.07679","last_updated":"2016-04-04T11:27:36Z","snapshot_observed_at":"2026-08-03T08:37:56.722087Z","submitted_at":"2015-12-24T01:31:40Z","title":"Deep Reinforcement Learning in Large Discrete Action Spaces","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1512.07679","snapshot_observed_at":"2026-08-05T00:19:12.929381Z","title":"arXiv preprint arXiv:1512.07679 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:12.929381Z"},"links":{"cited_paper":"/paper/1512.07679","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:eeb7b02181795c0771387ce18b4fe913133cddac5b6618cf3f5913fb20efe4e6","observation_id":"198f9939-aeb8-4fb9-b68a-a5813b2e7b4b","resolution":{"observed_at":"2026-08-05T00:19:12.929381Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:16.407679Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"630fa6d6-896f-487d-beea-a05525488c94","year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.010008Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:ff1198853c536bab42c333b4fed7c0ac22912de608df06aef657c73e69dee6e2","observation_id":"44e3588a-c9d6-446d-8267-5959ef8959d5","resolution":{"observed_at":"2026-08-05T00:19:16.460164Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1910.00177","last_updated":"2019-10-07T20:23:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2019-10-01T02:23:38Z","title":"Advantage-Weighted Regression: Simple and Scalable Off-Policy Reinforcement Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1910.00177","snapshot_observed_at":"2026-08-05T00:19:13.088764Z","title":"arXiv preprint arXiv:1910.00177 , year=","venue":null,"work_id":null,"year":1910},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.088764Z"},"links":{"cited_paper":"/paper/1910.00177","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:425437349020f50c66028701c0fa973d6b34bd518037af9089227a42f3122f28","observation_id":"44ed47ad-1fc1-4067-9902-6b0d347dee88","resolution":{"observed_at":"2026-08-05T00:19:13.088764Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2006.09359","last_updated":"2021-04-24T22:39:30Z","snapshot_observed_at":"2026-07-06T09:29:45.475911Z","submitted_at":"2020-06-16T17:54:41Z","title":"AWAC: Accelerating Online Reinforcement Learning with Offline Datasets","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2006.09359","snapshot_observed_at":"2026-08-05T00:19:13.185095Z","title":"arXiv preprint arXiv:2006.09359 , year=","venue":null,"work_id":null,"year":2006},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.185095Z"},"links":{"cited_paper":"/paper/2006.09359","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:43e641c0a6a3c309fd29126e5d4e5c342d93b4953b37f59ebe6429bcf52144ae","observation_id":"85636644-4893-436f-8d69-8016335f7797","resolution":{"observed_at":"2026-08-05T00:19:13.185095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:13.307778Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.307778Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:d34f9921297b51f8fd393122cda20cf906f43a5acc248fd93b5027b200aac21e","observation_id":"616e8515-de3d-4d1c-b101-39b9d14b4924","resolution":{"observed_at":"2026-08-05T00:19:13.307778Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:13.406677Z","title":"International Conference on Artificial Intelligence and Statistics , pages=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.406677Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:2d17f7605c04c2798ffee9be181b259d8a1d85852e8acddc344513a9b3e27a19","observation_id":"4d122eb5-4151-4b83-9c26-315927ce4600","resolution":{"observed_at":"2026-08-05T00:19:13.406677Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.14740","last_updated":"2024-02-26T18:26:25Z","snapshot_observed_at":"2026-08-09T14:30:33.899591Z","submitted_at":"2024-02-22T17:52:34Z","title":"Back to Basics: Revisiting REINFORCE Style Optimization for Learning from Human Feedback in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.14740","snapshot_observed_at":"2026-08-05T00:19:13.504276Z","title":"arXiv preprint arXiv:2402.14740 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.504276Z"},"links":{"cited_paper":"/paper/2402.14740","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:2fdfdbabfd8677c7ca6bf04df200fc56a910b5575b12ad574cce2a7bfe0cdc69","observation_id":"74d43e1a-d897-41df-8e2b-2c61583264dd","resolution":{"observed_at":"2026-08-05T00:19:13.504276Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14718","last_updated":"2024-04-20T00:34:59Z","snapshot_observed_at":"2026-07-06T15:32:14.701994Z","submitted_at":"2023-05-24T04:42:17Z","title":"Leftover Lunch: Advantage-based Offline Reinforcement Learning for Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14718","snapshot_observed_at":"2026-08-05T00:19:13.584966Z","title":"arXiv preprint arXiv:2305.14718 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.584966Z"},"links":{"cited_paper":"/paper/2305.14718","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:88bb48021607a441ed3711f66df5c61079991086ad5f209f3fc47c205d167eec","observation_id":"f4729e51-fd38-4421-92ca-9cad72d06ebc","resolution":{"observed_at":"2026-08-05T00:19:13.584966Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:13.674050Z","title":"The Thirty-ninth Annual Conference on Neural Information Processing Systems , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.674050Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:3a412c401906baa1705f547222b8dc0b3369e232c4b9fae2c45925ee51fe6049","observation_id":"fdba5dbc-4326-47d6-a6b0-5bebe7ac1b64","resolution":{"observed_at":"2026-08-05T00:19:13.674050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:13.787238Z","title":"Acm transactions on interactive intelligent systems (tiis) , volume=","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.787238Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:2f3656b7af5ea5115aad4f687ce1a1881eefad17a98c4605a93d0d6f4c2d7200","observation_id":"95bd82d2-ea9b-42cc-9b35-27aa58f7ffe5","resolution":{"observed_at":"2026-08-05T00:19:13.787238Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:16.191463Z","title":"Proceedings of the 38th international ACM SIGIR conference on research and development in information retrieval , pages=","venue":null,"work_id":"d5303f62-a92b-4013-a1fd-bf2b9bf618d1","year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.891636Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:6fe2bfec0cc864f972df920af7aedd132a4092256ff37c2c0194d4c4d7185bf7","observation_id":"c8fe50c8-b796-4a01-98bf-6f0a983112bb","resolution":{"observed_at":"2026-08-05T00:19:16.257418Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:16.009074Z","title":"Proceedings of the 16th ACM SIGKDD international conference on Knowledge discovery and data mining , pages=","venue":null,"work_id":"e61cdc54-90b6-462e-94d6-7481a954088f","year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:13.988717Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:e28c8652eab1233f450e6f5a51066da7db1b7fdca34f34551380d74d59ad715a","observation_id":"71be91dd-66ea-499d-9c11-fd1093792839","resolution":{"observed_at":"2026-08-05T00:19:16.114661Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:14.035733Z","title":"international conference on machine learning , pages=","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.035733Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:e99aa8714d7f328f5aeebfa1127ac5fe247b7b53adb24094b913e68dc13721ea","observation_id":"10a7a1bd-3bbb-48ac-a4c4-c44a9b1bad0b","resolution":{"observed_at":"2026-08-05T00:19:14.035733Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2110.06169","last_updated":"2021-10-12T17:05:05Z","snapshot_observed_at":"2026-08-06T15:42:21.967989Z","submitted_at":"2021-10-12T17:05:05Z","title":"Offline Reinforcement Learning with Implicit Q-Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.06169","snapshot_observed_at":"2026-08-05T00:19:14.216356Z","title":"arXiv preprint arXiv:2110.06169 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.216356Z"},"links":{"cited_paper":"/paper/2110.06169","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:3218095e052c023103f998650941101bb8c686b199f582d084147b4ccd885909","observation_id":"f922640e-028d-42a1-b6ec-6e3cb44bfbd1","resolution":{"observed_at":"2026-08-05T00:19:14.216356Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:15.822418Z","title":"2007 IEEE International Symposium on Approximate Dynamic Programming and Reinforcement Learning , pages=","venue":null,"work_id":"3bac2333-3dde-4e49-8db7-fa6ca25ce353","year":2007},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.300165Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:91bb008ec193d27013748dd3cfa6fd5a5a12b0556ae65acb89884c252a11d0a7","observation_id":"608e2981-d26b-4b75-9a27-e1d918845277","resolution":{"observed_at":"2026-08-05T00:19:15.884769Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:14.370525Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.370525Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:b8c83021fa738a295f3eedea94c6b8e5ed9dad344117e2bf889090edc7683f8f","observation_id":"9b7f5ebf-58c7-4f58-8ece-ecd70ef2026e","resolution":{"observed_at":"2026-08-05T00:19:14.370525Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:15.640110Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"38762667-815d-44bc-9e95-db37a996deb2","year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.429470Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:68b29cb7e6d92973955da30256c89ca30143246dc37113493ddb350a7db22d0e","observation_id":"8282fac1-639a-4950-b16e-894fea6c321c","resolution":{"observed_at":"2026-08-05T00:19:15.703841Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.00168","last_updated":"2024-02-02T03:41:50Z","snapshot_observed_at":"2026-07-06T16:41:27.099792Z","submitted_at":"2023-10-31T21:52:41Z","title":"The Alignment Ceiling: Objective Mismatch in Reinforcement Learning from Human Feedback","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.00168","snapshot_observed_at":"2026-08-05T00:19:14.540226Z","title":"arXiv preprint arXiv:2311.00168 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.540226Z"},"links":{"cited_paper":"/paper/2311.00168","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:1b85a3dafb79c7bd677a773cd649c9e5714514d7ebb21aacdc06d4f2e3536800","observation_id":"82d42bf6-2cca-4297-ae8f-a32229b733be","resolution":{"observed_at":"2026-08-05T00:19:14.540226Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2201.03544","last_updated":"2022-02-14T09:05:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-01-10T18:58:52Z","title":"The Effects of Reward Misspecification: Mapping and Mitigating Misaligned Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.03544","snapshot_observed_at":"2026-08-05T00:19:14.655825Z","title":"arXiv preprint arXiv:2201.03544 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.655825Z"},"links":{"cited_paper":"/paper/2201.03544","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:bd6b4a93148f8f10b2438a44960a11707a027e58ce7b4233739b8126ba965d27","observation_id":"686cfe5c-6d1e-450a-bbea-0f47829103fe","resolution":{"observed_at":"2026-08-05T00:19:14.655825Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:14.736605Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.736605Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:06cc9d55b66b4524e81aa3db365b7d63ac8d2e8768bf9765ba336a794c2e21de","observation_id":"038b2e80-d69c-482d-9f6b-104954b98d58","resolution":{"observed_at":"2026-08-05T00:19:14.736605Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T00:19:14.802419Z","title":"Machine learning , volume=","venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.802419Z"},"links":{"citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:f871bfdf116ca9260b070335fe8cb769ec52e9c4d2245ddc22d00650632d0ab9","observation_id":"9e2803f2-28f1-4b10-8431-2b1b00e6c358","resolution":{"observed_at":"2026-08-05T00:19:14.802419Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1606.06565","last_updated":"2016-07-25T17:23:29Z","snapshot_observed_at":"2026-07-06T05:00:46.434335Z","submitted_at":"2016-06-21T13:37:05Z","title":"Concrete Problems in AI Safety","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.06565","snapshot_observed_at":"2026-08-05T00:19:14.894582Z","title":"arXiv preprint arXiv:1606.06565 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-05T00:19:14.894582Z"},"links":{"cited_paper":"/paper/1606.06565","citing_paper":"/paper/2608.00816"},"observation_digest":"sha256:d8dc99123298770528152a72c86eaa4c5d8896c0357ce9ccb2fe9b0772b8776a","observation_id":"55dc147b-2d7b-4e08-a9d3-90bcf86d6fe8","resolution":{"observed_at":"2026-08-05T00:19:14.894582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2608.00816","last_updated":"2026-08-01T18:48:39Z","latest_version":1,"primary_category":"cs.IR","snapshot_observed_at":"2026-08-09T14:47:11.614545Z","submitted_at":"2026-08-01T18:48:39Z","title":"Exponential Reward Weighting for Fine-Tuning Generative Recommenders under Sparse and Noisy Feedback"},"reference_resolution":{"displayed":42,"state_counts":{"malformed_identifier":0,"metadata_mismatch":3,"parse_uncertain":0,"unresolved":29,"verified_exact":0,"verified_fuzzy":10},"total_outbound_references":42},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 42 of 42 outbound references and 0 inbound Pith citation observations for arXiv:2608.00816."}