{"as_of":"2026-08-14T00:34:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:c731782fdcb46350ef715ff7015fed04b6f9b8cba38250bd52b27232e98a388a","coverage":[{"denominator":34,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":34,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T01:10:58.285736Z","state":"measured"},{"denominator":34,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":34,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.11912/citation-record","integrity":"/paper/2506.11912/integrity","json":"/paper/2506.11912/citation-record.json","paper":"/paper/2506.11912"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.510215Z","title":null,"venue":null,"work_id":"177ae22e-810a-4e11-93f9-8ddaa1b467b5","year":1994},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.203368Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:f75962478aec0d21f6c7987d10f6fbe0feba9c9502fe37b143f2a295259ad2df","observation_id":"46dbb151-be58-4093-b9e6-3df113f9e72b","resolution":{"observed_at":"2026-08-07T01:11:01.513026Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.500967Z","title":null,"venue":null,"work_id":"2ccbdac1-9be0-41b6-9646-0d15f1f974b2","year":2001},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.291277Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:a9492e7e0058239b14bbadfcc260d4d141b08a7729e6fac962dae1eea1b5e4f9","observation_id":"7fa74118-4a6f-4053-85ca-dbe902c715d9","resolution":{"observed_at":"2026-08-07T01:11:01.503874Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.491129Z","title":null,"venue":null,"work_id":"9ad7ea5a-c0b4-44ed-bb94-8b0e543a0fe3","year":1999},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.396925Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:7785640ab911152071930ad349f933a0dd369299c2ef1546a13d18e8edd7bde8","observation_id":"a48bfeab-5b40-44cf-afa0-0288a6cf3763","resolution":{"observed_at":"2026-08-07T01:11:01.494683Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.481542Z","title":"C., and Le Roux, N","venue":null,"work_id":"9f6e1ed4-bb1d-4f78-bd4a-ad537201d97f","year":2021},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.440340Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:7fa400e07b0c7d4cfbd40b431221c6237b866b9ea660cf0a126444455543e075","observation_id":"472a5fd2-fc47-4978-b20a-eb5ec2ce3451","resolution":{"observed_at":"2026-08-07T01:11:01.484545Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.472409Z","title":"and Vicente, R","venue":null,"work_id":"5f2eb479-dea6-4774-be79-d59ecf55353b","year":2022},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.553966Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:81cab03cd7e43fd633e996b668ae40324166182d6312d546f240ec9985e1941d","observation_id":"774533e8-5f8b-42e7-8098-6098615d70c0","resolution":{"observed_at":"2026-08-07T01:11:01.475565Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.462661Z","title":null,"venue":null,"work_id":"251ea4fa-dbb9-4a52-bb7d-916abda65d5a","year":2019},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.654464Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:e6fbf9a3124d5706ac78292fa662673c16cc74e513fca8bb159a71ceedf85f46","observation_id":"17b54c4f-3a44-4c3d-9591-9fb04aa295dd","resolution":{"observed_at":"2026-08-07T01:11:01.465631Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.452675Z","title":null,"venue":null,"work_id":"e1327bd4-a4a9-45bf-81db-4c1fc5df5880","year":2023},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.748944Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:5d59917f8de32913c22183862508e6aa1a3b1b0aa4cf89c2f8b0aebfe0602c92","observation_id":"008efe3d-cae2-4101-bb9c-a4cd6fb11ebf","resolution":{"observed_at":"2026-08-07T01:11:01.455751Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.442660Z","title":"G., and Pineau, J","venue":null,"work_id":"4e1ccf9d-fc8f-4978-9d9e-6518b321fb70","year":2018},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.860801Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:61dfb35df2bcee0d5f54629074b58f7e40a1b49bfa5c35fe9c69f4bc889fde39","observation_id":"0d2581ef-e31a-46d5-97d7-76db9bcee0f1","resolution":{"observed_at":"2026-08-07T01:11:01.446104Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.433140Z","title":null,"venue":null,"work_id":"2e82aa65-2851-464d-be4b-26d03ce7241c","year":2001},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:55.947805Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:bc890f12b07dfbe7e9430fd634bc462a62b461e9192729d3a8c693a89074056a","observation_id":"72333d80-22e4-48dc-be10-5598b5e2e5f2","resolution":{"observed_at":"2026-08-07T01:11:01.436266Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.423770Z","title":"M., de Vries, J","venue":null,"work_id":"452186ae-b64c-4e9f-834b-86b1b0066d9d","year":2024},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.026758Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:b227c5cac74e4292033eccaa8844aceaae3e1763b38d3be176881d8b893aee74","observation_id":"0cffefbb-941c-42ae-b72d-98555e55efcd","resolution":{"observed_at":"2026-08-07T01:11:01.426839Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.413392Z","title":null,"venue":null,"work_id":"2f6bab60-badb-43a3-872b-3c0d1acf37d4","year":2017},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.120102Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:839fb868cbcaa02b52356868e76db5e17464c5abeb45f991a1b926ed582cae1d","observation_id":"0cb95045-92ff-4689-ad7b-f637f7d868e0","resolution":{"observed_at":"2026-08-07T01:11:01.416975Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.402628Z","title":null,"venue":null,"work_id":"350c5ab0-46d3-4713-8678-3af98f49f876","year":2023},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.240654Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:a9edc270744bc9b07ef51a3777457cb652a51d8dea837584f491e336c67d24cf","observation_id":"79f77830-87ff-4e15-a2d8-af1c07d25c23","resolution":{"observed_at":"2026-08-07T01:11:01.406013Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.392632Z","title":null,"venue":null,"work_id":"ae94af07-062b-4e2d-b348-733ad6584158","year":2023},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.297251Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:273d478c29f431fcbf85dcc2caf954f6ea9ee61fffd9bb71dd26f97c901e6d06","observation_id":"1e060f05-7894-4cfe-9ce1-b6b35a67ce2b","resolution":{"observed_at":"2026-08-07T01:11:01.396050Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.381649Z","title":"C., Bellemare, M","venue":null,"work_id":"081fda7c-be4f-47b2-93e5-5d8671913862","year":2018},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.401160Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:07fdfb63d98210b74566a6f854413bac627bd2401ff7cee3f47e7eb881619432","observation_id":"4f4a59dc-a71d-4fe4-99b9-c11ff7f435a2","resolution":{"observed_at":"2026-08-07T01:11:01.385105Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.371622Z","title":null,"venue":null,"work_id":"6f43bdd9-d83c-4ec9-a52a-7b7a6ea1aa12","year":2017},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.520833Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:9330daf4dc7566e4f8de29348efeefc63379b9c6214c818af5a83aee5dec6fc2","observation_id":"48fc7080-b5d6-4796-a17a-bb22c0c8f1d5","resolution":{"observed_at":"2026-08-07T01:11:01.374703Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.359605Z","title":"and Tsitsiklis, J","venue":null,"work_id":"61481729-a676-478c-bda6-b6968c884488","year":1999},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.607417Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:bf350f25bee69fd60785b965139fe7608108bc351e839d4700d40449bb37e9ed","observation_id":"0b379814-440b-41ee-880b-86b21970ceff","resolution":{"observed_at":"2026-08-07T01:11:01.363244Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:01.083368Z","title":null,"venue":null,"work_id":"d8c5ff14-ee94-40ea-99ec-aaec058fc546","year":1995},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.649485Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:273b46a81cf5a56aef867e3f1f72a363ace5753e0213a188230b051734b0da44","observation_id":"039bf625-d68f-45b1-9c64-7387d9bc8109","resolution":{"observed_at":"2026-08-07T01:11:01.230910Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:00.800353Z","title":null,"venue":null,"work_id":"f25904bc-ad48-4b42-bc14-e94ff4c56419","year":2022},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.755411Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:730738333f44fc7f159fec04bfa810ee8d137df2b2500007d0588307251cc3a6","observation_id":"626055d5-cb09-4b3d-8de2-b6c8b386213b","resolution":{"observed_at":"2026-08-07T01:11:00.918276Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:00.550564Z","title":"u rtler, N., Neitz, A., and Sch \\","venue":null,"work_id":"17f29777-3351-42f4-ad5e-eec18dfcc6e1","year":2022},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.868913Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:c2c31b930749d5d6116ef10a027cbfbd6495d48adea4506db963d45c46745741","observation_id":"d49e9354-15ef-4124-9211-41b03af524b9","resolution":{"observed_at":"2026-08-07T01:11:00.669845Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:00.324269Z","title":"and Sch \\\"o lkopf, B","venue":null,"work_id":"9a80fad2-78ac-4e48-89da-4c8918284acb","year":2024},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:56.938340Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:067ee72a89e0367062a94807b39ab32ee6413b8942d320a5b149ad2073810086","observation_id":"e56ff290-4238-4c07-af89-ff79ed58c5b3","resolution":{"observed_at":"2026-08-07T01:11:00.441155Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:11:00.141659Z","title":null,"venue":null,"work_id":"abbecb45-715c-46cf-8b08-645200974019","year":2016},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.021373Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:17a34c59c84a9ccf5a396fc3d5358e741f3fa7dddf3549c9f3fd149fc61b81a0","observation_id":"5f39d1e9-b2f7-4a8e-be39-4eba7c95d96b","resolution":{"observed_at":"2026-08-07T01:11:00.240560Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:57.146209Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.146209Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:5592bfafa5fc801c498804c42241957ea0fe4a421a54300333b1629e40e030b5","observation_id":"8aa180b3-4a14-459b-9713-3c19c869c926","resolution":{"observed_at":"2026-08-07T01:10:57.146209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.918876Z","title":"and Fergus, R","venue":null,"work_id":"c4729972-4f65-41a0-a677-4c8d9599bb9c","year":2021},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.288542Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:0319247d8fbe3c09e3157d09682a5cb57bd87676dbbd469297e32449e21ca4d0","observation_id":"41569b00-485d-4119-b71a-356d554fc190","resolution":{"observed_at":"2026-08-07T01:11:00.030784Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-07T01:10:57.370198Z","title":null,"venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.370198Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:88cfbeb5727b466a864e01556ed8911eb32664db18e28bb6413e755db2ed4412","observation_id":"7591ef08-5a22-4c0a-adc0-723e66f4365e","resolution":{"observed_at":"2026-08-07T01:10:57.370198Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.741513Z","title":null,"venue":null,"work_id":"e28c9ae2-2ce4-452b-aa73-7f45011e58c7","year":2020},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.503345Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:6136015b976903159dee8768284cbe68d05660b772851fc4338bda5a6c4bcd6c","observation_id":"97447029-0d80-4da3-8459-ab750a3cc7ba","resolution":{"observed_at":"2026-08-07T01:10:59.838039Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.522454Z","title":null,"venue":null,"work_id":"ef6dd7d2-6ad5-481d-91eb-5adad625b356","year":2024},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.574995Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:1a4d8f37d8f29bdc2ada99a0508c9b454047c3f04a24ba540e42bf3b349b9bc2","observation_id":"87f7db3e-fe1f-44e8-989c-25d07d12a3e9","resolution":{"observed_at":"2026-08-07T01:10:59.624083Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:57.647845Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.647845Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:dd009f4c631338330d3c54a1336393592ae01be5946d7304450df9b5e8277f90","observation_id":"396052ca-a9e3-49f0-b9cf-d44422daf8f3","resolution":{"observed_at":"2026-08-07T01:10:57.647845Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.230776Z","title":"S., McAllester, D., Singh, S., and Mansour, Y","venue":null,"work_id":"974821b0-2394-476b-87f2-2b55b8fa4752","year":1999},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.732333Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:8a083622b9acf8e2d525d8983920a2eabc65dd71c4d983e09a2ae3ae707a1de8","observation_id":"bb8be8b2-ca3f-430f-a7c7-428885eb974d","resolution":{"observed_at":"2026-08-07T01:10:59.373570Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:59.016376Z","title":"and Parr, R","venue":null,"work_id":"7e93fab5-09a4-4852-9815-52d7d835f736","year":2009},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.787899Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:37b0aff0a318c0310afca127f137581ed26dac1c50f5b337db1666612157d564","observation_id":"0f2b8c9d-558a-43f8-9073-42b7b6129aa6","resolution":{"observed_at":"2026-08-07T01:10:59.149889Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.06539","last_updated":"2021-10-13T07:31:31Z","snapshot_observed_at":"2026-08-13T17:54:06.610792Z","submitted_at":"2021-10-13T07:31:31Z","title":"On Covariate Shift of Latent Confounders in Imitation and Reinforcement Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.06539","snapshot_observed_at":"2026-08-07T01:10:57.900276Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:57.900276Z"},"links":{"cited_paper":"/paper/2110.06539","citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:8f682f84e30d6832c69a817982b8a7259e4d8c00d50ae587d0b9378670d34ac7","observation_id":"91a90401-e57a-4cc7-be2d-1249da362344","resolution":{"observed_at":"2026-08-07T01:10:57.900276Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03565","last_updated":"2026-07-13T10:54:37Z","snapshot_observed_at":"2026-08-12T22:30:36.866820Z","submitted_at":"2024-10-04T16:15:31Z","title":"Training on Irrelevant States Implies Data Augmentation: Generalization in Contextual MDPs","version":4},"cited_work":{"arxiv_id":"2410.03565","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.03565","snapshot_observed_at":"2026-08-07T01:10:58.387718Z","title":"Training on Irrelevant States Implies Data Augmentation: Generalization in Contextual MDPs","venue":"cs.LG","work_id":"b5b5094e-6af4-45b2-ba1e-a03d1b7d4a6c","year":2024},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:58.002804Z"},"links":{"cited_paper":"/paper/2410.03565","citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:de1497ae4b41b863ad5999aa6e0ab69347b1c2acad682672e3c4221382cd2228","observation_id":"a0f92a0c-aec2-4371-a52b-d10353729270","resolution":{"observed_at":"2026-08-07T01:10:58.465709Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:58.106871Z","title":null,"venue":null,"work_id":null,"year":1992},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:58.106871Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:347cfe2eaadf66a8b4c199b5288c2784361e5d8143ebec226493d9330f11fd54","observation_id":"91cc7418-1aae-4ac9-addb-042e5ee8f2ce","resolution":{"observed_at":"2026-08-07T01:10:58.106871Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:58.775263Z","title":null,"venue":null,"work_id":"9e004dea-b842-49bb-acbb-458168caa45d","year":2020},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:58.218410Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:3db57396743c70cd6f7bb60eac1f69bcbb5abd076461c986a70527e353c0510e","observation_id":"d9267978-3bed-40d2-b067-32ec803fd44e","resolution":{"observed_at":"2026-08-07T01:10:58.896922Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T01:10:58.597029Z","title":null,"venue":null,"work_id":"dec52605-b577-4a00-bb45-1caf2c513eae","year":2020},"citing_paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T01:10:58.285736Z"},"links":{"citing_paper":"/paper/2506.11912"},"observation_digest":"sha256:b018f06e5f74f603d1924dc5f708d9f8d3206b8b0e5a112d8b8b467800575b16","observation_id":"1367568e-2ac0-45e7-9829-ca83dbfddbc0","resolution":{"observed_at":"2026-08-07T01:10:58.684426Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.11912","last_updated":"2025-06-13T16:06:47Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-13T04:43:35.923055Z","submitted_at":"2025-06-13T16:06:47Z","title":"Breaking Habits: On the Role of the Advantage Function in Learning Causal State Representations"},"reference_resolution":{"displayed":34,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":22,"verified_exact":0,"verified_fuzzy":11},"total_outbound_references":34},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 34 of 34 outbound references and 0 inbound Pith citation observations for arXiv:2506.11912."}