{"as_of":"2026-08-14T06:49:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:adbccb1d4c8299929cd45757532ba0df854d6112d226eaeeddb23bf64f7f6c4a","coverage":[{"denominator":50,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":50,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T23:27:32.794315Z","state":"measured"},{"denominator":50,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":50,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.09109/citation-record","integrity":"/paper/2608.09109/integrity","json":"/paper/2608.09109/citation-record.json","paper":"/paper/2608.09109"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2022.findings-emnlp.296","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.044972Z","title":null,"venue":null,"work_id":"9170cb79-abbf-4a13-bbec-735e293da05a","year":2022},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.605729Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:dfcab81a2b07a76116618aa99a6921fb139d8ca9ce6d68d4f069a3966da28eae","observation_id":"3760904d-7100-4036-b05c-1f494382414b","resolution":{"observed_at":"2026-08-11T23:27:33.048510Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2510.17281","last_updated":"2026-06-03T04:15:09Z","snapshot_observed_at":"2026-08-11T13:20:24.725773Z","submitted_at":"2025-10-20T08:16:12Z","title":"MemoryBench: A Benchmark for Memory and Continual Learning in LLM Systems","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2510.17281","snapshot_observed_at":"2026-08-11T23:27:32.610040Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.610040Z"},"links":{"cited_paper":"/paper/2510.17281","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:2756bbd30179cf23e449890078b72ae6687a8fe6688885e9055ffc6b0e12fd7f","observation_id":"4fcf339f-deca-450b-9596-1f0e2439ffc2","resolution":{"observed_at":"2026-08-11T23:27:32.610040Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.614542Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.614542Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:771f0a0a813e242ee567083fdb4ee1ac617f1415b747aaf0ea6037904a172610","observation_id":"3d69960b-10b2-406f-a378-dcc050bc3ea2","resolution":{"observed_at":"2026-08-11T23:27:32.614542Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2026.acl-long.441","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.027940Z","title":null,"venue":null,"work_id":"e22ef603-1225-4dd0-87ff-1a451f8ec33a","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.617845Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:f26f18beed2b36e88a2cabf543f5adc143e30275d09e9c17ee1dd6b63e2195dc","observation_id":"9e44dc09-ade2-467d-88f0-60caa9a84b51","resolution":{"observed_at":"2026-08-11T23:27:33.033154Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2025.acl-long.1200","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.008650Z","title":null,"venue":null,"work_id":"609d43ff-eb88-41b6-ad07-27cbfb2a66d4","year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.621106Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:3aeae313c16c3fc7ab5193e090c35ba85890d76ef37ef5c2d162980afea91d86","observation_id":"c8b5931f-6879-4638-a8ef-5745600d941b","resolution":{"observed_at":"2026-08-11T23:27:33.012175Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.04475","last_updated":"2025-03-10T09:27:03Z","snapshot_observed_at":"2026-07-06T17:56:23.317089Z","submitted_at":"2024-04-06T02:29:02Z","title":"Length-Controlled AlpacaEval: A Simple Way to Debias Automatic Evaluators","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.04475","snapshot_observed_at":"2026-08-11T23:27:32.625100Z","title":"Hashimoto","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.625100Z"},"links":{"cited_paper":"/paper/2404.04475","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:a35ff63d98c80383722e5301e786d370c604564bd871b6128090819789640c0d","observation_id":"862f1f3d-120c-4225-a81f-ffb6d140653d","resolution":{"observed_at":"2026-08-11T23:27:32.625100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.09685","last_updated":"2021-10-16T18:40:34Z","snapshot_observed_at":"2026-08-11T08:20:29.798517Z","submitted_at":"2021-06-17T17:37:18Z","title":"LoRA: Low-Rank Adaptation of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.09685","snapshot_observed_at":"2026-08-11T23:27:32.628831Z","title":"Hu, Yelong Shen, Phillip Wallis, Zeyuan Allen-Zhu, Yuanzhi Li, Shean Wang, Lu Wang, and Weizhu Chen","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.628831Z"},"links":{"cited_paper":"/paper/2106.09685","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:62b3c4b2e39c486a5437e24c4b4aab4f560eebe7cb06946f73b109283db1c1f1","observation_id":"e68d4db8-a732-4991-9c3e-ac967e7fe43c","resolution":{"observed_at":"2026-08-11T23:27:32.628831Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2601.20802","last_updated":"2026-02-16T14:49:34Z","snapshot_observed_at":"2026-08-13T15:28:29.238641Z","submitted_at":"2026-01-28T17:45:12Z","title":"Reinforcement Learning via Self-Distillation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.20802","snapshot_observed_at":"2026-08-11T23:27:32.632709Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.632709Z"},"links":{"cited_paper":"/paper/2601.20802","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:98d24815b46d89a2e3b18acaed80d8503391cb2a0a4885bf0fa189e2cfdae18e","observation_id":"3262af93-f0ae-4b20-8a6f-7c34cc069583","resolution":{"observed_at":"2026-08-11T23:27:32.632709Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.636081Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.636081Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:399cc60fe04ec0a931721d5492446f09c0ea77c8e99507d5e8f01980e821bf39","observation_id":"f17ab736-15e5-40b3-b4d8-d1f9de2db81f","resolution":{"observed_at":"2026-08-11T23:27:32.636081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.639753Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.639753Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:afd1042f14261eebe3302604f169a6da2c93ab413e8cfd1d5d8f7b078f56e0f2","observation_id":"cb9034fa-85a7-449b-acee-a4d13f34b82d","resolution":{"observed_at":"2026-08-11T23:27:32.639753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2026.acl-long.1831","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.976005Z","title":null,"venue":null,"work_id":"5b540574-94a9-4fcd-9467-d4efba3c0069","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.642822Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:7c240181460bd93ce0025ea1a484258f0d66a79002e457286c0282dbda604998","observation_id":"07b71941-e55f-40cb-b877-9fc5233a429d","resolution":{"observed_at":"2026-08-11T23:27:32.979903Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2601.22900","last_updated":"2026-05-30T11:40:38Z","snapshot_observed_at":"2026-08-09T01:46:22.354879Z","submitted_at":"2026-01-30T12:19:54Z","title":"MulFeRL: Enhancing Reinforcement Learning with Verbal Feedback in a Multi-turn Loop","version":2},"cited_work":{"arxiv_id":"2601.22900","doi":null,"metadata_source":"pith","pith_arxiv_id":"2601.22900","snapshot_observed_at":"2026-08-11T23:27:33.239830Z","title":"MulFeRL: Enhancing Reinforcement Learning with Verbal Feedback in a Multi-turn Loop","venue":"cs.AI","work_id":"31a684ff-715a-4f20-a491-249bc316e15e","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.645757Z"},"links":{"cited_paper":"/paper/2601.22900","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:bf15fe720361072d1894ae400f174fb9032afab5fd46041e5d974f1f1726a42e","observation_id":"06f62b64-f665-46e8-a1f2-1cfca9c3763d","resolution":{"observed_at":"2026-08-11T23:27:33.244435Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.571950Z","title":null,"venue":null,"work_id":"3d7f71b8-52c0-4d98-a9dc-cc730e7f8d05","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.648978Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:a2d023082811b29106feac8fb86e7e36e9984f1cb79ec0ad6857132cf1c568b0","observation_id":"08e3b22d-4321-4978-8ba6-b884493639da","resolution":{"observed_at":"2026-08-11T23:27:33.575241Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2601.08584","last_updated":"2026-01-13T14:06:03Z","snapshot_observed_at":"2026-08-12T22:11:58.919874Z","submitted_at":"2026-01-13T14:06:03Z","title":"Ministral 3","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.08584","snapshot_observed_at":"2026-08-11T23:27:32.655467Z","title":"Liu, Kartik Khandelwal, Sandeep Subramanian, Victor Jouault, Abhinav Rastogi, et al","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.655467Z"},"links":{"cited_paper":"/paper/2601.08584","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:ea4eba8e78741ef0b61072896ac41c46bf44f55a6e273064738486c5dbe1a480","observation_id":"d275f545-eae2-41de-a217-de6461c842e9","resolution":{"observed_at":"2026-08-11T23:27:32.655467Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2302.02676","last_updated":"2023-10-18T07:11:12Z","snapshot_observed_at":"2026-08-13T12:50:20.733997Z","submitted_at":"2023-02-06T10:28:16Z","title":"Chain of Hindsight Aligns Language Models with Feedback","version":8},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.02676","snapshot_observed_at":"2026-08-11T23:27:32.658523Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.658523Z"},"links":{"cited_paper":"/paper/2302.02676","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:3380e0f5c443c59207474a7458aa74ddb39f24e5b43cd85fe131691a124dbfb5","observation_id":"20999843-4ea5-49a2-a9bb-a9e225dff3ff","resolution":{"observed_at":"2026-08-11T23:27:32.658523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.661712Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.661712Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:260bf9c780939a823aa34a3848422620cdc685414bba79e7774c318ffbb701ae","observation_id":"beb62dd7-0c5f-49a6-99f8-9d28f8a150d9","resolution":{"observed_at":"2026-08-11T23:27:32.661712Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.664633Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.664633Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:5056f5551fb920387f040285921be0209d17be426897d6b5f050bc7d9ee5ff39","observation_id":"cf0b829f-8c19-4c74-a1ac-610978094f78","resolution":{"observed_at":"2026-08-11T23:27:32.664633Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.667584Z","title":"McClelland, Bruce L","venue":null,"work_id":null,"year":1995},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.667584Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:c025c49e151010123511cbcb1eede2bef6f8a72b8f97438cd7e52f17cfde7ce8","observation_id":"d82ab8b9-1390-4d32-90d9-6fc0156394d1","resolution":{"observed_at":"2026-08-11T23:27:32.667584Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.562662Z","title":null,"venue":null,"work_id":"5585ab52-ed0d-4cbd-8b82-a267375c0fee","year":2022},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.670588Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:957015d2b7a0f9ee39eec93b2325fc66a2ada6ed3d89f72ce67b0b1173cfd4d1","observation_id":"792e1bdd-27c4-4b2b-bad8-0f50417b1dbb","resolution":{"observed_at":"2026-08-11T23:27:33.565778Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.673400Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.673400Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:c9e9d480d4a799b74204cff2ee47765fd89866e0cf66186db0cb92d87a2d6453","observation_id":"9483b4e2-e596-4b83-9c38-57ddb06fed9f","resolution":{"observed_at":"2026-08-11T23:27:32.673400Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.676269Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.676269Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:9f8dd5c295748bc81498c1bda681daed11c207644100b0722477ac9953c2ec50","observation_id":"fbdf20d6-8112-4a7b-91a0-fbf670c61672","resolution":{"observed_at":"2026-08-11T23:27:32.676269Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.547343Z","title":null,"venue":null,"work_id":"90165e89-76e5-417b-8ed2-0a3fbceb8e43","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.679064Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:701f890fbd0551b923afc02f0e727e47ce9e886a42133e39a9d2555f8963171e","observation_id":"561e00e5-2a54-4e99-b2fc-accaecab29e2","resolution":{"observed_at":"2026-08-11T23:27:33.550933Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.18290","last_updated":"2024-07-29T22:26:36Z","snapshot_observed_at":"2026-08-01T16:34:38.795326Z","submitted_at":"2023-05-29T17:57:46Z","title":"Direct Preference Optimization: Your Language Model is Secretly a Reward Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.18290","snapshot_observed_at":"2026-08-11T23:27:32.681796Z","title":"Manning, and Chelsea Finn","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.681796Z"},"links":{"cited_paper":"/paper/2305.18290","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:c737101a1d4fa06dc68960283dfa580c955f12558cb9c443fb333e648a069376","observation_id":"376331fb-3097-4a80-add6-5f2936b4c265","resolution":{"observed_at":"2026-08-11T23:27:32.681796Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.684861Z","title":"Robertson and Hugo Zaragoza","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.684861Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:85440f57a919495cd755d7c3a2e0b91adb90c9e0298c706ae2c698f29d52713a","observation_id":"6526d28b-bb1b-4420-80b3-f2b57548fbaf","resolution":{"observed_at":"2026-08-11T23:27:32.684861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2026.acl-long.1701","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.538081Z","title":null,"venue":null,"work_id":"ebe08e58-4173-4b30-b204-c4f95916cff5","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.688088Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:6618094bb0d1a5a469aea7d909c1ee88fa993c8f91548e5a4be6d9413606c9e3","observation_id":"2a503d5c-ec84-4285-bae6-e2445bf70d43","resolution":{"observed_at":"2026-08-11T23:27:33.541137Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.528863Z","title":null,"venue":null,"work_id":"a43e2b2b-1664-4cb8-9388-3186c08f7afd","year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.691135Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:0e1ea0daad7be7ed53a0167bcd2a55fe74dfd7806d48b30f634eff21445c4d3e","observation_id":"77eea6f1-0ce0-4e04-950d-194db880b5e5","resolution":{"observed_at":"2026-08-11T23:27:33.531952Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.05727","last_updated":"2025-08-04T02:14:13Z","snapshot_observed_at":"2026-08-10T21:05:12.257951Z","submitted_at":"2025-01-10T05:51:52Z","title":"Self-Evolving Critique Abilities in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.05727","snapshot_observed_at":"2026-08-11T23:27:32.694101Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.694101Z"},"links":{"cited_paper":"/paper/2501.05727","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:a79b4aff2da449adaeca1090870b00e88eaa2c0820e9647f0eeacb631495e9ac","observation_id":"7822ba08-53ed-4e57-9333-65c5b676de04","resolution":{"observed_at":"2026-08-11T23:27:32.694101Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.519417Z","title":null,"venue":null,"work_id":"e9dea157-d3f1-4abb-9667-b9a0935c8cad","year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.697254Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:1d327e3983f8f23d13a2c84a976f9c3ca5de6d1e17fa21a3ed742414dd4b20f7","observation_id":"cf9efc98-6fed-4169-9f6c-1a8855bb1b6f","resolution":{"observed_at":"2026-08-11T23:27:33.522787Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.510590Z","title":null,"venue":null,"work_id":"ece07ab8-b873-4f59-be58-18c45d4b6cb0","year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.705215Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:4d6e2af0f75a78de6698cc9fa25641dbfc176182b1109b25ecc82f8b19d9770d","observation_id":"d90a5be7-9302-4bda-9959-a2b3fcc63802","resolution":{"observed_at":"2026-08-11T23:27:33.513581Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.20857","last_updated":"2026-05-18T16:18:02Z","snapshot_observed_at":"2026-08-11T09:44:01.811561Z","submitted_at":"2025-11-25T21:08:07Z","title":"Evo-Memory: Benchmarking LLM Agent Test-time Learning with Self-Evolving Memory","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.20857","snapshot_observed_at":"2026-08-11T23:27:32.708285Z","title":"Chi, Chi Wang, Shuo Chen, Fernando Pereira, Wang-Cheng Kang, and Derek Zhiyuan Cheng","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.708285Z"},"links":{"cited_paper":"/paper/2511.20857","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:b502d9db00562d212551d736637bec679aecd0e1682509182e4c20c6b95247de","observation_id":"415fb5b4-650a-4e1b-bfd7-4e921de5505c","resolution":{"observed_at":"2026-08-11T23:27:32.708285Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.711965Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.711965Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:e2f06e4ce5ff60064f3bbb421eff6a04b09de7f4f5ad3462c744d8fec846213b","observation_id":"a0a85de5-5111-4429-bd86-64ebbaf5692a","resolution":{"observed_at":"2026-08-11T23:27:32.711965Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.findings-acl.818","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.880135Z","title":null,"venue":null,"work_id":"4edd77da-3c37-499c-a9b6-b4ca2d80408f","year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.718549Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:3b98fc4489c5f6cb91d381c8f95474bf54b996e7abef86212f185e370aa93790","observation_id":"f4c62592-924e-4c26-b660-556558500d25","resolution":{"observed_at":"2026-08-11T23:27:32.884316Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.12110","last_updated":"2025-10-08T01:46:37Z","snapshot_observed_at":"2026-08-03T02:27:06.991396Z","submitted_at":"2025-02-17T18:36:14Z","title":"A-MEM: Agentic Memory for LLM Agents","version":11},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.12110","snapshot_observed_at":"2026-08-11T23:27:32.715114Z","title":"arXiv:2502.12110 [cs.CL] https://arxiv.org/abs/2502.12110","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.715114Z"},"links":{"cited_paper":"/paper/2502.12110","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:c6d7a23350a145fcbb9e017548910230b85978ab8236be92d6fd6d1b9fadf736","observation_id":"fc7a41c0-480e-4f1d-87d6-03607d2196b1","resolution":{"observed_at":"2026-08-11T23:27:32.715114Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2026.findings-acl.22","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.849535Z","title":null,"venue":null,"work_id":"a0e42ba4-5f01-4346-8db9-ebedbf30b6b4","year":2026},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.725029Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:1ff26936514d65fa1183babd6ea69924a5b698cc04136ebfede681264f1c53e9","observation_id":"a388f0ef-b094-4aee-acf6-25a7a29ca21e","resolution":{"observed_at":"2026-08-11T23:27:32.855611Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-08-11T23:27:32.721644Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.721644Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:ddca075e328733cf83a9c1f88a2a332312b9f5abe0ded200881748da514ca4e5","observation_id":"4ebb84fc-8be7-4135-8d28-d3f571206e98","resolution":{"observed_at":"2026-08-11T23:27:32.721644Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.01470","last_updated":"2024-05-02T17:00:02Z","snapshot_observed_at":"2026-08-13T13:43:58.816757Z","submitted_at":"2024-05-02T17:00:02Z","title":"WildChat: 1M ChatGPT Interaction Logs in the Wild","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.01470","snapshot_observed_at":"2026-08-11T23:27:32.740590Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.740590Z"},"links":{"cited_paper":"/paper/2405.01470","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:b144b0563853728d623c20222f316e3abda81f04103d38338c67c3e8bbeb3733","observation_id":"23d260cc-f4c9-4466-acdf-142244cb54d3","resolution":{"observed_at":"2026-08-11T23:27:32.740590Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:32.728105Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.728105Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:ad033c769cd131822e67d4e38219ddf554c9ec947e5bec40df52ec2e600fbdf0","observation_id":"a2726180-dc30-412d-9334-cdfd54036779","resolution":{"observed_at":"2026-08-11T23:27:32.728105Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.05176","last_updated":"2025-06-11T02:54:49Z","snapshot_observed_at":"2026-08-13T07:28:32.447439Z","submitted_at":"2025-06-05T15:49:48Z","title":"Qwen3 Embedding: Advancing Text Embedding and Reranking Through Foundation Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.05176","snapshot_observed_at":"2026-08-11T23:27:32.737128Z","title":"doi:10.48550/arXiv","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.737128Z"},"links":{"cited_paper":"/paper/2506.05176","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:a554f9bbff4c16e82ca73635daace2299f3380d9bcc20d9a2db9d764bcb1093f","observation_id":"47d7215b-6d51-4bfb-b525-ee48dc605ca6","resolution":{"observed_at":"2026-08-11T23:27:32.737128Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.11998","last_updated":"2024-03-10T19:34:57Z","snapshot_observed_at":"2026-08-13T10:08:28.290453Z","submitted_at":"2023-09-21T12:13:55Z","title":"LMSYS-Chat-1M: A Large-Scale Real-World LLM Conversation Dataset","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.11998","snapshot_observed_at":"2026-08-11T23:27:32.744082Z","title":"Xing, Joseph E","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.744082Z"},"links":{"cited_paper":"/paper/2309.11998","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:2377607f7f23c8d1922624a21fecd26f104b1647ef0acec11bd02ca034766e75","observation_id":"dc7d04a6-0cd6-473b-bfce-59818bc5591e","resolution":{"observed_at":"2026-08-11T23:27:32.744082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.07911","last_updated":"2023-11-14T05:13:55Z","snapshot_observed_at":"2026-07-06T16:47:08.877195Z","submitted_at":"2023-11-14T05:13:55Z","title":"Instruction-Following Evaluation for Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.07911","snapshot_observed_at":"2026-08-11T23:27:32.754578Z","title":"components","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.754578Z"},"links":{"cited_paper":"/paper/2311.07911","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:563348b33d2b72f8b55bfadecbd38fa414111a51946b696643463018f0ea3c1b","observation_id":"c47baf38-2a0e-4547-ac64-d94e32b777bd","resolution":{"observed_at":"2026-08-11T23:27:32.754578Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.485802Z","title":"This includes ambiguity, irrelevance, conflict with TASK, pure reaction without a requested change, and a separate deliverable","venue":null,"work_id":"d3201be2-6ed9-4c36-aa55-b067c8b076a7","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.766082Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:4a179415f78fce574f418e1de44a7819820a0a639574e95fac03aeb13449bc7d","observation_id":"e89f3539-2c63-46e0-8a9b-57d2bd83bb41","resolution":{"observed_at":"2026-08-11T23:27:33.491012Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.474755Z","title":null,"venue":null,"work_id":"8fbb450d-fd9f-4f0c-8cec-117c7270652f","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.772120Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:1826eb72fb646c77b1e9f43b9f99991115a4b090c43212dc7707c5eecb0b19fa","observation_id":"78226881-846d-4006-a881-93d01b07ee8f","resolution":{"observed_at":"2026-08-11T23:27:33.479635Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.458480Z","title":"role\": \"FIX","venue":null,"work_id":"4c0027d8-4a77-44c6-9c46-07ad7ea18690","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.775505Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:8a15b57430fe440646792d2f68a84b415ff859ecf97dd2d160f126384007af28","observation_id":"93bc17d9-e6b9-4cab-a2ba-55ca31263524","resolution":{"observed_at":"2026-08-11T23:27:33.462154Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.448203Z","title":"action\":","venue":null,"work_id":"4d6c446c-0200-4a8b-8c50-6cc137dff782","year":2018},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.778468Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:4eaa9473f768c21149b41bb911dde5c0f1d1f64bd82ca91d86b91a792faa503d","observation_id":"d5a520ba-7b9a-46e5-aa57-86bd69e418b3","resolution":{"observed_at":"2026-08-11T23:27:33.452113Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.437663Z","title":null,"venue":null,"work_id":"924c6014-d010-4574-b66a-72b21a30a6b3","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.782294Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:a0eeb247f5cb6588b6f269e28726f34a8817ae688f721f5cbde0112bd986c7bd","observation_id":"f8d15c49-df33-4187-92c7-0206922ec122","resolution":{"observed_at":"2026-08-11T23:27:33.441027Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.427764Z","title":"I’m following up on the dashboard I sent for your use and would appreciate your feedback on its clarity, usefulness, and any areas that could be improved","venue":null,"work_id":"81404e79-895d-4cc4-a642-7a7e7857f716","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.785187Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:e2402c1b92eaaa20ab68337d70c6c1ec6b29db9dbc937b26d078e59c3c986d4a","observation_id":"f5705954-ade0-4a35-9309-f2d607e95d4f","resolution":{"observed_at":"2026-08-11T23:27:33.431401Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.398137Z","title":"action\":","venue":null,"work_id":"a128af09-93f5-4eb8-b631-007ebec0f95c","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.788305Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:05852c7118e7c38cd4e379caa2ac3e9a9d94ab3b7c352b74ca508ecf24caf76e","observation_id":"1a3b1f5d-66b4-4109-be39-1f2bca14106f","resolution":{"observed_at":"2026-08-11T23:27:33.417328Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.385229Z","title":"Thank you for your guidance","venue":null,"work_id":"ea091ef5-71d7-491c-9a63-e2f31a5bc0a9","year":null},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.791422Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:eb449a0db253481606a6b3c3300fce9152da99bc7c64717e35dd24c623688580","observation_id":"6aa74e32-9a8e-4e0f-bb3e-f4d4b87e70de","resolution":{"observed_at":"2026-08-11T23:27:33.388717Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T23:27:33.357582Z","title":"Target- consistent","venue":null,"work_id":"1058b4bf-b6d5-481b-8994-1feadcc92adf","year":2018},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.794315Z"},"links":{"citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:c98479846edfc990ac807b9750f06252e3b1dbe7cebf582e5bd957b57d4e239b","observation_id":"c4cb33b6-c72e-483b-8b7b-fe5e7589b7e4","resolution":{"observed_at":"2026-08-11T23:27:33.360911Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.03724","last_updated":"2025-12-03T03:19:27Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-04T17:21:46Z","title":"MemOS: A Memory OS for AI System","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.03724","snapshot_observed_at":"2026-08-11T23:27:32.652089Z","title":"doi:10.48550/arXiv.2507.03724","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-11T23:27:32.652089Z"},"links":{"cited_paper":"/paper/2507.03724","citing_paper":"/paper/2608.09109"},"observation_digest":"sha256:b8e1003d957f028ba902a9d922317187e766020b2f34d159a291f180f51a47cf","observation_id":"5b328f84-959e-4e16-af4c-1bd5bcf88fb3","resolution":{"observed_at":"2026-08-11T23:27:32.652089Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2608.09109","last_updated":"2026-08-10T04:27:58Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-08-13T23:41:27.460718Z","submitted_at":"2026-08-10T04:27:58Z","title":"Different Feedback, Different Updates: Selective Self-Learning from User Interactions for Large Language Models"},"reference_resolution":{"displayed":50,"state_counts":{"malformed_identifier":2,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":33,"verified_exact":8,"verified_fuzzy":7},"total_outbound_references":50},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 14 August 2026, this Paper Citation Record lists 50 of 50 outbound references and 0 inbound Pith citation observations for arXiv:2608.09109."}