{"as_of":"2026-08-10T04:39:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:51b54e981a263a24f52c0c62a8294f863a8b5e4812a6f86f4fed1bff99a03d95","coverage":[{"denominator":26,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":26,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T19:25:15.516775Z","state":"measured"},{"denominator":59,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":59,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":33,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":33,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T21:40:39.936409Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T10:59:46.413359Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2509.02544","last_updated":"2025-09-05T14:59:27Z","snapshot_observed_at":"2026-08-05T09:41:26.544360Z","submitted_at":"2025-09-02T17:44:45Z","title":"UI-TARS-2 Technical Report: Advancing GUI Agent with Multi-Turn Reinforcement Learning","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-05-13T10:13:58.774968Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2509.02544"},"observation_digest":"sha256:d647422963e302ffe5a727c3e8df421a0f0e57487ba5c25f450d19f4ea00d268","observation_id":"63f560cc-b7a7-414c-9742-3d60d8d1ffd6","resolution":{"observed_at":"2026-05-13T10:13:59.172725Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-08-04T21:28:49.949465Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment.arXiv preprint arXiv:2507.05720,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2509.07980","last_updated":"2025-09-12T17:15:56Z","snapshot_observed_at":"2026-08-07T14:28:58.369966Z","submitted_at":"2025-09-09T17:59:35Z","title":"Parallel-R1: Towards Parallel Thinking via Reinforcement Learning","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-04T21:28:49.949465Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2509.07980"},"observation_digest":"sha256:10cebd8b51e2610256118cddc5e231aebe12141cece1bd992fafd97082dfa653","observation_id":"3a8e54aa-775a-42e5-89c2-6ac17b513180","resolution":{"observed_at":"2026-08-04T21:28:49.949465Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-08-04T17:40:08.471718Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.10723","last_updated":"2025-09-12T22:26:31Z","snapshot_observed_at":"2026-08-04T17:40:07.040846Z","submitted_at":"2025-09-12T22:26:31Z","title":"Dark Patterns Meet GUI Agents: LLM Agent Susceptibility to Manipulative Interfaces and the Role of Human Oversight","version":1},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-04T17:40:08.471718Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2509.10723"},"observation_digest":"sha256:857f286459ab722aa8cf578149e6330a339b0c6eef9ea0eadee9afdb61e53b1c","observation_id":"5ae22897-4d5e-41b7-a0f6-50398cbe2574","resolution":{"observed_at":"2026-08-04T17:40:08.471718Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2509.21982","last_updated":"2026-04-13T03:14:30Z","snapshot_observed_at":"2026-07-06T22:30:50.642382Z","submitted_at":"2025-09-26T07:05:01Z","title":"RISK: A Framework for GUI Agents in E-commerce Risk Management","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-18T13:28:21.961995Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2509.21982"},"observation_digest":"sha256:0829e4baeec62475b5f22aa72deb4aad32f231a37b3f2b103df06d00753a86ef","observation_id":"0275ec03-9afc-4f46-92b8-d3a7355f8473","resolution":{"observed_at":"2026-05-18T13:31:25.053135Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-15T12:42:28.672749Z","title":"arXiv preprint arXiv:2507.05720 (2025) 2, 9","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.08316","last_updated":"2026-07-01T07:09:02Z","snapshot_observed_at":"2026-08-04T04:57:22.852231Z","submitted_at":"2026-03-09T12:38:28Z","title":"SlowBA: An efficiency backdoor attack towards VLM-based GUI agents","version":3},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-07-15T12:42:28.672749Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2603.08316"},"observation_digest":"sha256:904a18bc12f45305b9d1cc615bec1522f86689e826cf1b4d9fe213a958019a13","observation_id":"efb470d7-759a-483f-8b25-e9a0731b127a","resolution":{"observed_at":"2026-07-15T12:42:28.672749Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2604.22558","last_updated":"2026-04-24T13:53:39Z","snapshot_observed_at":"2026-08-05T08:11:55.563542Z","submitted_at":"2026-04-24T13:53:39Z","title":"SOLAR-RL: Semi-Online Long-horizon Assignment Reinforcement Learning","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-08T12:09:24.371878Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2604.22558"},"observation_digest":"sha256:6f7da6c51b4591e93878eb86b3158fdaee9d31564f889870d9a78c1e11031c16","observation_id":"1a3ebfeb-a48c-4ad7-9c76-b2aaf3e30458","resolution":{"observed_at":"2026-05-11T19:21:08.806694Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2604.24348","last_updated":"2026-04-27T11:44:26Z","snapshot_observed_at":"2026-07-31T20:25:37.534799Z","submitted_at":"2026-04-27T11:44:26Z","title":"OS-SPEAR: A Toolkit for the Safety, Performance,Efficiency, and Robustness Analysis of OS Agents","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-08T03:51:54.310805Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2604.24348"},"observation_digest":"sha256:2752a804ace1a8c7f184afa18ad36b6e4a79f63fcd9a4ebde54b91f0226dfba3","observation_id":"fbfbe4dd-dae8-4ec8-9efa-c94f4318c101","resolution":{"observed_at":"2026-05-11T21:56:11.716277Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2604.27859","last_updated":"2026-05-15T06:25:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-30T13:43:25Z","title":"Rethinking Agentic Reinforcement Learning In Large Language Models","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-05-07T06:30:09.945371Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2604.27859"},"observation_digest":"sha256:246649883184d986752ab8a8e5d2216854b1377ff1e401a1ce6cae335b11145b","observation_id":"a6998ae5-9bde-4c9d-a8e8-e95c95c4faf7","resolution":{"observed_at":"2026-05-12T10:21:29.723892Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2604.27859","last_updated":"2026-05-15T06:25:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-30T13:43:25Z","title":"Rethinking Agentic Reinforcement Learning In Large Language Models","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-05-08T03:12:19.414358Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2604.27859"},"observation_digest":"sha256:893fa51c3c853f5e5046f6750499affc829ee756a040577fccfa3caeb7435d18","observation_id":"0a30f749-90ca-4ef4-84fa-e6a68e7ddc80","resolution":{"observed_at":"2026-05-11T22:11:17.105373Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2604.27859","last_updated":"2026-05-15T06:25:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-30T13:43:25Z","title":"Rethinking Agentic Reinforcement Learning In Large Language Models","version":3},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-05-19T16:58:41.558250Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2604.27859"},"observation_digest":"sha256:7c6391c1ce45eae8923e01df659a0d5e8d9c612e2acb224ed8139a5b29913bd1","observation_id":"7c510af3-1617-4057-b248-c2c0fa308fff","resolution":{"observed_at":"2026-05-19T17:02:40.767243Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2604.27955","last_updated":"2026-04-30T14:51:49Z","snapshot_observed_at":"2026-07-06T23:13:20.516759Z","submitted_at":"2026-04-30T14:51:49Z","title":"GUI Agents with Reinforcement Learning: Toward Digital Inhabitants","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-07T05:48:00.486572Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2604.27955"},"observation_digest":"sha256:ad2e6cab1f3e3fd0655644b2ff06f480c8a8cb55c5a04c271d2e5c55932e3ec3","observation_id":"2863d9bf-e91d-4564-bd58-260ce13cf7cd","resolution":{"observed_at":"2026-05-12T10:31:29.543329Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2605.10347","last_updated":"2026-05-22T05:43:30Z","snapshot_observed_at":"2026-07-06T23:22:18.704329Z","submitted_at":"2026-05-11T10:49:31Z","title":"How Mobile World Model Guides GUI Agents?","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-12T04:28:34.562344Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2605.10347"},"observation_digest":"sha256:687d8d04fe381cf069d1736cdcfa09731b97fb23dea3a8cd0ca091443b3d1995","observation_id":"21044062-974a-4928-a17c-e81906e752b0","resolution":{"observed_at":"2026-05-12T06:16:24.120408Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2605.10347","last_updated":"2026-05-22T05:43:30Z","snapshot_observed_at":"2026-07-06T23:22:18.704329Z","submitted_at":"2026-05-11T10:49:31Z","title":"How Mobile World Model Guides GUI Agents?","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-25T06:00:34.560405Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2605.10347"},"observation_digest":"sha256:0e532bf5379b1b91da4709cdce7c51dbbe843a5160f52451fe1508e5ab97cd86","observation_id":"87ae547c-8f90-40d8-a953-6911361e1925","resolution":{"observed_at":"2026-05-25T06:05:27.174882Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2605.14311","last_updated":"2026-05-15T06:09:21Z","snapshot_observed_at":"2026-07-06T23:25:44.508166Z","submitted_at":"2026-05-14T03:23:44Z","title":"Beyond Binary: Reframing GUI Critique as Continuous Semantic Alignment","version":1},"reference_index":107,"source":"arxiv_source","source_observed_at":"2026-05-15T01:58:19.295247Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2605.14311"},"observation_digest":"sha256:bf16d5e2b5364dc8481d85e7a0a094850355d86ded4861e580b4cf34ec28184b","observation_id":"15c3040c-152b-4a25-97bf-2794e2f92f60","resolution":{"observed_at":"2026-05-15T01:58:28.798857Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2605.14311","last_updated":"2026-05-15T06:09:21Z","snapshot_observed_at":"2026-07-06T23:25:44.508166Z","submitted_at":"2026-05-14T03:23:44Z","title":"Beyond Binary: Reframing GUI Critique as Continuous Semantic Alignment","version":2},"reference_index":107,"source":"arxiv_source","source_observed_at":"2026-05-19T16:45:16.963802Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2605.14311"},"observation_digest":"sha256:e0c4178204c26ea47ef188a8d1344e65cbd9c9cacc611171d31563104cecce80","observation_id":"4bd35668-b932-429c-9454-db734f1fe880","resolution":{"observed_at":"2026-05-19T16:47:40.433083Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2605.15963","last_updated":"2026-05-15T13:55:05Z","snapshot_observed_at":"2026-07-06T23:27:11.118592Z","submitted_at":"2026-05-15T13:55:05Z","title":"PAGER: Bridging the Semantic-Execution Gap in Point-Precise Geometric GUI Control","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-20T18:46:20.417039Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2605.15963"},"observation_digest":"sha256:d4b2c040242dcae98a118d85e295e6827c5444221379e900c9354d85597816cb","observation_id":"26d1e9dc-e6af-4611-b782-43e1d92e7e53","resolution":{"observed_at":"2026-05-20T18:48:53.309473Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2605.20246","last_updated":"2026-05-21T05:35:18Z","snapshot_observed_at":"2026-08-02T08:14:48.016576Z","submitted_at":"2026-05-18T04:56:59Z","title":"GROW: Aligning GRPO with State-Action Modeling for Open-World VLM Agents","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-21T08:45:10.914047Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2605.20246"},"observation_digest":"sha256:23185f6226c61e3a281cb3d801390e4baa53f482d02e7406d5b5ca77867c68ea","observation_id":"bd8e26d0-56d1-46bb-9b71-b322cf9db82e","resolution":{"observed_at":"2026-05-21T08:49:53.857367Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2605.20246","last_updated":"2026-05-21T05:35:18Z","snapshot_observed_at":"2026-08-02T08:14:48.016576Z","submitted_at":"2026-05-18T04:56:59Z","title":"GROW: Aligning GRPO with State-Action Modeling for Open-World VLM Agents","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-22T09:06:03.221498Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2605.20246"},"observation_digest":"sha256:bd879a1b85fdb180443d6f5013f5d5f52aca6fe410124835e55bd106f0b81a5c","observation_id":"d7c09be8-2028-4ab9-b67d-f3049fe923ed","resolution":{"observed_at":"2026-05-22T09:06:20.173865Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2605.29486","last_updated":"2026-05-28T07:14:15Z","snapshot_observed_at":"2026-08-07T22:27:39.670483Z","submitted_at":"2026-05-28T07:14:15Z","title":"PhoneWorld: Scaling Phone-Use Agent Environments","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-29T07:41:20.001845Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2605.29486"},"observation_digest":"sha256:55259b87ab00243cb8a2c2b2b43719e213ef7e45d28cd6ed8d4dea6f1a4aaaa2","observation_id":"f42979d3-b9e7-40a2-8235-e7dfb4472da6","resolution":{"observed_at":"2026-06-29T07:43:13.465174Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.01249","last_updated":"2026-06-17T04:44:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-31T14:04:51Z","title":"Trust Region On-Policy Distillation","version":3},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-06-28T17:38:50.313305Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.01249"},"observation_digest":"sha256:64518b354f00b9c9cb546ab41280cdd674d56e95aa2748d517cb7b62a43cd85f","observation_id":"76350695-f64d-4dbd-bbed-099933d38905","resolution":{"observed_at":"2026-07-01T20:56:13.468062Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.07027","last_updated":"2026-06-05T08:17:28Z","snapshot_observed_at":"2026-08-01T08:01:00.499508Z","submitted_at":"2026-06-05T08:17:28Z","title":"StainFlow: Entity-Stain Tracking and Evidence Linking for Process Rewards in GUI Agents","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-27T22:24:41.353303Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.07027"},"observation_digest":"sha256:5baaa1d1f98380afc8fbd198912be6df414c59acdcfa9ab4fe95f0da66dca630","observation_id":"e2aac60f-1170-4956-b3de-3f4baa7949bb","resolution":{"observed_at":"2026-07-02T16:47:09.682368Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.13316","last_updated":"2026-07-31T08:57:20Z","snapshot_observed_at":"2026-08-05T23:10:47.877661Z","submitted_at":"2026-06-11T13:10:48Z","title":"ReSum: Synergizing LLM Reasoning and Summarization with Reinforcement Learning","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-06-27T06:39:34.199607Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.13316"},"observation_digest":"sha256:63a899d8090f7b441a2acb2b204113dbe4d8c8074d97e92ad4fd39611283a112","observation_id":"708df9eb-6a98-416f-b902-80cfd915f4b9","resolution":{"observed_at":"2026-07-03T15:08:33.447615Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-08-03T02:12:22.599076Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.13316","last_updated":"2026-07-31T08:57:20Z","snapshot_observed_at":"2026-08-05T23:10:47.877661Z","submitted_at":"2026-06-11T13:10:48Z","title":"ReSum: Synergizing LLM Reasoning and Summarization with Reinforcement Learning","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-03T02:12:22.599076Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.13316"},"observation_digest":"sha256:83e6f3795416b5d907d3a2d6212ab5535e4e7c6da0fe2e11bf5abfe5ce0402e5","observation_id":"523ffd92-8128-4ad0-bb08-a339d02f10f2","resolution":{"observed_at":"2026-08-03T02:12:22.599076Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.19930","last_updated":"2026-08-07T03:54:59Z","snapshot_observed_at":"2026-08-10T04:09:40.716236Z","submitted_at":"2026-06-18T08:29:33Z","title":"MobileForge: Annotation-Free Adaptation for Mobile GUI Agents with Hierarchical Feedback-Guided Policy Optimization","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-26T15:58:59.248141Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.19930"},"observation_digest":"sha256:56033f3d4dbf1249e952677fd2539d47a290d9f3932e6503f796dbcd874221b3","observation_id":"57f0ba21-c697-463f-bc2f-5afc9becc450","resolution":{"observed_at":"2026-07-04T05:29:35.726985Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.23049","last_updated":"2026-06-24T02:20:35Z","snapshot_observed_at":"2026-07-06T23:57:53.101659Z","submitted_at":"2026-06-22T08:57:54Z","title":"PhoneBuddy: Training Open Models for Agentic Phone Use","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-06-26T08:15:49.428124Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.23049"},"observation_digest":"sha256:9609c2e353b96076ddbe9144729ff86334bb781ff8c04d1ffc79b704629497bf","observation_id":"e40903f4-d89d-4bc7-96be-26b568f4e670","resolution":{"observed_at":"2026-07-04T10:59:46.414991Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.29705","last_updated":"2026-06-29T02:16:21Z","snapshot_observed_at":"2026-08-01T21:20:53.464369Z","submitted_at":"2026-06-29T02:16:21Z","title":"GUICrafter: Weakly-Supervised GUI Agent Leveraging Massive Unannotated Screenshots","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-06-30T06:39:12.591090Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.29705"},"observation_digest":"sha256:2ef7d3fbbd3cdd51dbddd669fd191401df4ffd181265bd428c017283994aa31c","observation_id":"261967c5-1828-4791-b59d-5bb38452d8ae","resolution":{"observed_at":"2026-06-30T06:44:19.182992Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.31410","last_updated":"2026-07-01T03:59:21Z","snapshot_observed_at":"2026-07-07T00:05:07.912609Z","submitted_at":"2026-06-30T09:36:35Z","title":"Xiaomi-GUI-0 Technical Report","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-07-01T06:07:48.137351Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.31410"},"observation_digest":"sha256:fdee77c2708613d002a3cc7d1657b644f5787fc22174751146f992d4c548f9c5","observation_id":"62da6060-afcd-4fd2-9fb9-8a4a31ff7234","resolution":{"observed_at":"2026-07-01T09:55:40.273011Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.31410","last_updated":"2026-07-01T03:59:21Z","snapshot_observed_at":"2026-07-07T00:05:07.912609Z","submitted_at":"2026-06-30T09:36:35Z","title":"Xiaomi-GUI-0 Technical Report","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-07-02T19:37:49.661596Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.31410"},"observation_digest":"sha256:05ab75fcbb6977cef91fd9033c527c3e6a23c0c092af75283662f65313da138f","observation_id":"ec939f23-337e-41e1-83e4-ab99c1f7fb6a","resolution":{"observed_at":"2026-07-02T19:47:19.031664Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.31612","last_updated":"2026-07-02T06:42:39Z","snapshot_observed_at":"2026-08-02T11:25:34.061644Z","submitted_at":"2026-06-30T13:01:19Z","title":"What Memory Do GUI Agents Really Need? From Passive Records to Active Task-Driving States","version":1},"reference_index":110,"source":"arxiv_source","source_observed_at":"2026-07-01T05:26:53.441833Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.31612"},"observation_digest":"sha256:48b034d94e9c9946a9f092f8350ee78e4b667bf4d6da05f395c04bdc519f6158","observation_id":"602e9c74-b509-42a8-96db-10e91dd21869","resolution":{"observed_at":"2026-07-01T10:25:42.439626Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":"2507.05720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-04T10:59:46.413359Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment","venue":null,"work_id":"4b1aecd7-add8-45d1-b46c-5c1a9c30fe03","year":2025},"citing_paper":{"arxiv_id":"2606.31612","last_updated":"2026-07-02T06:42:39Z","snapshot_observed_at":"2026-08-02T11:25:34.061644Z","submitted_at":"2026-06-30T13:01:19Z","title":"What Memory Do GUI Agents Really Need? From Passive Records to Active Task-Driving States","version":2},"reference_index":110,"source":"arxiv_source","source_observed_at":"2026-07-03T22:03:26.625209Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2606.31612"},"observation_digest":"sha256:6aaa4c3fe20bd6dbd6a35cfc5b429a43f118312312ab4db71dee1dde9a0ae0f9","observation_id":"53026ef0-790f-482a-b9ef-8fa5ff3e5ed3","resolution":{"observed_at":"2026-07-03T22:08:58.675212Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment.arXiv preprint arXiv:2507.05720,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:db819596c5c1f751f0928223b07dd03f1eb38ceabcd04011a9956b3edd1456db","observation_id":"f9893b93-b9d0-42f5-8f8a-aa23860c5100","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-31T23:04:32.783554Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment.arXiv preprint arXiv:2507.05720,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24112","last_updated":"2026-07-27T07:54:08Z","snapshot_observed_at":"2026-08-08T19:47:37.701014Z","submitted_at":"2026-07-27T07:54:08Z","title":"Scaling GUI Agents with Visual State Transitions","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-31T23:04:32.783554Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2607.24112"},"observation_digest":"sha256:1ab6ce3f2821622d838e6d7681afbe3f08dda9bbaf97584e84fad9faa1b20382","observation_id":"c7f73eeb-98a9-4b3f-a109-656ac294941f","resolution":{"observed_at":"2026-07-31T23:04:32.783554Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-08-07T21:40:39.936409Z","title":"Sun, H.; Xu, W.; Liu, W.; Luan, J.; Wang, B.; Shang, S.; Wen, J.-R.; and Yan, R","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.05891","last_updated":"2026-08-06T11:15:41Z","snapshot_observed_at":"2026-08-10T01:15:02.149227Z","submitted_at":"2026-08-06T11:15:41Z","title":"AppDeltaWorld: Transition-Grounded Delta Code World Model for Mobile GUI Agents","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T21:40:39.936409Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2608.05891"},"observation_digest":"sha256:348454e6c7d1923f5afa86d65fce3d857c97aab8c6d3609e65c17542e104b3bc","observation_id":"14cf9d17-7cd4-4318-86cc-9c162a8ecd3b","resolution":{"observed_at":"2026-08-07T21:40:39.936409Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2507.05720/citation-record","integrity":"/paper/2507.05720/integrity","json":"/paper/2507.05720/citation-record.json","paper":"/paper/2507.05720"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:25:16.161564Z","title":"Anthropic","venue":null,"work_id":"04475044-f8c7-4418-b24e-355be4943b5c","year":2025},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:13.327952Z"},"links":{"citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:2afea4f922b6f8537357cd3a30ab30067bda580aa1e869918522a8266d9aa3a4","observation_id":"c6dcb8f5-1007-49e2-8bb5-515fce356d95","resolution":{"observed_at":"2026-08-06T19:25:16.291161Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-08-06T19:25:13.431735Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:13.431735Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:8e844704a191bf0f35171f200fb6bf8fe8e0a0f583dcbc3d84323bdabfc4e4be","observation_id":"576aa0c9-9596-4b60-a2e1-ec8566a2860c","resolution":{"observed_at":"2026-08-06T19:25:13.431735Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.10978","last_updated":"2025-10-28T15:11:36Z","snapshot_observed_at":"2026-07-29T19:20:21.974239Z","submitted_at":"2025-05-16T08:26:59Z","title":"Group-in-Group Policy Optimization for LLM Agent Training","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.10978","snapshot_observed_at":"2026-08-06T19:25:13.581491Z","title":"Group-in-group policy optimization for llm agent training","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:13.581491Z"},"links":{"cited_paper":"/paper/2505.10978","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:b072d4a49c618c258b81e14ed3c04db84049556176e964359a1c4f0c3a2f876f","observation_id":"ca850bb0-9c2f-4c4d-950c-0094caa3da06","resolution":{"observed_at":"2026-08-06T19:25:13.581491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-08-06T19:25:13.669050Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:13.669050Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:f7a6c8d2878d73c1b91f9d64ca5ada800fa569fbbe751254f89742b4ff48c4d4","observation_id":"f46be06e-2ed8-4fa3-963e-6116923fbf40","resolution":{"observed_at":"2026-08-06T19:25:13.669050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.12856","last_updated":"2024-02-25T16:17:43Z","snapshot_observed_at":"2026-07-06T15:57:53.178833Z","submitted_at":"2023-07-24T14:56:30Z","title":"A Real-World WebAgent with Planning, Long Context Understanding, and Program Synthesis","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.12856","snapshot_observed_at":"2026-08-06T19:25:13.752798Z","title":"A real-world webagent with planning, long context understanding, and program synthesis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:13.752798Z"},"links":{"cited_paper":"/paper/2307.12856","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:8617b37cfbb4b42b4dbe32a7a39de092d43d8d3bc50567b04cb580679fb9114c","observation_id":"ea2e2c8e-5f4a-45fb-8318-28831894267c","resolution":{"observed_at":"2026-08-06T19:25:13.752798Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:25:13.829334Z","title":"Webcot: Enhancing web agent reasoning by recon- structing chain-of-thought in reflection, branching, and rollback","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:13.829334Z"},"links":{"citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:42f46ca3cb4c71bf43bd65833c89738850bbbc1d0093c97fab9247b6a69204de","observation_id":"1972ed27-6b21-4592-8db5-70eda511947a","resolution":{"observed_at":"2026-08-06T19:25:13.829334Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.02716","last_updated":"2024-02-05T04:25:24Z","snapshot_observed_at":"2026-08-09T14:47:51.558187Z","submitted_at":"2024-02-05T04:25:24Z","title":"Understanding the planning of LLM agents: A survey","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.02716","snapshot_observed_at":"2026-08-06T19:25:13.947753Z","title":"Understanding the planning of llm agents: A survey","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:13.947753Z"},"links":{"cited_paper":"/paper/2402.02716","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:b00436d31740f2e9dc6af6d93159425b193ff1f5b75db929e433ab2475da1842","observation_id":"969098ea-d8b0-4762-a8f5-1729a0d8092e","resolution":{"observed_at":"2026-08-06T19:25:13.947753Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-08-06T19:25:14.053149Z","title":"Gpt-4o system card","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.053149Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:01ee92de785766bab4c48d0119e5578d86542c752a1485b056310618bf1b7b24","observation_id":"7ff5e137-d65a-4927-9de1-7b55855348f0","resolution":{"observed_at":"2026-08-06T19:25:14.053149Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.14239","last_updated":"2025-04-19T09:25:55Z","snapshot_observed_at":"2026-08-05T23:52:52.919433Z","submitted_at":"2025-04-19T09:25:55Z","title":"InfiGUI-R1: Advancing Multimodal GUI Agents from Reactive Actors to Deliberative Reasoners","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.14239","snapshot_observed_at":"2026-08-06T19:25:14.164095Z","title":"Infigui-r1: Advancing multimodal gui agents from reactive actors to deliberative reasoners","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.164095Z"},"links":{"cited_paper":"/paper/2504.14239","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:5925fed95526cee4f4fa37695a88fde8aa1da0f813d7e5ce238292cd7f6a0d08","observation_id":"34ae667f-c7b6-4cf1-a984-7e9718c6209c","resolution":{"observed_at":"2026-08-06T19:25:14.164095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.16282","last_updated":"2025-05-22T06:24:32Z","snapshot_observed_at":"2026-08-09T18:08:58.661516Z","submitted_at":"2025-05-22T06:24:32Z","title":"ARPO:End-to-End Policy Optimization for GUI Agents with Experience Replay","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.16282","snapshot_observed_at":"2026-08-06T19:25:14.276099Z","title":"Arpo: End-to-end policy opti- mization for gui agents with experience replay","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.276099Z"},"links":{"cited_paper":"/paper/2505.16282","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:dfe5aa7b16a4eddc6ef350fb41cc455e7604a6b380672620471f33416432c662","observation_id":"c7df2476-8c0a-48d9-8a9f-2cf60b09acd7","resolution":{"observed_at":"2026-08-06T19:25:14.276099Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12326","last_updated":"2025-01-21T17:48:10Z","snapshot_observed_at":"2026-07-06T20:23:58.426780Z","submitted_at":"2025-01-21T17:48:10Z","title":"UI-TARS: Pioneering Automated GUI Interaction with Native Agents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12326","snapshot_observed_at":"2026-08-06T19:25:14.363133Z","title":"Ui-tars: Pioneering automated gui interaction with native agents","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.363133Z"},"links":{"cited_paper":"/paper/2501.12326","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:ac74f9615f17de8e5b4246941e8e586f86cabeac5e97eaa7a1ece26d6206ce90","observation_id":"f4db16d2-c08d-4299-a4b6-645f8904d456","resolution":{"observed_at":"2026-08-06T19:25:14.363133Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14573","last_updated":"2025-04-06T20:37:50Z","snapshot_observed_at":"2026-08-07T08:11:20.950548Z","submitted_at":"2024-05-23T13:48:54Z","title":"AndroidWorld: A Dynamic Benchmarking Environment for Autonomous Agents","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.14573","snapshot_observed_at":"2026-08-06T19:25:14.413197Z","title":"Androidworld: A dynamic benchmarking environment for autonomous agents","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.413197Z"},"links":{"cited_paper":"/paper/2405.14573","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:419de088d4a597547b897c193eaf2208eb961b986a30bac6b216c27b6693d393","observation_id":"27403f17-307e-4b33-920a-ae00166372af","resolution":{"observed_at":"2026-08-06T19:25:14.413197Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06347","last_updated":"2017-08-28T09:20:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-07-20T02:32:33Z","title":"Proximal Policy Optimization Algorithms","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1707.06347","snapshot_observed_at":"2026-08-06T19:25:14.500949Z","title":"Proximal policy optimization algorithms","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.500949Z"},"links":{"cited_paper":"/paper/1707.06347","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:7164465e39f87448b9e42348b01db393c8b04ee879d34f8dad15c1983a53780d","observation_id":"e414f609-f21c-4afd-a036-6395d8f8a743","resolution":{"observed_at":"2026-08-06T19:25:14.500949Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.19723","last_updated":"2025-06-27T08:25:48Z","snapshot_observed_at":"2026-08-07T22:32:37.258577Z","submitted_at":"2024-12-27T16:21:58Z","title":"OS-Genesis: Automating GUI Agent Trajectory Construction via Reverse Task Synthesis","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.19723","snapshot_observed_at":"2026-08-06T19:25:14.659875Z","title":"Os-genesis: Automating gui agent trajectory construction via reverse task synthesis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.659875Z"},"links":{"cited_paper":"/paper/2412.19723","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:1a8dd421f9b73809173be7258b90ec630e9e8a18259ee2faaeb6825cc807122a","observation_id":"c9d985b2-e4d4-4630-a71a-c02dc279676f","resolution":{"observed_at":"2026-08-06T19:25:14.659875Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.04890","last_updated":"2025-02-13T10:09:37Z","snapshot_observed_at":"2026-08-06T08:17:46.881583Z","submitted_at":"2024-11-07T17:28:10Z","title":"GUI Agents with Foundation Models: A Comprehensive Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.04890","snapshot_observed_at":"2026-08-06T19:25:14.763290Z","title":"Gui agents with foundation models: A comprehensive survey","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.763290Z"},"links":{"cited_paper":"/paper/2411.04890","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:593f2ae7dae39ccedff7824bdcbaf73e9e33cd81ce8cff83bc5d642631c9d9f1","observation_id":"07fae943-36b9-42f7-bd29-663d8f9a71aa","resolution":{"observed_at":"2026-08-06T19:25:14.763290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:25:14.815524Z","title":"net/forum?id=LPG8pPSfQD","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.815524Z"},"links":{"citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:935874b1d67c17be04970502fe45f3c4bc4e8c91b6cb9bd9fe497a5a4e1efbb4","observation_id":"c87c052a-2cc4-4289-b899-1ada62990ee1","resolution":{"observed_at":"2026-08-06T19:25:14.815524Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.23218","last_updated":"2024-10-30T17:10:19Z","snapshot_observed_at":"2026-08-05T06:30:40.974506Z","submitted_at":"2024-10-30T17:10:19Z","title":"OS-ATLAS: A Foundation Action Model for Generalist GUI Agents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.23218","snapshot_observed_at":"2026-08-06T19:25:14.898449Z","title":"Os-atlas: A foundation action model for generalist gui agents","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.898449Z"},"links":{"cited_paper":"/paper/2410.23218","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:abc315b38690d973f05230c938e42619b030fd2ef1a9820f532f5e2d3e04129b","observation_id":"7c4bb803-937f-4e66-a8e7-e1ed7125e31d","resolution":{"observed_at":"2026-08-06T19:25:14.898449Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.23762","last_updated":"2025-05-29T17:59:51Z","snapshot_observed_at":"2026-08-08T14:02:10.449991Z","submitted_at":"2025-05-29T17:59:51Z","title":"ZeroGUI: Automating Online GUI Learning at Zero Human Cost","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.23762","snapshot_observed_at":"2026-08-06T19:25:14.987098Z","title":"Zerogui: Automating online gui learning at zero human cost","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.987098Z"},"links":{"cited_paper":"/paper/2505.23762","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:fe5f0e26e22d02badf50954957cb132c0bd805522bae091d552e338be3bb9f36","observation_id":"f8971c1c-42bb-4ec5-be13-e442f3de724c","resolution":{"observed_at":"2026-08-06T19:25:14.987098Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.14476","last_updated":"2025-05-20T01:37:34Z","snapshot_observed_at":"2026-08-02T01:40:54.187278Z","submitted_at":"2025-03-18T17:49:06Z","title":"DAPO: An Open-Source LLM Reinforcement Learning System at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.14476","snapshot_observed_at":"2026-08-06T19:25:15.059491Z","title":"Dapo: An open-source llm reinforcement learning system at scale","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:15.059491Z"},"links":{"cited_paper":"/paper/2503.14476","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:8a1bdda4d58e30cc1ce2514a41497d88d4cc58c7367eee0ea68733b21983f42d","observation_id":"9b759d12-6ff1-4189-8a1c-7a4a11f82a21","resolution":{"observed_at":"2026-08-06T19:25:15.059491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.18279","last_updated":"2025-05-06T15:08:00Z","snapshot_observed_at":"2026-07-06T19:57:55.925634Z","submitted_at":"2024-11-27T12:13:39Z","title":"Large Language Model-Brained GUI Agents: A Survey","version":12},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.18279","snapshot_observed_at":"2026-08-06T19:25:15.167234Z","title":"Large language model-brained gui agents: A survey","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:15.167234Z"},"links":{"cited_paper":"/paper/2411.18279","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:2b06f662947fe2bf64e5b5e2127f1f21f28b10bdb2f76aa33e9c40d56ba0f939","observation_id":"99de16da-ae23-4fe6-b540-ba5d18d94c12","resolution":{"observed_at":"2026-08-06T19:25:15.167234Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.01614","last_updated":"2024-03-12T23:14:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-03T08:33:09Z","title":"GPT-4V(ision) is a Generalist Web Agent, if Grounded","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.01614","snapshot_observed_at":"2026-08-06T19:25:15.259521Z","title":"Gpt-4v (ision) is a generalist web agent, if grounded","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:15.259521Z"},"links":{"cited_paper":"/paper/2401.01614","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:69cc987b5de18a615ebafc1dce0ca824f2ef880f7b48a867e97caf53fbf93231","observation_id":"3baabd57-cca2-4b7d-8a47-8eaaf5a6d485","resolution":{"observed_at":"2026-08-06T19:25:15.259521Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.15810","last_updated":"2025-05-22T11:15:23Z","snapshot_observed_at":"2026-08-09T07:15:17.774887Z","submitted_at":"2025-05-21T17:59:09Z","title":"GUI-G1: Understanding R1-Zero-Like Training for Visual Grounding in GUI Agents","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.15810","snapshot_observed_at":"2026-08-06T19:25:15.375077Z","title":"Gui-g1: Understanding r1-zero-like training for visual grounding in gui agents","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:15.375077Z"},"links":{"cited_paper":"/paper/2505.15810","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:7c2802cae8cd4136ef9d5884c730314d97e5db0102f25212119026098fb10f17","observation_id":"69bdcee3-0121-4858-af68-24f41777ac12","resolution":{"observed_at":"2026-08-06T19:25:15.375077Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:25:16.054042Z","title":"Create a calendar event for tomorrow at 20h with the title ’Call with the Team’ and the description ’We will prepare for team roles.’. The event should last for 30 mins","venue":null,"work_id":"e0e72d2c-f408-41cc-951c-c5348e1a205b","year":2024},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:15.516775Z"},"links":{"citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:656d61ee30dee6426691c0d5607f2f3df0ec13284be979d2be2a4e92267a8712","observation_id":"67bf531f-2d18-4b56-8d1a-aa57fe5fcc82","resolution":{"observed_at":"2026-08-06T19:25:16.121157Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-06T19:25:14.593433Z","title":"Deepseekmath: Pushing the limits of mathematical reasoning in open language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:14.593433Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:f7cb418b30120b8230938aff0e11f90c901419bda49efe01ffc07ca6ffe7232e","observation_id":"690966ed-c226-4dbc-834f-5658b79c412c","resolution":{"observed_at":"2026-08-06T19:25:14.593433Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:25:16.427847Z","title":"Anthropic","venue":null,"work_id":"94558af1-479d-4590-8618-8ffd3af9a6aa","year":2025},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:13.262977Z"},"links":{"citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:1ab8f0231aa0f6ac194d1f4af271cc0b39bc33327eb523791e530cee0cf876eb","observation_id":"a1dfcd8c-5dc4-45fe-b9f4-d98706c0ab84","resolution":{"observed_at":"2026-08-06T19:25:16.564804Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.21024","last_updated":"2025-08-21T05:55:43Z","snapshot_observed_at":"2026-08-09T21:07:02.669130Z","submitted_at":"2025-04-23T02:54:31Z","title":"WebEvolver: Enhancing Web Agent Self-Improvement with Coevolving World Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.21024","snapshot_observed_at":"2026-08-06T19:25:13.505455Z","title":"Webevolver: Enhancing web agent self-improvement with coevolving world model","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-06T19:25:13.505455Z"},"links":{"cited_paper":"/paper/2504.21024","citing_paper":"/paper/2507.05720"},"observation_digest":"sha256:93d0ef40bde1b44142266b5b5a775949c6952986fd965f9ecde2fe4c7d99e85b","observation_id":"abc42732-20d2-48ca-8db9-f658d6e8d679","resolution":{"observed_at":"2026-08-06T19:25:13.505455Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment"},"reference_resolution":{"displayed":26,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":23,"verified_exact":0,"verified_fuzzy":3},"total_outbound_references":26},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 26 of 26 outbound references and 33 inbound Pith citation observations for arXiv:2507.05720."}