{"as_of":"2026-08-13T18:03:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:01fdc54fcfa40e7499965b05202f09570e871550af3c77708093257824884881","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":47,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":47,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":47,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":47,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T20:54:45.673130Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T09:47:59.545083Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-12T20:54:45.673130Z","title":"arXiv preprint arXiv:2406.10100","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.09301","last_updated":"2024-11-14T09:23:40Z","snapshot_observed_at":"2026-08-13T06:30:46.098333Z","submitted_at":"2024-11-14T09:23:40Z","title":"LHRS-Bot-Nova: Improved Multimodal Large Language Model for Remote Sensing Vision-Language Interpretation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T20:54:45.673130Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2411.09301"},"observation_digest":"sha256:2693f9a2d97ef6cea936679326593342e8a6cff65a745c2b9dfd968ee3e1848e","observation_id":"e4bea15d-f33a-47cd-8be0-ccecfb522928","resolution":{"observed_at":"2026-08-12T20:54:45.673130Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-12T10:21:22.567192Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.19325","last_updated":"2025-03-12T19:28:05Z","snapshot_observed_at":"2026-08-12T10:15:33.525314Z","submitted_at":"2024-11-28T18:59:56Z","title":"GEOBench-VLM: Benchmarking Vision-Language Models for Geospatial Tasks","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-12T10:21:22.567192Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2411.19325"},"observation_digest":"sha256:202890352522924d1c99137a8a95a4e1779ce9028bbb0f12d7e0afc3074ff175","observation_id":"d9ea57bc-fb00-41f8-ac97-5cc1b089929e","resolution":{"observed_at":"2026-08-12T10:21:22.567192Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-11T20:32:31.655715Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.05679","last_updated":"2024-12-10T02:23:30Z","snapshot_observed_at":"2026-08-12T05:19:12.753205Z","submitted_at":"2024-12-07T15:11:21Z","title":"RSUniVLM: A Unified Vision Language Model for Remote Sensing via Granularity-oriented Mixture of Experts","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-11T20:32:31.655715Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2412.05679"},"observation_digest":"sha256:5468cd86d3c355c1c831946a9b8c688a21dc2230fa02a7a96825ec74218909bc","observation_id":"bbd86159-5cad-4f12-b6ca-539211b748b8","resolution":{"observed_at":"2026-08-11T20:32:31.655715Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-11T11:36:35.321568Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.15190","last_updated":"2025-04-07T06:19:02Z","snapshot_observed_at":"2026-08-12T17:01:58.965430Z","submitted_at":"2024-12-19T18:57:13Z","title":"EarthDial: Turning Multi-sensory Earth Observations to Interactive Dialogues","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T11:36:35.321568Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2412.15190"},"observation_digest":"sha256:c08f77763890d391e9cd0f3be7cbe1c5f71666fd2e3ab68c9e693b035bff427e","observation_id":"26bfe0da-a9b9-42a9-92fd-197a29d63697","resolution":{"observed_at":"2026-08-11T11:36:35.321568Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-11T10:30:49.625539Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.16583","last_updated":"2024-12-21T11:17:15Z","snapshot_observed_at":"2026-08-13T06:30:47.102049Z","submitted_at":"2024-12-21T11:17:15Z","title":"REO-VLM: Transforming VLM to Meet Regression Challenges in Earth Observation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T10:30:49.625539Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2412.16583"},"observation_digest":"sha256:2ed13a45f60f8df16f15d668ef9d43d638dde98b8933990354603e28fd10e741","observation_id":"04f4fd80-60f9-4520-a6cf-78a289fd86a7","resolution":{"observed_at":"2026-08-11T10:30:49.625539Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-10T23:18:21.355874Z","title":"Skysensegpt: A fine-grained instruction tuning dataset and model for remote sensing vision-language understanding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.20742","last_updated":"2024-12-30T06:34:18Z","snapshot_observed_at":"2026-08-13T13:55:29.561149Z","submitted_at":"2024-12-30T06:34:18Z","title":"UniRS: Unifying Multi-temporal Remote Sensing Tasks through Vision Language Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-10T23:18:21.355874Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2412.20742"},"observation_digest":"sha256:8af11767e43d33f9d83e39cc2a2d07bee2056e69bdfb1531cfeaebd1b20b5c9a","observation_id":"f3867446-ec3d-4bc5-b5ed-6bd479515b0e","resolution":{"observed_at":"2026-08-10T23:18:21.355874Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-10T20:53:48.229191Z","title":"Skysensegpt: A fine- grained instruction tuning dataset and model for remote sensing vision-language understanding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.06828","last_updated":"2025-03-13T08:16:01Z","snapshot_observed_at":"2026-08-13T16:56:56.343853Z","submitted_at":"2025-01-12T14:45:27Z","title":"GeoPix: Multi-Modal Large Language Model for Pixel-level Image Understanding in Remote Sensing","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T20:53:48.229191Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2501.06828"},"observation_digest":"sha256:29bdf74fa3e5398d0b9f8f32be021dff993ed1675551a58fe484abda733508bd","observation_id":"ed8fa80b-baed-4c4e-bb80-6eaef91a6e75","resolution":{"observed_at":"2026-08-10T20:53:48.229191Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-10T19:48:22.426502Z","title":"Skysensegpt: A fine-grained instruction tun- ing dataset and model for remote sensing vision-language understanding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09720","last_updated":"2025-01-31T21:29:40Z","snapshot_observed_at":"2026-08-13T17:28:25.975897Z","submitted_at":"2025-01-16T18:09:22Z","title":"A Simple Aerial Detection Baseline of Multimodal Language Models","version":3},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T19:48:22.426502Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2501.09720"},"observation_digest":"sha256:3a8c46ec6d13a7cda8e44bd617cfcf4c49b0a5c08f86cc9733fcc21d1bce1679","observation_id":"ba9150a4-093c-4345-9f98-20a6603184d5","resolution":{"observed_at":"2026-08-10T19:48:22.426502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-10T19:26:44.887132Z","title":"Skysensegpt: A fine-grained instruction tuning dataset and model for remote sensing vision-language understanding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.10144","last_updated":"2025-01-17T12:12:33Z","snapshot_observed_at":"2026-08-10T19:21:17.993926Z","submitted_at":"2025-01-17T12:12:33Z","title":"A Vision-Language Framework for Multispectral Scene Representation Using Language-Grounded Features","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T19:26:44.887132Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2501.10144"},"observation_digest":"sha256:0194c0374b19093518f5213808010e26e716f961d664bd96809f6e8821a2dddd","observation_id":"ca4e9245-c9bb-4459-b843-cc1052393104","resolution":{"observed_at":"2026-08-10T19:26:44.887132Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-07T15:38:44.401424Z","title":"Skysensegpt: A fine-grained instruction tuning dataset and model for remote sensing vision-language understanding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.14361","last_updated":"2025-05-20T13:47:40Z","snapshot_observed_at":"2026-08-10T20:27:13.425753Z","submitted_at":"2025-05-20T13:47:40Z","title":"Vision-Language Modeling Meets Remote Sensing: Models, Datasets and Perspectives","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T15:38:44.401424Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2505.14361"},"observation_digest":"sha256:ba11b11a14da50fe6924fa95e1ef9014a271abd9a94c774b33ec4cb3eb7e3994","observation_id":"9ae29f85-5d1f-4171-9a1e-bd4620c1b678","resolution":{"observed_at":"2026-08-07T15:38:44.401424Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-06T22:22:38.288515Z","title":"Skysensegpt: A fine-grained instruction tuning dataset and model for remote sensing vision-language understanding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.21863","last_updated":"2025-06-27T02:31:37Z","snapshot_observed_at":"2026-08-11T11:40:41.776672Z","submitted_at":"2025-06-27T02:31:37Z","title":"Remote Sensing Large Vision-Language Model: Semantic-augmented Multi-level Alignment and Semantic-aware Expert Modeling","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T22:22:38.288515Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2506.21863"},"observation_digest":"sha256:cdac8db4d8596fa11e737f5488079fc3667ebaaf772ae1328d8043a50397593a","observation_id":"f16d1259-155d-4449-8054-a9a99b154a5a","resolution":{"observed_at":"2026-08-06T22:22:38.288515Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-06T21:52:22.909176Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.23219","last_updated":"2025-06-29T13:04:27Z","snapshot_observed_at":"2026-08-11T00:58:38.306843Z","submitted_at":"2025-06-29T13:04:27Z","title":"UrbanLLaVA: A Multi-modal Large Language Model for Urban Intelligence with Spatial Reasoning and Understanding","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T21:52:22.909176Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2506.23219"},"observation_digest":"sha256:a2fdc16b80671dd7b10017407e5290775c3a84443f3e06216abc42cbe45f21ea","observation_id":"cf9c645e-b765-4ad7-85d1-1def9183b5d5","resolution":{"observed_at":"2026-08-06T21:52:22.909176Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-06T21:49:44.779870Z","title":"Skysensegpt: A fine-grained instruction tuning dataset and model for re- mote sensing vision-language understanding","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.23352","last_updated":"2025-06-29T18:03:03Z","snapshot_observed_at":"2026-08-09T05:37:38.871712Z","submitted_at":"2025-06-29T18:03:03Z","title":"GeoProg3D: Compositional Visual Reasoning for City-Scale 3D Language Fields","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T21:49:44.779870Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2506.23352"},"observation_digest":"sha256:475e9289ac55ee4e9b8482eb4823302d0867d9e83387f5c0bb34dbedb48a9d3f","observation_id":"42f8f1c8-c23b-4d8e-bec0-fae4bce9c14e","resolution":{"observed_at":"2026-08-06T21:49:44.779870Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-06T19:46:12.290250Z","title":"et al., 2024a","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.04664","last_updated":"2025-07-07T05:10:15Z","snapshot_observed_at":"2026-08-12T16:47:36.907901Z","submitted_at":"2025-07-07T05:10:15Z","title":"VectorLLM: Human-like Extraction of Structured Building Contours vis Multimodal LLMs","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-06T19:46:12.290250Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2507.04664"},"observation_digest":"sha256:3b2fd671f9d50e98fe0c91e921a2e8e67dbf28a9fed2e6492ff45376522ee9d1","observation_id":"b13b024e-9a19-4236-ac5f-2454e73df9d4","resolution":{"observed_at":"2026-08-06T19:46:12.290250Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2511.23332","last_updated":"2026-05-11T14:21:31Z","snapshot_observed_at":"2026-08-10T23:36:32.005563Z","submitted_at":"2025-11-28T16:40:08Z","title":"UniGeoSeg: Towards Unified Open-World Segmentation for Geospatial Scenes","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-17T04:37:12.944874Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2511.23332"},"observation_digest":"sha256:eec600199f71415fe4c1cf9ed6fa99f705d1fda9039a69c47896e8f4a642f4c6","observation_id":"8ffa5418-9196-4717-bc6e-b0d5641b41d8","resolution":{"observed_at":"2026-05-17T04:39:03.328784Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2512.17492","last_updated":"2026-04-28T06:26:52Z","snapshot_observed_at":"2026-08-13T01:52:24.307850Z","submitted_at":"2025-12-19T12:03:05Z","title":"MMLANDMARKS: a Cross-View Instance-Level Benchmark for Geo-Spatial Understanding","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-16T20:49:11.961079Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2512.17492"},"observation_digest":"sha256:d359f7d61e4b5cb0f4d49986b374025126e9d6496670f877e20c0c10f86962ae","observation_id":"8e25808f-3e2e-4e14-b1cd-d3587b79009a","resolution":{"observed_at":"2026-05-16T20:51:15.264448Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2601.01891","last_updated":"2026-01-05T08:34:17Z","snapshot_observed_at":"2026-08-12T14:20:38.723123Z","submitted_at":"2026-01-05T08:34:17Z","title":"Agentic AI in Remote Sensing: Foundations, Taxonomy, and Emerging Systems","version":1},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-05-16T18:28:33.277442Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2601.01891"},"observation_digest":"sha256:9945e626ea6f545446bd22818d3215f5c69d9440f3268e29c810fde2b67fec65","observation_id":"d6c47bf0-3361-4e32-8429-b492462d7ee9","resolution":{"observed_at":"2026-05-16T18:31:10.820324Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2602.07045","last_updated":"2026-05-14T08:15:47Z","snapshot_observed_at":"2026-08-12T13:46:50.835837Z","submitted_at":"2026-02-04T08:21:33Z","title":"VLRS-Bench: A Vision-Language Reasoning Benchmark for Remote Sensing","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-16T08:06:43.928029Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2602.07045"},"observation_digest":"sha256:cda0d87d3e2641d6c7bd72cb8fe11167d64f4a3927d9d975bb750c591b23c4f4","observation_id":"c7413130-d60f-4b5b-8638-cbe42b902252","resolution":{"observed_at":"2026-05-16T08:07:34.083932Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-02T22:10:51.236088Z","title":"arXiv preprint arXiv:2406.10100 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2602.17665","last_updated":"2026-07-12T17:48:41Z","snapshot_observed_at":"2026-08-09T06:01:13.885489Z","submitted_at":"2026-02-19T18:59:54Z","title":"OpenEarthAgent: A Unified Framework for Tool-Augmented Geospatial Agents","version":4},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-02T22:10:51.236088Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2602.17665"},"observation_digest":"sha256:ec71d51c80b9c9487aa0fcf08742650feed1f7ba7c9135f4ee774234969d8c65","observation_id":"c49e2ad0-fe82-4a21-b751-ff1ce7c7b4cb","resolution":{"observed_at":"2026-08-02T22:10:51.236088Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-04T05:53:31.189135Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.11804","last_updated":"2026-08-03T09:02:14Z","snapshot_observed_at":"2026-08-09T17:28:56.114205Z","submitted_at":"2026-03-12T11:08:30Z","title":"OSMDA: OpenStreetMap-based Domain Adaptation for Remote Sensing VLMs","version":3},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-04T05:53:31.189135Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2603.11804"},"observation_digest":"sha256:81f8e21ba6b09a6355adefc163d0b6da157d77a4df01ec5669049428f221936e","observation_id":"4300c672-2360-426e-8515-a8d1c6608488","resolution":{"observed_at":"2026-08-04T05:53:31.189135Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-15T11:59:49.461107Z","title":"arxiv 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.16307","last_updated":"2026-07-15T03:58:39Z","snapshot_observed_at":"2026-08-07T14:04:41.514571Z","submitted_at":"2026-03-17T09:43:00Z","title":"NeSy-Route: A Neuro-Symbolic Benchmark for Constrained Route Planning in Remote Sensing","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-15T11:59:49.461107Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2603.16307"},"observation_digest":"sha256:09d2dba2708c26c1b96a56bf3f8c5012e9c21e8a0fe925d9e21787416f05be15","observation_id":"450e51f5-f484-4cf7-b594-5856e21ca3fe","resolution":{"observed_at":"2026-07-15T11:59:49.461107Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-02T18:07:48.194772Z","title":"arxiv 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.16307","last_updated":"2026-07-15T03:58:39Z","snapshot_observed_at":"2026-08-07T14:04:41.514571Z","submitted_at":"2026-03-17T09:43:00Z","title":"NeSy-Route: A Neuro-Symbolic Benchmark for Constrained Route Planning in Remote Sensing","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-02T18:07:48.194772Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2603.16307"},"observation_digest":"sha256:abcc8e00ba8db930428e45a1818c63151b5674ec321ee747722607a922d863a0","observation_id":"6f99e2ba-13c9-414c-a598-e36c177bdcfa","resolution":{"observed_at":"2026-08-02T18:07:48.194772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2604.07765","last_updated":"2026-04-12T05:49:10Z","snapshot_observed_at":"2026-08-11T11:43:35.089308Z","submitted_at":"2026-04-09T03:40:46Z","title":"RemoteAgent: Bridging Vague Human Intents and Earth Observation with RL-based Agentic MLLMs","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-10T18:00:20.216268Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2604.07765"},"observation_digest":"sha256:48466be1aa7f892a1d31bd38dfcaaa3b25bf9ac1d43663cc8a17708ce1f1c875","observation_id":"b1f1243f-6e52-4753-9425-84c88c52c1f1","resolution":{"observed_at":"2026-05-11T05:40:59.510236Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2604.08896","last_updated":"2026-04-10T02:59:38Z","snapshot_observed_at":"2026-08-06T17:54:52.229005Z","submitted_at":"2026-04-10T02:59:38Z","title":"GeoMMBench and GeoMMAgent: Toward Expert-Level Multimodal Intelligence in Geoscience and Remote Sensing","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-10T17:56:12.097628Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2604.08896"},"observation_digest":"sha256:e4bf104a02326c6029964ccb219935a1d96a2b62eee8c9f9998031d7248b88c9","observation_id":"443405c3-476e-4674-a7a0-c1b1b6d3de44","resolution":{"observed_at":"2026-05-11T05:51:00.393334Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2604.10591","last_updated":"2026-04-12T11:47:11Z","snapshot_observed_at":"2026-08-12T07:28:11.685251Z","submitted_at":"2026-04-12T11:47:11Z","title":"GeoMeld: Toward Semantically Grounded Foundation Models for Remote Sensing","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T15:16:55.294813Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2604.10591"},"observation_digest":"sha256:469726bd434c11c0b8bb45ac48507a947fa85a341610ce5c0ab138ba3b874f08","observation_id":"5bdf4a5b-3a28-4d47-a502-6c783599d33d","resolution":{"observed_at":"2026-05-11T10:56:01.594799Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2604.13654","last_updated":"2026-04-15T09:20:02Z","snapshot_observed_at":"2026-07-06T23:01:37.207817Z","submitted_at":"2026-04-15T09:20:02Z","title":"Vision-and-Language Navigation for UAVs: Progress, Challenges, and a Research Roadmap","version":1},"reference_index":235,"source":"pdf_text","source_observed_at":"2026-05-10T13:48:08.135538Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2604.13654"},"observation_digest":"sha256:5f7dedfdd569905ce10c8677b4bf13400702335e2f7623080f9309b52abf43fe","observation_id":"b9501892-3512-436c-bff3-0e9aa784107b","resolution":{"observed_at":"2026-05-11T11:36:02.129268Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2604.17243","last_updated":"2026-04-19T04:04:44Z","snapshot_observed_at":"2026-07-06T23:04:23.730478Z","submitted_at":"2026-04-19T04:04:44Z","title":"RemoteShield: Enable Robust Multimodal Large Language Models for Earth Observation","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-10T06:48:11.250689Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2604.17243"},"observation_digest":"sha256:614fe2b68fac8c481afd97f8b63f6e3ce1263c74d83bbe76a1b2cc32ff1faff2","observation_id":"2d956cbd-59cb-4345-ad40-a36a2df62840","resolution":{"observed_at":"2026-05-10T06:51:46.253973Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2604.22855","last_updated":"2026-04-22T12:28:04Z","snapshot_observed_at":"2026-08-12T23:02:10.325214Z","submitted_at":"2026-04-22T12:28:04Z","title":"Evaluating Remote Sensing Image Captions Beyond Metric Biases","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-10T01:06:35.604862Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2604.22855"},"observation_digest":"sha256:e38bdff41c8f922f495a90931f3e52a9715cd4e2ef29d340169fce920de87dab","observation_id":"d7593d96-a0dd-47b4-bb8b-6d26fde39839","resolution":{"observed_at":"2026-05-10T01:10:09.411869Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2604.24919","last_updated":"2026-06-01T00:50:05Z","snapshot_observed_at":"2026-08-12T15:37:50.007543Z","submitted_at":"2026-04-27T18:59:49Z","title":"Agentic AI for Remote Sensing: Technical Challenges and Research Directions","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-05-08T04:29:22.477531Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2604.24919"},"observation_digest":"sha256:ff7d4e9d10a0a9829d4fc468a2e1e75b75630b74431f65f4c25cd29e77e41087","observation_id":"efb24cf2-d62d-453d-a9cf-ca80abde23d6","resolution":{"observed_at":"2026-05-11T21:46:29.767853Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2604.24919","last_updated":"2026-06-01T00:50:05Z","snapshot_observed_at":"2026-08-12T15:37:50.007543Z","submitted_at":"2026-04-27T18:59:49Z","title":"Agentic AI for Remote Sensing: Technical Challenges and Research Directions","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-05-14T20:55:38.841743Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2604.24919"},"observation_digest":"sha256:0cb3e1e049709630edf5981cf7afa849fc893ff3af01a9706f2ddb53a0a2d78e","observation_id":"3c3763a1-f7ee-4ae4-b93a-4ef327cb6dfe","resolution":{"observed_at":"2026-05-14T20:59:27.730514Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2605.04451","last_updated":"2026-08-07T02:21:45Z","snapshot_observed_at":"2026-08-13T14:56:42.822561Z","submitted_at":"2026-05-06T03:25:41Z","title":"RemoteZero: Geospatial Reasoning with Zero Labels","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-08T18:24:46.030608Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2605.04451"},"observation_digest":"sha256:d2753704e733ade6e56793ddd31c61ee3e6937ef7e3b982de9dcd091214883f5","observation_id":"64c4b8d6-ce6c-44ca-a0e8-363d3804e94b","resolution":{"observed_at":"2026-05-09T06:30:40.509035Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2605.04777","last_updated":"2026-05-06T11:30:21Z","snapshot_observed_at":"2026-08-10T21:11:42.125344Z","submitted_at":"2026-05-06T11:30:21Z","title":"Bridging Perception and Action: A Lightweight Multimodal Meta-Planner Framework for Robust Earth Observation Agents","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-05-08T15:45:27.700503Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2605.04777"},"observation_digest":"sha256:84b858c4556be5c963f19d023ec21a4ae6d9b9ae756a1693a1f0c5a56ce110c3","observation_id":"82ef7a8d-7944-41d3-a245-1b32360b0e4e","resolution":{"observed_at":"2026-05-11T18:31:12.207277Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2605.05594","last_updated":"2026-05-07T02:27:04Z","snapshot_observed_at":"2026-07-06T23:18:12.631820Z","submitted_at":"2026-05-07T02:27:04Z","title":"The Cost of Context: Mitigating Textual Bias in Multimodal Retrieval-Augmented Generation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-08T11:17:44.998703Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2605.05594"},"observation_digest":"sha256:9d74c98cd8651980b4045307a79b5f4b4c98eae17c7e84b931b7671ee726e846","observation_id":"c9c6ef9a-5bcd-4a12-9590-14b03a842ed3","resolution":{"observed_at":"2026-05-11T19:41:08.805651Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2605.07562","last_updated":"2026-05-08T10:35:11Z","snapshot_observed_at":"2026-07-06T23:19:56.415523Z","submitted_at":"2026-05-08T10:35:11Z","title":"Beyond GSD-as-Token: Continuous Scale Conditioning for Remote Sensing VLMs","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-11T01:52:51.777228Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2605.07562"},"observation_digest":"sha256:3e88a7425094b469f509ce581c7370f1b840690416733cdd1fdc385394ff9c05","observation_id":"d93d75bb-a9a4-45eb-bb15-73d515e680ed","resolution":{"observed_at":"2026-05-11T04:15:56.932162Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2605.12542","last_updated":"2026-06-11T05:51:05Z","snapshot_observed_at":"2026-08-11T15:09:57.550432Z","submitted_at":"2026-05-09T08:34:30Z","title":"Earth Science Foundation Models: From Perception to Reasoning and Discovery","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-14T22:07:40.242567Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2605.12542"},"observation_digest":"sha256:49faee9ed761a6f7c328a12c1154fb7286a850dce6c2cbe6b18adc038015f9a3","observation_id":"80f53e9d-461b-4795-9936-35392faf677d","resolution":{"observed_at":"2026-05-14T22:08:03.561292Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2605.12542","last_updated":"2026-06-11T05:51:05Z","snapshot_observed_at":"2026-08-11T15:09:57.550432Z","submitted_at":"2026-05-09T08:34:30Z","title":"Earth Science Foundation Models: From Perception to Reasoning and Discovery","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-06-30T23:07:21.558834Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2605.12542"},"observation_digest":"sha256:117a8fafb172142d40a893bd1329208d262cdb7322bb1fb0f84417399c1da53a","observation_id":"21b28714-d104-4a97-bf3f-6eeda9257b1b","resolution":{"observed_at":"2026-07-01T13:35:46.137194Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2606.10819","last_updated":"2026-06-09T13:01:51Z","snapshot_observed_at":"2026-08-11T19:39:55.065170Z","submitted_at":"2026-06-09T13:01:51Z","title":"Earth-OneVision: Extending Remote Sensing Multimodal Large Language Models to More Sensor Modalities and Tasks","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-27T13:16:12.768793Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2606.10819"},"observation_digest":"sha256:a57dbe34f1e33e14f75874e44b6da98706e789cf946a8c9ca993b989e834ade8","observation_id":"cb299905-9bb4-4fc5-a604-aad2fde7a52f","resolution":{"observed_at":"2026-07-03T05:27:40.010548Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2606.11740","last_updated":"2026-06-10T07:16:27Z","snapshot_observed_at":"2026-08-02T02:17:16.809620Z","submitted_at":"2026-06-10T07:16:27Z","title":"UniReason-Med: A Shared Grounded Reasoning Interface for 2D-to-3D Transfer in Medical VQA","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-06-27T10:21:12.782864Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2606.11740"},"observation_digest":"sha256:99a3e529a2d7a81eb376ab4876234969ac5254dbb5db8adbdcf68cdfb024f53e","observation_id":"3f8b395c-be29-46d6-ab56-cd058a08a398","resolution":{"observed_at":"2026-07-03T09:47:59.546380Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":"2406.10100","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-03T09:47:59.545083Z","title":"Skysensegpt: A fine-grained in- struction tuning dataset and model for remote sensing vision- language understanding.arXiv preprint arXiv:2406.10100","venue":null,"work_id":"7e943dbe-53ff-448e-9058-3b37224aadbd","year":2024},"citing_paper":{"arxiv_id":"2607.01050","last_updated":"2026-07-01T15:12:51Z","snapshot_observed_at":"2026-07-07T00:06:39.908636Z","submitted_at":"2026-07-01T15:12:51Z","title":"GeoSearcher: Anchor-Guided Progressive Reasoning for Remote Sensing Visual Grounding with Process Supervision","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-02T13:52:29.466227Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2607.01050"},"observation_digest":"sha256:17ff1f60599076548897a78e072f1b0f84670f22d5969193e6738fe0485d313c","observation_id":"f52bf4f1-e34c-43fa-9946-b403cf99ed20","resolution":{"observed_at":"2026-07-02T13:56:59.140862Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-07-11T22:49:02.844739Z","title":"SkySenseGPT: A fine-grained instruction tuning dataset and model for remote sensing vision-language understanding.arXiv preprint arXiv:2406.10100,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.03949","last_updated":"2026-08-06T06:36:23Z","snapshot_observed_at":"2026-08-09T23:09:50.971564Z","submitted_at":"2026-07-04T16:52:34Z","title":"TESSERA v2: Scaling Pixel-wise Earth Foundation Models","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-11T22:49:02.844739Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2607.03949"},"observation_digest":"sha256:2c97281390bbaafd0948c2c8370bc66457803f13715f486dc6b3b6ac66874efc","observation_id":"af3c0b0d-a4c6-4bc0-97a6-af42e3955630","resolution":{"observed_at":"2026-07-11T22:49:02.844739Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-01T22:27:21.208472Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.15768","last_updated":"2026-07-17T08:58:29Z","snapshot_observed_at":"2026-08-10T21:53:54.246008Z","submitted_at":"2026-07-17T08:58:29Z","title":"GeoChrono: Benchmarking and Rethinking Long-Term Temporal Understanding in Remote Sensing","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-01T22:27:21.208472Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2607.15768"},"observation_digest":"sha256:34c1fa1a95514e04a2e2962042d680147590529e9422a5c1d6ff2ca4a9899757","observation_id":"0edfb1e6-1e5f-4514-94b7-1f2a17e3322b","resolution":{"observed_at":"2026-08-01T22:27:21.208472Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-01T21:55:56.609929Z","title":"arXiv preprint arXiv:2406.10100 (2024).https://doi.org/10.48550/arXiv.2406.101002, 3, 4","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.15942","last_updated":"2026-07-17T13:25:44Z","snapshot_observed_at":"2026-08-08T05:33:02.171627Z","submitted_at":"2026-07-17T13:25:44Z","title":"More with Less: a Large Scale Remote Sensing VLM with a Simple Recipe","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-01T21:55:56.609929Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2607.15942"},"observation_digest":"sha256:392db47c7f175dec700c390e90f4dbd9e0f485a034b2983d65048ebe44db4de4","observation_id":"9499125f-1f8c-4a9b-9d54-7f74a8375d32","resolution":{"observed_at":"2026-08-01T21:55:56.609929Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-01T18:11:02.304478Z","title":"Skysensegpt: A fine-grained instruction tuning dataset and model for remote sensing vision-language understanding.arXiv preprint arXiv:2406.10100, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.17386","last_updated":"2026-07-19T19:07:19Z","snapshot_observed_at":"2026-08-09T17:28:55.542980Z","submitted_at":"2026-07-19T19:07:19Z","title":"SkyVLaM: Multimodal Large Language Model for UAV Video Understanding in Remote Sensing","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-01T18:11:02.304478Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2607.17386"},"observation_digest":"sha256:0432eb80a9ae76d4fdaea1c9689d7e553abf7c5c594ca6d26a1d9a0d36aa52b3","observation_id":"7da24c0e-c78c-49c0-915a-672bacdb4bf6","resolution":{"observed_at":"2026-08-01T18:11:02.304478Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-01T11:33:43.613930Z","title":"Skysensegpt: A fine-grained instruction tuning dataset and model for remote sensing vision-language understanding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.19857","last_updated":"2026-07-22T07:45:50Z","snapshot_observed_at":"2026-08-12T13:31:25.598171Z","submitted_at":"2026-07-22T07:45:50Z","title":"Memory-Augmented Multimodal Large Language Models for Small Object Understanding in Streaming Aerial Videos","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-01T11:33:43.613930Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2607.19857"},"observation_digest":"sha256:a64dc9dd44a84ddf43ed86d1df1cd78cd191b422a9634c04c18efbe95337c2eb","observation_id":"aa2add77-efb9-4878-be89-fdcd1b629294","resolution":{"observed_at":"2026-08-01T11:33:43.613930Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-01T10:20:57.231744Z","title":"arXiv preprint arXiv:2406.10100 , primaryclass =","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.20284","last_updated":"2026-07-22T15:30:08Z","snapshot_observed_at":"2026-08-09T18:33:19.700255Z","submitted_at":"2026-07-22T15:30:08Z","title":"Multimodal Large Language Models for Remote Sensing Image Understanding: Domain-Specific or General-Purpose?","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-01T10:20:57.231744Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2607.20284"},"observation_digest":"sha256:79e2f076d227104534daa1b4f8479d329f7a06c02975da76b549f0d4a35acc8a","observation_id":"9015ada6-7997-4178-b8f4-91ef0f9baea2","resolution":{"observed_at":"2026-08-01T10:20:57.231744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-01T05:32:43.036448Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.22205","last_updated":"2026-07-28T11:05:01Z","snapshot_observed_at":"2026-08-13T08:37:12.676504Z","submitted_at":"2026-07-24T11:21:25Z","title":"Filling Before Advancing: Capability-Gap-Driven Post-Training for Scenario-Specialized Remote Sensing MLLMs","version":2},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-08-01T05:32:43.036448Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2607.22205"},"observation_digest":"sha256:9658158b614c0ca82a0625f54f859e33ef14ef8b7ff75d5d71f2f15bfea68ff7","observation_id":"21ce53e2-4ee3-4fff-9c18-b689f17bf94d","resolution":{"observed_at":"2026-08-01T05:32:43.036448Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10100","snapshot_observed_at":"2026-08-01T00:59:20.911720Z","title":"Skysensegpt: A fine-grained instruction tuning dataset and model for remote sensing vision-language understanding.arXiv preprint arXiv:2406.10100, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.25993","last_updated":"2026-07-28T17:09:16Z","snapshot_observed_at":"2026-08-12T12:27:46.074162Z","submitted_at":"2026-07-28T17:09:16Z","title":"Beyond Zooming: Learning Multi-Tool Visual Reasoning for Ultra-High-Resolution Remote Sensing","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-01T00:59:20.911720Z"},"links":{"cited_paper":"/paper/2406.10100","citing_paper":"/paper/2607.25993"},"observation_digest":"sha256:8c81f06bf753e603f71ce264fc13f033021be81250f0d6d27a18467390f16bf1","observation_id":"efde5a10-6ec6-4f5a-a206-f4331c1d89bd","resolution":{"observed_at":"2026-08-01T00:59:20.911720Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2406.10100/citation-record","integrity":"/paper/2406.10100/integrity","json":"/paper/2406.10100/citation-record.json","paper":"/paper/2406.10100"},"outbound":[],"paper":{"arxiv_id":"2406.10100","last_updated":"2024-07-08T04:33:37Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-12T23:42:36.797128Z","submitted_at":"2024-06-14T14:57:07Z","title":"SkySenseGPT: A Fine-Grained Instruction Tuning Dataset and Model for Remote Sensing Vision-Language Understanding"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 47 inbound Pith citation observations for arXiv:2406.10100."}