{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2026:Q7E5PPRS25U77MVLH6SZXUHEJT","short_pith_number":"pith:Q7E5PPRS","schema_version":"1.0","canonical_sha256":"87c9d7be32d769ffb2ab3fa59bd0e44cf8c41f544734e0d6579133ec837d40bc","source":{"kind":"arxiv","id":"2606.27872","version":1},"attestation_state":"computed","paper":{"title":"S$^2$-VLA: State-Space Guided Vision-Language-Action Models for Long-Horizon Manipulation","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.RO","authors_text":"Jing Zhao, Shiliang Sun, Xiangyi Wei, Yang Li, Zhipeng Xie, Zongyi Han","submitted_at":"2026-06-26T09:13:16Z","abstract_excerpt":"Vision-Language-Action (VLA) models have demonstrated strong capabilities in robotic manipulation, but their performance degrades significantly in long-horizon tasks due to cumulative error propagation. This limitation largely arises from static feature fusion mechanisms that rely on fixed weights to combine visual, language, and action representations, preventing the model from adapting to different phases of task execution. To address this limitation, we propose S$^2$-VLA, a framework that introduces a State-Space Guided Adaptive Attention (SSGAA) mechanism. SSGAA maintains a belief state th"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2606.27872","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.RO","submitted_at":"2026-06-26T09:13:16Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"4d5f11fee75ebdd60e238e99e4c138c2805618edd19f5180b9d0b3047a13fd09","abstract_canon_sha256":"cea2251767319b05adb2ffa3b4ebf4a43bcc4dfc17a946affddf36719500a384"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-06-29T01:14:51.315140Z","signature_b64":"qg5z/RcV2EmqIdtz+BT23j0tqgBHMIQNLyLzneeh9MlwnsdqsDr4NXlozOukbARpy0rxftXrJ/oEnTpHYh+YCA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"87c9d7be32d769ffb2ab3fa59bd0e44cf8c41f544734e0d6579133ec837d40bc","last_reissued_at":"2026-06-29T01:14:51.314737Z","signature_status":"signed_v1","first_computed_at":"2026-06-29T01:14:51.314737Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"S$^2$-VLA: State-Space Guided Vision-Language-Action Models for Long-Horizon Manipulation","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.RO","authors_text":"Jing Zhao, Shiliang Sun, Xiangyi Wei, Yang Li, Zhipeng Xie, Zongyi Han","submitted_at":"2026-06-26T09:13:16Z","abstract_excerpt":"Vision-Language-Action (VLA) models have demonstrated strong capabilities in robotic manipulation, but their performance degrades significantly in long-horizon tasks due to cumulative error propagation. This limitation largely arises from static feature fusion mechanisms that rely on fixed weights to combine visual, language, and action representations, preventing the model from adapting to different phases of task execution. To address this limitation, we propose S$^2$-VLA, a framework that introduces a State-Space Guided Adaptive Attention (SSGAA) mechanism. SSGAA maintains a belief state th"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2606.27872","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2606.27872/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2606.27872","created_at":"2026-06-29T01:14:51.314797+00:00"},{"alias_kind":"arxiv_version","alias_value":"2606.27872v1","created_at":"2026-06-29T01:14:51.314797+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2606.27872","created_at":"2026-06-29T01:14:51.314797+00:00"},{"alias_kind":"pith_short_12","alias_value":"Q7E5PPRS25U7","created_at":"2026-06-29T01:14:51.314797+00:00"},{"alias_kind":"pith_short_16","alias_value":"Q7E5PPRS25U77MVL","created_at":"2026-06-29T01:14:51.314797+00:00"},{"alias_kind":"pith_short_8","alias_value":"Q7E5PPRS","created_at":"2026-06-29T01:14:51.314797+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/Q7E5PPRS25U77MVLH6SZXUHEJT","json":"https://pith.science/pith/Q7E5PPRS25U77MVLH6SZXUHEJT.json","graph_json":"https://pith.science/api/pith-number/Q7E5PPRS25U77MVLH6SZXUHEJT/graph.json","events_json":"https://pith.science/api/pith-number/Q7E5PPRS25U77MVLH6SZXUHEJT/events.json","paper":"https://pith.science/paper/Q7E5PPRS"},"agent_actions":{"view_html":"https://pith.science/pith/Q7E5PPRS25U77MVLH6SZXUHEJT","download_json":"https://pith.science/pith/Q7E5PPRS25U77MVLH6SZXUHEJT.json","view_paper":"https://pith.science/paper/Q7E5PPRS","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2606.27872&json=true","fetch_graph":"https://pith.science/api/pith-number/Q7E5PPRS25U77MVLH6SZXUHEJT/graph.json","fetch_events":"https://pith.science/api/pith-number/Q7E5PPRS25U77MVLH6SZXUHEJT/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/Q7E5PPRS25U77MVLH6SZXUHEJT/action/timestamp_anchor","attest_storage":"https://pith.science/pith/Q7E5PPRS25U77MVLH6SZXUHEJT/action/storage_attestation","attest_author":"https://pith.science/pith/Q7E5PPRS25U77MVLH6SZXUHEJT/action/author_attestation","sign_citation":"https://pith.science/pith/Q7E5PPRS25U77MVLH6SZXUHEJT/action/citation_signature","submit_replication":"https://pith.science/pith/Q7E5PPRS25U77MVLH6SZXUHEJT/action/replication_record"}},"created_at":"2026-06-29T01:14:51.314797+00:00","updated_at":"2026-06-29T01:14:51.314797+00:00"}