{"as_of":"2026-08-21T14:36:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b4539c4a625a6f2889fea75f1366a6676d6d730e3cbb8cd43d86df066ed872a6","coverage":[{"denominator":43,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":43,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:16:12.146074Z","state":"measured"},{"denominator":43,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":43,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.15703/citation-record","integrity":"/paper/2505.15703/integrity","json":"/paper/2505.15703/citation-record.json","paper":"/paper/2505.15703"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2301.00493","last_updated":"2023-01-02T00:36:22Z","snapshot_observed_at":"2026-08-12T21:55:36.747659Z","submitted_at":"2023-01-02T00:36:22Z","title":"Argoverse 2: Next Generation Datasets for Self-Driving Perception and Forecasting","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.00493","snapshot_observed_at":"2026-08-07T15:16:10.721562Z","title":"Argoverse 2: Next generation datasets for self-driving perception and forecasting,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.721562Z"},"links":{"cited_paper":"/paper/2301.00493","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:7627473106bf432845cc0775e528c93addc39c70fe3caa524a505a2b5fdbe277","observation_id":"eb7f8e35-be2f-4d6f-bfa7-96cf1ebedeae","resolution":{"observed_at":"2026-08-07T15:16:10.721562Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.636531Z","title":"A survey on trajectory-prediction methods for autonomous driving,","venue":null,"work_id":"87c2ae12-9b52-4ca9-82dd-31a35d31888c","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.773288Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:5b49827d3c78559ccd69652d23121dcc04f5e005125b5d23e9812b2ccf987287","observation_id":"06e5b162-8cbe-46e6-90d8-50add44f764d","resolution":{"observed_at":"2026-08-07T15:16:13.644669Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.611757Z","title":"Vectornet: Encoding hd maps and agent dynamics from vectorized representation,","venue":null,"work_id":"b7a52fae-2348-4064-8bed-3065ec1a1a71","year":2020},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.814015Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:52fb933625f453b301ad2b7b66c5ccdc1cd367524411c0d82b8877102f221b9c","observation_id":"2473a5d5-517a-4ec4-a8c6-8c522a1b46a8","resolution":{"observed_at":"2026-08-07T15:16:13.618951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.589855Z","title":"Learning lane graph representations for motion forecasting,","venue":null,"work_id":"c6c0434d-944e-400c-90f5-aeaaf5b1c762","year":2020},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.851157Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:adb7987a32d2578bc21fe522cee1edf13a68b3ab69d8ec9c465337b3ad78dd83","observation_id":"da0a02a5-3c41-4902-9ab3-d6dfcae877fb","resolution":{"observed_at":"2026-08-07T15:16:13.596594Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.568681Z","title":"Simpl: A simple and efficient multi-agent motion prediction baseline for autonomous driving,","venue":null,"work_id":"592561c8-82d6-4527-b2cd-5b1b4d19062b","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.873717Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:80e8a48eafcd352197908f6074bc6dbd0a94d63a18d2d7cac499aadfc93d49ce","observation_id":"23c89954-c98d-4559-b397-c7d0456c3e65","resolution":{"observed_at":"2026-08-07T15:16:13.576163Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2106.08417","last_updated":"2022-03-04T20:25:25Z","snapshot_observed_at":"2026-08-18T06:43:45.571895Z","submitted_at":"2021-06-15T20:20:44Z","title":"Scene Transformer: A unified architecture for predicting multiple agent trajectories","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.08417","snapshot_observed_at":"2026-08-07T15:16:10.889911Z","title":"Scene transformer: A unified architecture for predicting multiple agent tra- jectories,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.889911Z"},"links":{"cited_paper":"/paper/2106.08417","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:fa2320f30cbf7da0cc713ed37631d65a7593fb0d401f16d5335301d934524111","observation_id":"56fd1171-cf96-4879-8ffd-493d4f656378","resolution":{"observed_at":"2026-08-07T15:16:10.889911Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.548359Z","title":"Forecast-mae: Self-supervised pre- training for motion forecasting with masked autoencoders,","venue":null,"work_id":"90432211-560f-4bd8-81e4-e2c602d73116","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.917439Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:b43bf8f2583b4ef89942d0d18e5a4e3183ca994cedfce43228b9072468e887f4","observation_id":"9fcb46f8-a407-4cb3-9072-b0c84ae37c97","resolution":{"observed_at":"2026-08-07T15:16:13.555570Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.523172Z","title":"Gorela: Go relative for viewpoint-invariant motion forecasting,","venue":null,"work_id":"208d4d60-0ca2-41b2-af41-a61967f39159","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.937366Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:ebaba39ec65db02681c4399e052f46eb02680c1900d4b768a3684fc07843d1e2","observation_id":"c3ad91c7-2881-46ae-83dd-a7020e38b0ad","resolution":{"observed_at":"2026-08-07T15:16:13.531818Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.499582Z","title":"Multipath++: Efficient information fusion and trajectory aggregation for behavior prediction,","venue":null,"work_id":"f815bff8-fa34-46c5-b2b7-55ffd0b37b95","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.969322Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:7c3bf692ac48ac6421be9ab4f015291be38e265fe9e81d48c1c1ee903e6104d1","observation_id":"e52c2686-ec22-46c7-9844-3f42dcd4ec64","resolution":{"observed_at":"2026-08-07T15:16:13.510111Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.476421Z","title":"Motion transformer with global intention localization and local movement refinement,","venue":null,"work_id":"21175e45-05e7-4398-a960-60b03130dcc3","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.015045Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:355d3d58dba81bd8fbaa399c0aa5e047326106bd2b1bf03a7d862e98986baeb7","observation_id":"6968e5e9-2cd4-4bd9-b6a1-b48cffe71cbc","resolution":{"observed_at":"2026-08-07T15:16:13.483053Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.446008Z","title":"Query-centric trajectory prediction,","venue":null,"work_id":"2cf05071-cd19-4e84-b687-7b7a6a6e62e1","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.044080Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:84f5bccdb2f5d8dd92011a8c7d12f8e9444d809c3e8fcf1a9086d980c37dbd0e","observation_id":"894c0a48-f34e-4885-8c54-f3c6e63716f4","resolution":{"observed_at":"2026-08-07T15:16:13.459520Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05982","last_updated":"2024-10-08T12:27:49Z","snapshot_observed_at":"2026-08-16T13:11:31.759434Z","submitted_at":"2024-10-08T12:27:49Z","title":"DeMo: Decoupling Motion Forecasting into Directional Intentions and Dynamic States","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05982","snapshot_observed_at":"2026-08-07T15:16:11.068312Z","title":"Decoupling motion forecast- ing into directional intentions and dynamic states,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.068312Z"},"links":{"cited_paper":"/paper/2410.05982","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:9a395de40009a159b75e9c81e5226ce83fa58d676467669e04f4b920452ea516","observation_id":"fad70c9f-247d-4988-94b2-87ac46e902c1","resolution":{"observed_at":"2026-08-07T15:16:11.068312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.418190Z","title":"Prophnet: Efficient agent-centric motion forecasting with anchor-informed proposals,","venue":null,"work_id":"1e0cf202-b862-4182-a588-d3d25b463aa0","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.115406Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:bcd539c1d7ff0f2a73d4644459f94d68c45d9523c11c2d76b933b9f080b6118e","observation_id":"30ffecdc-fa20-4294-b8a7-c4e50b92e82c","resolution":{"observed_at":"2026-08-07T15:16:13.425641Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1910.05449","last_updated":"2019-10-12T00:34:37Z","snapshot_observed_at":"2026-08-10T19:51:16.852667Z","submitted_at":"2019-10-12T00:34:37Z","title":"MultiPath: Multiple Probabilistic Anchor Trajectory Hypotheses for Behavior Prediction","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1910.05449","snapshot_observed_at":"2026-08-07T15:16:11.155330Z","title":"Multipath: Multiple probabilistic anchor trajectory hypotheses for behavior prediction,","venue":null,"work_id":null,"year":1910},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.155330Z"},"links":{"cited_paper":"/paper/1910.05449","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:80b12fd63a9f8edc8ba9db8c645ad940a95c7a0eeb6e57b47858c8b921c68744","observation_id":"e8881259-4ea7-4196-834e-29a720aa3a1f","resolution":{"observed_at":"2026-08-07T15:16:11.155330Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.382808Z","title":"Multimodal trajectory prediction conditioned on lane-graph traversals,","venue":null,"work_id":"8700d2c4-9184-46e1-8e18-5ffe9afe2984","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.184025Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:efb39c339eebd3a4150694a51a3a4ef168b7748347ad8f369a1a832e3f83d3dc","observation_id":"8b5e768b-37b5-43f7-bfc1-9b8aa1ee6ccd","resolution":{"observed_at":"2026-08-07T15:16:13.395039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.356856Z","title":"Tnt: Target-driven trajectory prediction,","venue":null,"work_id":"cd182c6b-b187-4229-9b2c-fef64ba86c4e","year":2021},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.230670Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:7592202157bb91fe396a09ee70632ca39ef19383d6a34453418532e3621b3734","observation_id":"ca13b4f1-0f21-49e3-bb89-22c88293066d","resolution":{"observed_at":"2026-08-07T15:16:13.365022Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.327854Z","title":"Densetnt: End-to-end trajectory pre- diction from dense goal sets,","venue":null,"work_id":"d2768aff-a5ec-43ba-a797-6248b145bd73","year":2021},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.266397Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:612d8da71ec2fd685a3fc3c848ad4d171789fffeaccbfaacd1066486f6671efe","observation_id":"c7e928de-cdc1-4b26-9836-a1b366b9632c","resolution":{"observed_at":"2026-08-07T15:16:13.337487Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.302756Z","title":"Learning to predict vehicle trajectories with model-based planning,","venue":null,"work_id":"48924f16-f1fa-467f-9b04-40599b01a153","year":2021},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.311399Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:06fab3e082ddb1c64ea4e54582b3d282bcaa5d08212c1db02e3111851290cae7","observation_id":"6126020c-5d16-462b-8dc1-aa1cb14d777e","resolution":{"observed_at":"2026-08-07T15:16:13.312469Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1810.04805","last_updated":"2019-05-24T20:37:26Z","snapshot_observed_at":"2026-08-14T18:16:28.847993Z","submitted_at":"2018-10-11T00:50:01Z","title":"BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.04805","snapshot_observed_at":"2026-08-07T15:16:11.365866Z","title":"Bert: Pre-training of deep bidirectional transformers for language understanding,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.365866Z"},"links":{"cited_paper":"/paper/1810.04805","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:daa349052e7319fdf4ccf1d21cf6d7ba90531eab4bb59843a36137b788d98651","observation_id":"9b361510-2aa2-4046-81f0-c1337e6887ea","resolution":{"observed_at":"2026-08-07T15:16:11.365866Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-16T09:25:53.087782Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-07T15:16:11.405944Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale,","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.405944Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:340b06bb65007c1049d250e97e75fc2af300408cf74decc2693932b50e5f890a","observation_id":"62cc7bf0-284b-46b0-9407-4ba955f98159","resolution":{"observed_at":"2026-08-07T15:16:11.405944Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.273655Z","title":"Taskprompter: Spatial-channel multi-task prompt- ing for dense scene understanding,","venue":null,"work_id":"c4624b5b-8e28-478c-b82e-a898885041d6","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.447683Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:471e47c76f849cd8ab013f6305e263cb807c4cad3779f54704a671bd34d4c3c6","observation_id":"94168420-6d8e-4f4c-a1fc-7bd56dfe2d2b","resolution":{"observed_at":"2026-08-07T15:16:13.281956Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.242020Z","title":"Attention is all you need,","venue":null,"work_id":"e714542e-802d-4e71-bae0-83ec2160683f","year":2017},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.476920Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:72ada405b8a62f2325946f6c4d357e58ba0c0dfd2b5c0ab1b523caed88a7d706","observation_id":"13b64dd8-aac4-4bd8-9d93-c83e022343ec","resolution":{"observed_at":"2026-08-07T15:16:13.251061Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.221637Z","title":"Multiple futures prediction,","venue":null,"work_id":"08b63617-3dca-454e-96d0-52c2ff8de07d","year":2019},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.545441Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:561ecd8daee4c1ba0d1abb7f1d6ebc9ecaacd27f61d7e9c405da2860b965d54e","observation_id":"fa198b32-480e-41d4-9fc9-fced882391cc","resolution":{"observed_at":"2026-08-07T15:16:13.228167Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.199010Z","title":"Hgcn-gjs: Hierar- chical graph convolutional network with groupwise joint sampling for trajectory prediction,","venue":null,"work_id":"b2082a75-7085-4ca1-a16f-31760221c66d","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.576814Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:22134933e70ae09dfe1e0d912d3efff944f1fa8395b6921cfd53d2ed389c690c","observation_id":"b8c24edc-1a37-4852-adcd-a549d928b496","resolution":{"observed_at":"2026-08-07T15:16:13.205795Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.168294Z","title":"Hdgt: Heterogeneous driving graph transformer for multi-agent trajectory prediction via scene encoding,","venue":null,"work_id":"628fe6a6-6c11-4dbd-a333-ad36cbb2d248","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.595038Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:f2aaaee3290a99844b3d257061ccf148c71a2137b02008c4da0836c7946ac1d6","observation_id":"74a0e1a7-5252-4116-a7fc-e729a1ac0b37","resolution":{"observed_at":"2026-08-07T15:16:13.174485Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.145267Z","title":"Real-time motion prediction via heterogeneous polyline transformer with relative pose encoding,","venue":null,"work_id":"b9c7edff-d404-4c22-8336-31f9b1a111a8","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.620839Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:0d561ae272294268f10c6f01007385c47225461569ba7bb7fb23676525cf28bc","observation_id":"c38d1294-9b40-4b59-b811-88e112ffd9b7","resolution":{"observed_at":"2026-08-07T15:16:13.153604Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.117105Z","title":"Smartrefine: A scenario-adaptive refinement framework for efficient motion prediction,","venue":null,"work_id":"5f50ea9a-71f1-411f-825b-5c3e057359e8","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.662120Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:990ecc3b5b7043ade51c94d20854290f0e9cca926137ff40199a8cc32619df76","observation_id":"d90c0dec-c110-428d-b120-cd0377572537","resolution":{"observed_at":"2026-08-07T15:16:13.126867Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.092973Z","title":"End-to-end object detection with transformers,","venue":null,"work_id":"eede876f-8232-40b8-91ed-6ee36d5338b6","year":2020},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.688906Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:555bbdcb8af1589874d7c7591bc9171633184ee8f7dc030535500ce2c750fc95","observation_id":"3867940e-8332-4105-abae-65a3e39fe629","resolution":{"observed_at":"2026-08-07T15:16:13.100116Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.064619Z","title":"Rethinking imitation-based planners for autonomous driving,","venue":null,"work_id":"0370dc75-7776-4bed-8bf7-61ff4e46b902","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.726077Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:69bbb20560decc7b23b1b9103768086a4353f1889c30778657be28f689b44bf2","observation_id":"783bbc65-13f4-4503-8c85-fc5874b4b8af","resolution":{"observed_at":"2026-08-07T15:16:13.071523Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.14327","last_updated":"2024-04-22T16:38:41Z","snapshot_observed_at":"2026-08-16T13:58:49.202638Z","submitted_at":"2024-04-22T16:38:41Z","title":"PLUTO: Pushing the Limit of Imitation Learning-based Planning for Autonomous Driving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.14327","snapshot_observed_at":"2026-08-07T15:16:11.756863Z","title":"Pluto: Pushing the limit of imita- tion learning-based planning for autonomous driving,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.756863Z"},"links":{"cited_paper":"/paper/2404.14327","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:eb76893ebb420bc81a74f1b23e5bcbb177574d6ea99885a3432686ad3b3b1044","observation_id":"1c0db970-c1e0-45e5-a04a-1e2763a8a121","resolution":{"observed_at":"2026-08-07T15:16:11.756863Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.037386Z","title":"Vmamba: Visual state space model,","venue":null,"work_id":"63a6df43-f8bd-430d-b00b-b848be2a1e6b","year":2025},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.789035Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:e13353dab80c200279723d2ca6d92c07d8bd2647c90509cbd64f3b24c105294e","observation_id":"2c8175b6-5c78-4b43-9b5f-28ba53079cc4","resolution":{"observed_at":"2026-08-07T15:16:13.045979Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.016494Z","title":"Videomamba: State space model for efficient video understanding,","venue":null,"work_id":"83a65ad1-803d-486a-9046-558025513f71","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.828704Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:e3ea02e4f7e7b0f48d8a6dbaa6ccbbc0f017c21d76b23f98d8cce7f23c74f8a7","observation_id":"0c446726-881e-4f10-99c4-0c48deaea239","resolution":{"observed_at":"2026-08-07T15:16:13.023454Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.14520","last_updated":"2025-01-08T11:03:00Z","snapshot_observed_at":"2026-08-16T14:07:37.083558Z","submitted_at":"2024-03-21T16:17:57Z","title":"Cobra: Extending Mamba to Multi-Modal Large Language Model for Efficient Inference","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.14520","snapshot_observed_at":"2026-08-07T15:16:11.859602Z","title":"Cobra: Extending mamba to multi-modal large language model for efficient inference,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.859602Z"},"links":{"cited_paper":"/paper/2403.14520","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:17ac96ef531c4184ccd8fbce2dbd4c64f6a3008e2812a097efa6936f10828df0","observation_id":"ebd480ea-2b44-49b9-b6f3-cdd3af81ccf8","resolution":{"observed_at":"2026-08-07T15:16:11.859602Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.00752","last_updated":"2024-05-31T17:55:27Z","snapshot_observed_at":"2026-08-17T20:47:46.242385Z","submitted_at":"2023-12-01T18:01:34Z","title":"Mamba: Linear-Time Sequence Modeling with Selective State Spaces","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.00752","snapshot_observed_at":"2026-08-07T15:16:11.886964Z","title":"Mamba: Linear-time sequence modeling with selective state spaces,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.886964Z"},"links":{"cited_paper":"/paper/2312.00752","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:5fd6304c28a85eae8b50648254b5191e87ecc27892e01c9cbdc1a8d61cd57706","observation_id":"5233c138-c85a-4100-af16-2034d7a6bd01","resolution":{"observed_at":"2026-08-07T15:16:11.886964Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1607.06450","last_updated":"2016-07-21T19:57:52Z","snapshot_observed_at":"2026-08-15T04:53:45.483331Z","submitted_at":"2016-07-21T19:57:52Z","title":"Layer Normalization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1607.06450","snapshot_observed_at":"2026-08-07T15:16:11.913465Z","title":"Layer normalization,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.913465Z"},"links":{"cited_paper":"/paper/1607.06450","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:8498b396080b3143edf5a235ebe0f23b5898955f88bfc385fc19c75dd91c555e","observation_id":"35ae53b7-348b-4064-98d5-447debed0952","resolution":{"observed_at":"2026-08-07T15:16:11.913465Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.994371Z","title":"Ganet: Goal area network for motion forecasting,","venue":null,"work_id":"fa6c7f04-6c20-41f6-960e-1ab5048cd21c","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.939371Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:fbf7543d9a16b267cc8b12c8d7c04b81b55b2972e162f1d34780f5b9684e1d80","observation_id":"aff6c694-da51-425a-a2ad-86ed7553b2b9","resolution":{"observed_at":"2026-08-07T15:16:13.000164Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.928943Z","title":"Motion forecasting in continuous driving,","venue":null,"work_id":"65351a29-d8b1-4cfd-960b-76d28fb6d117","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.962427Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:4c9420b1d4cc1e6a66fdb071340b1f2e354e35ba5d6c08ebfbb119bb99985f76","observation_id":"9d0c5e0a-ed45-4ea5-8046-901d346cc9c6","resolution":{"observed_at":"2026-08-07T15:16:12.957652Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2207.06553","last_updated":"2022-07-13T23:25:30Z","snapshot_observed_at":"2026-08-16T16:46:09.089907Z","submitted_at":"2022-07-13T23:25:30Z","title":"QML for Argoverse 2 Motion Forecasting Challenge","version":1},"cited_work":{"arxiv_id":"2207.06553","doi":null,"metadata_source":"pith","pith_arxiv_id":"2207.06553","snapshot_observed_at":"2026-08-07T15:16:12.307406Z","title":"QML for Argoverse 2 Motion Forecasting Challenge","venue":"cs.CV","work_id":"4e27976b-b9f9-4d0c-a8ab-0848ac0d4f9b","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.984473Z"},"links":{"cited_paper":"/paper/2207.06553","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:64ba16319920fcbad4a72dd26373a77bf2e53f6b076d0caeb0d5c7f64a80c43d","observation_id":"f7c9a6e4-5626-4599-9dcc-cae74f1b7701","resolution":{"observed_at":"2026-08-07T15:16:12.360966Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.855508Z","title":"Macformer: Map-agent coupled transformer for real- time and robust trajectory prediction,","venue":null,"work_id":"49fd4440-88b2-42c7-826e-f9a67703d369","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.007974Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:159490e7511a76526b3ffbd9ebacb934d4a6399b439cc87c701559b4b519c691","observation_id":"3120dcf8-6e82-4c0f-b977-3bab8c5d0693","resolution":{"observed_at":"2026-08-07T15:16:12.881223Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.07934","last_updated":"2022-07-01T03:19:40Z","snapshot_observed_at":"2026-08-16T16:52:37.928972Z","submitted_at":"2022-06-16T05:56:24Z","title":"BANet: Motion Forecasting with Boundary Aware Network","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.07934","snapshot_observed_at":"2026-08-07T15:16:12.034224Z","title":"Banet: Motion forecasting with boundary aware network,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.034224Z"},"links":{"cited_paper":"/paper/2206.07934","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:08ee7064e50df18679a716a4aac53183b9d930d12f858091449753f8cb167112","observation_id":"28ca253b-09b0-40c5-8984-955f004770a8","resolution":{"observed_at":"2026-08-07T15:16:12.034224Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.780005Z","title":"Dynamic scenario representation learning for motion forecasting with heterogeneous graph convolu- tional recurrent networks,","venue":null,"work_id":"c875af78-9063-4b0f-857b-ae9edf2ec9ff","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.078017Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:a7e6f21b7895ec43c79ee855fa7d83d24f557ea7d09d1d27d2a8a1b21b3d12e6","observation_id":"427eabd1-37a0-46b3-a9c6-e57f6a1584af","resolution":{"observed_at":"2026-08-07T15:16:12.808959Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1711.05101","last_updated":"2019-01-04T21:01:49Z","snapshot_observed_at":"2026-08-14T20:13:52.872565Z","submitted_at":"2017-11-14T14:24:06Z","title":"Decoupled Weight Decay Regularization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.05101","snapshot_observed_at":"2026-08-07T15:16:12.111237Z","title":"Decoupled weight decay regularization,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.111237Z"},"links":{"cited_paper":"/paper/1711.05101","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:38ddb324c67fc1ef799fef608f002e551a610c78c291e29b60cddc2f56fc8092","observation_id":"87e038bc-c546-4f4c-9bed-4e41eac07fb9","resolution":{"observed_at":"2026-08-07T15:16:12.111237Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.706608Z","title":"Vision mamba: Efficient visual representation learning with bidirectional state space model,","venue":null,"work_id":"ca313464-e1bd-4ce0-b85a-88c15f3aa00a","year":null},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.146074Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:53af7de8e95039c1c8477435c2a33eb1c17f4b419781dbff450f21822df93e0b","observation_id":"b840834f-67af-4c36-b31d-1e2d63a854ae","resolution":{"observed_at":"2026-08-07T15:16:12.743804Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-14T20:21:52.049965Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning"},"reference_resolution":{"displayed":43,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":12,"verified_exact":1,"verified_fuzzy":30},"total_outbound_references":43},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 21 August 2026, this Paper Citation Record lists 43 of 43 outbound references and 0 inbound Pith citation observations for arXiv:2505.15703."}