{"as_of":"2026-08-08T11:51:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:90633b3d06b196512ac65397fc03b8f89010f28bd8b21ca6ff0351d643e84fed","coverage":[{"denominator":43,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":43,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T15:16:12.146074Z","state":"measured"},{"denominator":43,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":43,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.15703/citation-record","integrity":"/paper/2505.15703/integrity","json":"/paper/2505.15703/citation-record.json","paper":"/paper/2505.15703"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2301.00493","last_updated":"2023-01-02T00:36:22Z","snapshot_observed_at":"2026-07-06T14:36:37.805357Z","submitted_at":"2023-01-02T00:36:22Z","title":"Argoverse 2: Next Generation Datasets for Self-Driving Perception and Forecasting","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.00493","snapshot_observed_at":"2026-08-07T15:16:10.721562Z","title":"Argoverse 2: Next generation datasets for self-driving perception and forecasting,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.721562Z"},"links":{"cited_paper":"/paper/2301.00493","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:1bd67f799d995f49d523e0f5b9cfe43d7c58903b23be2699b22e6faccfd9bce7","observation_id":"eb7f8e35-be2f-4d6f-bfa7-96cf1ebedeae","resolution":{"observed_at":"2026-08-07T15:16:10.721562Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.636531Z","title":"A survey on trajectory-prediction methods for autonomous driving,","venue":null,"work_id":"87c2ae12-9b52-4ca9-82dd-31a35d31888c","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.773288Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:670596111f5964124fecc31c08f4d39bd544b855e56feab43e1008aba6530f5b","observation_id":"06e5b162-8cbe-46e6-90d8-50add44f764d","resolution":{"observed_at":"2026-08-07T15:16:13.644669Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.611757Z","title":"Vectornet: Encoding hd maps and agent dynamics from vectorized representation,","venue":null,"work_id":"b7a52fae-2348-4064-8bed-3065ec1a1a71","year":2020},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.814015Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:687d0caa89766d13ec3bc46149bc56badb6d4f649ed536ba432c6429baf8f0a0","observation_id":"2473a5d5-517a-4ec4-a8c6-8c522a1b46a8","resolution":{"observed_at":"2026-08-07T15:16:13.618951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.589855Z","title":"Learning lane graph representations for motion forecasting,","venue":null,"work_id":"c6c0434d-944e-400c-90f5-aeaaf5b1c762","year":2020},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.851157Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:54e6eccd7361d6377ac462a211a3b5ab053c6446f78742dbf6c8f71227f7f3b3","observation_id":"da0a02a5-3c41-4902-9ab3-d6dfcae877fb","resolution":{"observed_at":"2026-08-07T15:16:13.596594Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.568681Z","title":"Simpl: A simple and efficient multi-agent motion prediction baseline for autonomous driving,","venue":null,"work_id":"592561c8-82d6-4527-b2cd-5b1b4d19062b","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.873717Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:09a0a40592d4a873a573d295d4f464eaaff6494af0c1e72f5df9aeaed84c7df7","observation_id":"23c89954-c98d-4559-b397-c7d0456c3e65","resolution":{"observed_at":"2026-08-07T15:16:13.576163Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2106.08417","last_updated":"2022-03-04T20:25:25Z","snapshot_observed_at":"2026-07-06T11:19:38.321360Z","submitted_at":"2021-06-15T20:20:44Z","title":"Scene Transformer: A unified architecture for predicting multiple agent trajectories","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.08417","snapshot_observed_at":"2026-08-07T15:16:10.889911Z","title":"Scene transformer: A unified architecture for predicting multiple agent tra- jectories,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.889911Z"},"links":{"cited_paper":"/paper/2106.08417","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:b13c65665f3862d1b46de4b939e96983ed4d8bf79f84627f998e2894ebdcf6b6","observation_id":"56fd1171-cf96-4879-8ffd-493d4f656378","resolution":{"observed_at":"2026-08-07T15:16:10.889911Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.548359Z","title":"Forecast-mae: Self-supervised pre- training for motion forecasting with masked autoencoders,","venue":null,"work_id":"90432211-560f-4bd8-81e4-e2c602d73116","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.917439Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:740753914bc4aa0dbb822765dc3240f39d23fb1f44f84c7abf0ccb2ad011a299","observation_id":"9fcb46f8-a407-4cb3-9072-b0c84ae37c97","resolution":{"observed_at":"2026-08-07T15:16:13.555570Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.523172Z","title":"Gorela: Go relative for viewpoint-invariant motion forecasting,","venue":null,"work_id":"208d4d60-0ca2-41b2-af41-a61967f39159","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.937366Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:de06ff12c60e9ef541148acd1561301dee928f4e758c8676a46f44e52fbbacf3","observation_id":"c3ad91c7-2881-46ae-83dd-a7020e38b0ad","resolution":{"observed_at":"2026-08-07T15:16:13.531818Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.499582Z","title":"Multipath++: Efficient information fusion and trajectory aggregation for behavior prediction,","venue":null,"work_id":"f815bff8-fa34-46c5-b2b7-55ffd0b37b95","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:10.969322Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:3aa7a27ad01c1e3283f670e479f51a6d818d3f9ab3d759a323732eccf5e34c80","observation_id":"e52c2686-ec22-46c7-9844-3f42dcd4ec64","resolution":{"observed_at":"2026-08-07T15:16:13.510111Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.476421Z","title":"Motion transformer with global intention localization and local movement refinement,","venue":null,"work_id":"21175e45-05e7-4398-a960-60b03130dcc3","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.015045Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:512f9b5eba9d51c7086bb515e0216c59ce9d38295d04329da8d86ca03f3e1402","observation_id":"6968e5e9-2cd4-4bd9-b6a1-b48cffe71cbc","resolution":{"observed_at":"2026-08-07T15:16:13.483053Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.446008Z","title":"Query-centric trajectory prediction,","venue":null,"work_id":"2cf05071-cd19-4e84-b687-7b7a6a6e62e1","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.044080Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:6e61c5727dcad0c366d348d25a43be988a91ade0739f276f9c25b8efe5d9b2ef","observation_id":"894c0a48-f34e-4885-8c54-f3c6e63716f4","resolution":{"observed_at":"2026-08-07T15:16:13.459520Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05982","last_updated":"2024-10-08T12:27:49Z","snapshot_observed_at":"2026-08-03T07:10:48.183309Z","submitted_at":"2024-10-08T12:27:49Z","title":"DeMo: Decoupling Motion Forecasting into Directional Intentions and Dynamic States","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05982","snapshot_observed_at":"2026-08-07T15:16:11.068312Z","title":"Decoupling motion forecast- ing into directional intentions and dynamic states,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.068312Z"},"links":{"cited_paper":"/paper/2410.05982","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:aa87f0ddd18b58618b9e81e8791a7483fb4ef19f56aba5279f11482958bf0ed3","observation_id":"fad70c9f-247d-4988-94b2-87ac46e902c1","resolution":{"observed_at":"2026-08-07T15:16:11.068312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.418190Z","title":"Prophnet: Efficient agent-centric motion forecasting with anchor-informed proposals,","venue":null,"work_id":"1e0cf202-b862-4182-a588-d3d25b463aa0","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.115406Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:83c9ca6a1837cb101c3126c0ff2cf1c19a7e9e255701df8e9b94263ef06dec74","observation_id":"30ffecdc-fa20-4294-b8a7-c4e50b92e82c","resolution":{"observed_at":"2026-08-07T15:16:13.425641Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1910.05449","last_updated":"2019-10-12T00:34:37Z","snapshot_observed_at":"2026-08-04T23:16:21.521818Z","submitted_at":"2019-10-12T00:34:37Z","title":"MultiPath: Multiple Probabilistic Anchor Trajectory Hypotheses for Behavior Prediction","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1910.05449","snapshot_observed_at":"2026-08-07T15:16:11.155330Z","title":"Multipath: Multiple probabilistic anchor trajectory hypotheses for behavior prediction,","venue":null,"work_id":null,"year":1910},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.155330Z"},"links":{"cited_paper":"/paper/1910.05449","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:2fa47113c890d5d172fa806d305000f614f5873ee2585bf98fd514dc93b662f5","observation_id":"e8881259-4ea7-4196-834e-29a720aa3a1f","resolution":{"observed_at":"2026-08-07T15:16:11.155330Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.382808Z","title":"Multimodal trajectory prediction conditioned on lane-graph traversals,","venue":null,"work_id":"8700d2c4-9184-46e1-8e18-5ffe9afe2984","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.184025Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:f54c56c84c0c0830a05f7c7503c2fcc4fb511c9bbe7a1215e06a12090e3e3fc6","observation_id":"8b5e768b-37b5-43f7-bfc1-9b8aa1ee6ccd","resolution":{"observed_at":"2026-08-07T15:16:13.395039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.356856Z","title":"Tnt: Target-driven trajectory prediction,","venue":null,"work_id":"cd182c6b-b187-4229-9b2c-fef64ba86c4e","year":2021},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.230670Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:3dfef8a035cffffb13b7d79b6ff40622bfd3aaef412456fe3607e1f1f3fea6e6","observation_id":"ca13b4f1-0f21-49e3-bb89-22c88293066d","resolution":{"observed_at":"2026-08-07T15:16:13.365022Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.327854Z","title":"Densetnt: End-to-end trajectory pre- diction from dense goal sets,","venue":null,"work_id":"d2768aff-a5ec-43ba-a797-6248b145bd73","year":2021},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.266397Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:1b88f5c53d5af9b1ce7910f88aa1cec5e96d5a8ccc8b11ec9a2ec10d5db56176","observation_id":"c7e928de-cdc1-4b26-9836-a1b366b9632c","resolution":{"observed_at":"2026-08-07T15:16:13.337487Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.302756Z","title":"Learning to predict vehicle trajectories with model-based planning,","venue":null,"work_id":"48924f16-f1fa-467f-9b04-40599b01a153","year":2021},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.311399Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:b4b6a9ec76f11dba62501e76e5a4a7b16abdd8507903b50652569aef20202098","observation_id":"6126020c-5d16-462b-8dc1-aa1cb14d777e","resolution":{"observed_at":"2026-08-07T15:16:13.312469Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1810.04805","last_updated":"2019-05-24T20:37:26Z","snapshot_observed_at":"2026-07-30T09:12:38.100527Z","submitted_at":"2018-10-11T00:50:01Z","title":"BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.04805","snapshot_observed_at":"2026-08-07T15:16:11.365866Z","title":"Bert: Pre-training of deep bidirectional transformers for language understanding,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.365866Z"},"links":{"cited_paper":"/paper/1810.04805","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:16404f128594e591c8fa996faa2a8e560ef80928b7b3de0ba98a5d0a103441fa","observation_id":"9b361510-2aa2-4046-81f0-c1337e6887ea","resolution":{"observed_at":"2026-08-07T15:16:11.365866Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-07T15:16:11.405944Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale,","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.405944Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:74079d7124c54a6573e5f38f676af6efb63849dca9e6700a5f292350d8f58116","observation_id":"62cc7bf0-284b-46b0-9407-4ba955f98159","resolution":{"observed_at":"2026-08-07T15:16:11.405944Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.273655Z","title":"Taskprompter: Spatial-channel multi-task prompt- ing for dense scene understanding,","venue":null,"work_id":"c4624b5b-8e28-478c-b82e-a898885041d6","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.447683Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:7e7febd8e67616750c8adb46054875d0c3f403088d8a4f45c49e7f2d31c9e8e6","observation_id":"94168420-6d8e-4f4c-a1fc-7bd56dfe2d2b","resolution":{"observed_at":"2026-08-07T15:16:13.281956Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.242020Z","title":"Attention is all you need,","venue":null,"work_id":"e714542e-802d-4e71-bae0-83ec2160683f","year":2017},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.476920Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:649fec093769238829dc77628bc3dacbb998d61e636a435bf635e2756706acb5","observation_id":"13b64dd8-aac4-4bd8-9d93-c83e022343ec","resolution":{"observed_at":"2026-08-07T15:16:13.251061Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.221637Z","title":"Multiple futures prediction,","venue":null,"work_id":"08b63617-3dca-454e-96d0-52c2ff8de07d","year":2019},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.545441Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:b07e6143d0fc4382f810332fefdb13402a5fad5938fdb4643c2fffb9999e3141","observation_id":"fa198b32-480e-41d4-9fc9-fced882391cc","resolution":{"observed_at":"2026-08-07T15:16:13.228167Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.199010Z","title":"Hgcn-gjs: Hierar- chical graph convolutional network with groupwise joint sampling for trajectory prediction,","venue":null,"work_id":"b2082a75-7085-4ca1-a16f-31760221c66d","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.576814Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:8c75d3681aa56482877e4fa0d2563fda8bcafcd2d5e77d44a8b5b57d543af28b","observation_id":"b8c24edc-1a37-4852-adcd-a549d928b496","resolution":{"observed_at":"2026-08-07T15:16:13.205795Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.168294Z","title":"Hdgt: Heterogeneous driving graph transformer for multi-agent trajectory prediction via scene encoding,","venue":null,"work_id":"628fe6a6-6c11-4dbd-a333-ad36cbb2d248","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.595038Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:4652bea0ef400044c933f5b5c77a0dfae81652cd46b0fcff31a11057a082ddd5","observation_id":"74a0e1a7-5252-4116-a7fc-e729a1ac0b37","resolution":{"observed_at":"2026-08-07T15:16:13.174485Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.145267Z","title":"Real-time motion prediction via heterogeneous polyline transformer with relative pose encoding,","venue":null,"work_id":"b9c7edff-d404-4c22-8336-31f9b1a111a8","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.620839Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:7a8efe72e92fb0de890abdeac1765b3130ab9d4b5c5e4153d8b932bf638169f6","observation_id":"c38d1294-9b40-4b59-b811-88e112ffd9b7","resolution":{"observed_at":"2026-08-07T15:16:13.153604Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.117105Z","title":"Smartrefine: A scenario-adaptive refinement framework for efficient motion prediction,","venue":null,"work_id":"5f50ea9a-71f1-411f-825b-5c3e057359e8","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.662120Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:8e561682cc7b8a10db0118f69382359ab2e759465eb567295dd0fa37c3681d63","observation_id":"d90c0dec-c110-428d-b120-cd0377572537","resolution":{"observed_at":"2026-08-07T15:16:13.126867Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.092973Z","title":"End-to-end object detection with transformers,","venue":null,"work_id":"eede876f-8232-40b8-91ed-6ee36d5338b6","year":2020},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.688906Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:8c957406c9d123204b2c5573107ac4fb1d6f8e7ae95d5bcb0d31553c9ef329d4","observation_id":"3867940e-8332-4105-abae-65a3e39fe629","resolution":{"observed_at":"2026-08-07T15:16:13.100116Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.064619Z","title":"Rethinking imitation-based planners for autonomous driving,","venue":null,"work_id":"0370dc75-7776-4bed-8bf7-61ff4e46b902","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.726077Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:5c841b03b1ea55cc008510cc4c7ecd28e368325daa43393daf2346b5281b33dd","observation_id":"783bbc65-13f4-4503-8c85-fc5874b4b8af","resolution":{"observed_at":"2026-08-07T15:16:13.071523Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.14327","last_updated":"2024-04-22T16:38:41Z","snapshot_observed_at":"2026-08-08T03:52:57.869549Z","submitted_at":"2024-04-22T16:38:41Z","title":"PLUTO: Pushing the Limit of Imitation Learning-based Planning for Autonomous Driving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.14327","snapshot_observed_at":"2026-08-07T15:16:11.756863Z","title":"Pluto: Pushing the limit of imita- tion learning-based planning for autonomous driving,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.756863Z"},"links":{"cited_paper":"/paper/2404.14327","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:c31e8135360dc9750dc55181e22f8157ef8ef49677b8b70d41d781fb66c9362d","observation_id":"1c0db970-c1e0-45e5-a04a-1e2763a8a121","resolution":{"observed_at":"2026-08-07T15:16:11.756863Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.037386Z","title":"Vmamba: Visual state space model,","venue":null,"work_id":"63a6df43-f8bd-430d-b00b-b848be2a1e6b","year":2025},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.789035Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:3333dd886d27623749958c0363d2e3204d90b2d023d6bbe176e9ddb3feaa4a03","observation_id":"2c8175b6-5c78-4b43-9b5f-28ba53079cc4","resolution":{"observed_at":"2026-08-07T15:16:13.045979Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:13.016494Z","title":"Videomamba: State space model for efficient video understanding,","venue":null,"work_id":"83a65ad1-803d-486a-9046-558025513f71","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.828704Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:1908c339a3c88c6620f927d6add621f3386c5e30109ce0f9d6b4f7200bc2c77a","observation_id":"0c446726-881e-4f10-99c4-0c48deaea239","resolution":{"observed_at":"2026-08-07T15:16:13.023454Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.14520","last_updated":"2025-01-08T11:03:00Z","snapshot_observed_at":"2026-07-06T17:48:24.076444Z","submitted_at":"2024-03-21T16:17:57Z","title":"Cobra: Extending Mamba to Multi-Modal Large Language Model for Efficient Inference","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.14520","snapshot_observed_at":"2026-08-07T15:16:11.859602Z","title":"Cobra: Extending mamba to multi-modal large language model for efficient inference,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.859602Z"},"links":{"cited_paper":"/paper/2403.14520","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:f407e8f12a3a15dd41f43d851d6bf46f710989f2c5616392e63ed0eb49dc85a4","observation_id":"ebd480ea-2b44-49b9-b6f3-cdd3af81ccf8","resolution":{"observed_at":"2026-08-07T15:16:11.859602Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.00752","last_updated":"2024-05-31T17:55:27Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-01T18:01:34Z","title":"Mamba: Linear-Time Sequence Modeling with Selective State Spaces","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.00752","snapshot_observed_at":"2026-08-07T15:16:11.886964Z","title":"Mamba: Linear-time sequence modeling with selective state spaces,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.886964Z"},"links":{"cited_paper":"/paper/2312.00752","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:ecff1ba7f00d1533873d3763daeb0312cc7b9e9505201f33cc55fe782cff7208","observation_id":"5233c138-c85a-4100-af16-2034d7a6bd01","resolution":{"observed_at":"2026-08-07T15:16:11.886964Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1607.06450","last_updated":"2016-07-21T19:57:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2016-07-21T19:57:52Z","title":"Layer Normalization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1607.06450","snapshot_observed_at":"2026-08-07T15:16:11.913465Z","title":"Layer normalization,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.913465Z"},"links":{"cited_paper":"/paper/1607.06450","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:b3801d19e6d051a966abc0877433e9f0bd1144e67aaeb17cc945ffc7f408f53a","observation_id":"35ae53b7-348b-4064-98d5-447debed0952","resolution":{"observed_at":"2026-08-07T15:16:11.913465Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.994371Z","title":"Ganet: Goal area network for motion forecasting,","venue":null,"work_id":"fa6c7f04-6c20-41f6-960e-1ab5048cd21c","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.939371Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:c86b98dc879818a5988d9694644d42625c8fe22df76c9ba87d447f79531a516e","observation_id":"aff6c694-da51-425a-a2ad-86ed7553b2b9","resolution":{"observed_at":"2026-08-07T15:16:13.000164Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.928943Z","title":"Motion forecasting in continuous driving,","venue":null,"work_id":"65351a29-d8b1-4cfd-960b-76d28fb6d117","year":2024},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.962427Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:7411407ffe114a10a6d7d0023a4dafa889071f1678501c8c79e76fcc0c4d4806","observation_id":"9d0c5e0a-ed45-4ea5-8046-901d346cc9c6","resolution":{"observed_at":"2026-08-07T15:16:12.957652Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2207.06553","last_updated":"2022-07-13T23:25:30Z","snapshot_observed_at":"2026-08-05T00:39:40.376772Z","submitted_at":"2022-07-13T23:25:30Z","title":"QML for Argoverse 2 Motion Forecasting Challenge","version":1},"cited_work":{"arxiv_id":"2207.06553","doi":null,"metadata_source":"pith","pith_arxiv_id":"2207.06553","snapshot_observed_at":"2026-08-07T15:16:12.307406Z","title":"QML for Argoverse 2 Motion Forecasting Challenge","venue":"cs.CV","work_id":"4e27976b-b9f9-4d0c-a8ab-0848ac0d4f9b","year":2022},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:11.984473Z"},"links":{"cited_paper":"/paper/2207.06553","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:7fb304dcbc9f78ad68b98b622e1866ea0d70e753284d6d1c1674f971b12ea91d","observation_id":"f7c9a6e4-5626-4599-9dcc-cae74f1b7701","resolution":{"observed_at":"2026-08-07T15:16:12.360966Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.855508Z","title":"Macformer: Map-agent coupled transformer for real- time and robust trajectory prediction,","venue":null,"work_id":"49fd4440-88b2-42c7-826e-f9a67703d369","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.007974Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:f0b9f0cfc93d858a4d1335775653bbcf91ec8dfd2cde5dcef15b5ba4f6728e0c","observation_id":"3120dcf8-6e82-4c0f-b977-3bab8c5d0693","resolution":{"observed_at":"2026-08-07T15:16:12.881223Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.07934","last_updated":"2022-07-01T03:19:40Z","snapshot_observed_at":"2026-07-06T13:21:28.206982Z","submitted_at":"2022-06-16T05:56:24Z","title":"BANet: Motion Forecasting with Boundary Aware Network","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2206.07934","snapshot_observed_at":"2026-08-07T15:16:12.034224Z","title":"Banet: Motion forecasting with boundary aware network,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.034224Z"},"links":{"cited_paper":"/paper/2206.07934","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:fa093a60befa51908ce11219e1a3d5d93eb6a479d166b802eb9465287f270d8a","observation_id":"28ca253b-09b0-40c5-8984-955f004770a8","resolution":{"observed_at":"2026-08-07T15:16:12.034224Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.780005Z","title":"Dynamic scenario representation learning for motion forecasting with heterogeneous graph convolu- tional recurrent networks,","venue":null,"work_id":"c875af78-9063-4b0f-857b-ae9edf2ec9ff","year":2023},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.078017Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:df15a4d45460c53cf16dfaaa9b3e3940510e6cc7834321ad89bd3b6525fd938a","observation_id":"427eabd1-37a0-46b3-a9c6-e57f6a1584af","resolution":{"observed_at":"2026-08-07T15:16:12.808959Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1711.05101","last_updated":"2019-01-04T21:01:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2017-11-14T14:24:06Z","title":"Decoupled Weight Decay Regularization","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1711.05101","snapshot_observed_at":"2026-08-07T15:16:12.111237Z","title":"Decoupled weight decay regularization,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.111237Z"},"links":{"cited_paper":"/paper/1711.05101","citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:bad754eeaa111b21e5ed88866aebe53543f2b3d6e410aad33e12ef8cb1b1a969","observation_id":"87e038bc-c546-4f4c-9bed-4e41eac07fb9","resolution":{"observed_at":"2026-08-07T15:16:12.111237Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T15:16:12.706608Z","title":"Vision mamba: Efficient visual representation learning with bidirectional state space model,","venue":null,"work_id":"ca313464-e1bd-4ce0-b85a-88c15f3aa00a","year":null},"citing_paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T15:16:12.146074Z"},"links":{"citing_paper":"/paper/2505.15703"},"observation_digest":"sha256:51990cd6a0806c6b509ee5a7b0117c451bb356ce535e9bbe5b476347565aeba8","observation_id":"b840834f-67af-4c36-b31d-1e2d63a854ae","resolution":{"observed_at":"2026-08-07T15:16:12.743804Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.15703","last_updated":"2025-05-21T16:16:52Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-07T15:10:30.247783Z","submitted_at":"2025-05-21T16:16:52Z","title":"HAMF: A Hybrid Attention-Mamba Framework for Joint Scene Context Understanding and Future Motion Representation Learning"},"reference_resolution":{"displayed":43,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":12,"verified_exact":1,"verified_fuzzy":30},"total_outbound_references":43},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 43 of 43 outbound references and 0 inbound Pith citation observations for arXiv:2505.15703."}