{"as_of":"2026-08-15T00:32:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:79cc3acd7b76ca24d3387cc6ec1b6bd8dd32eced761d034f54c87498fe33567e","coverage":[{"denominator":50,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":50,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T23:01:08.557671Z","state":"measured"},{"denominator":60,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":60,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-14T06:32:32.682623+00:00","state":"measured"},{"denominator":10,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":10,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T00:25:04.948807Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T21:10:08.552391Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":"2501.09024","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-07-04T21:10:08.552391Z","title":"Social-llava: Enhancing robot navigation through human-language reasoning in social spaces","venue":null,"work_id":"8821b021-c012-4b4b-a6f3-a9004e6a5b56","year":2024},"citing_paper":{"arxiv_id":"2503.07557","last_updated":"2026-05-02T15:43:17Z","snapshot_observed_at":"2026-08-03T10:34:14.204884Z","submitted_at":"2025-03-10T17:27:17Z","title":"AutoSpatial: Visual-Language Reasoning for Social Robot Navigation through Efficient Spatial Reasoning Learning","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-23T00:26:58.273861Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2503.07557"},"observation_digest":"sha256:1aea2f18966cc60dc43623cd2d55a20ce3b0fed4d57b7b609c70381f1e15fc45","observation_id":"a735bf98-5ff6-4bb9-93db-4398d6f382d8","resolution":{"observed_at":"2026-05-23T00:27:17.796539Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":"2501.09024","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-07-04T21:10:08.552391Z","title":"Social-llava: Enhancing robot navigation through human-language reasoning in social spaces","venue":null,"work_id":"8821b021-c012-4b4b-a6f3-a9004e6a5b56","year":2024},"citing_paper":{"arxiv_id":"2504.05477","last_updated":"2025-04-07T20:16:00Z","snapshot_observed_at":"2026-07-06T21:05:44.765030Z","submitted_at":"2025-04-07T20:16:00Z","title":"Trust Through Transparency: Explainable Social Navigation for Autonomous Mobile Robots via Vision-Language Models","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-22T20:24:46.225249Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2504.05477"},"observation_digest":"sha256:893adcc972f7897015f48abaa0e40eba1d9679d4a48f365aa069abd7e847848b","observation_id":"8030906a-0d0b-48d4-acb5-a22bdabea35b","resolution":{"observed_at":"2026-05-22T20:25:05.406392Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-08-07T00:25:04.948807Z","title":"Payandeh, D","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14233","last_updated":"2025-06-17T06:46:15Z","snapshot_observed_at":"2026-08-14T06:53:50.124735Z","submitted_at":"2025-06-17T06:46:15Z","title":"Narrate2Nav: Real-Time Visual Navigation with Implicit Language Reasoning in Human-Centric Environments","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T00:25:04.948807Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2506.14233"},"observation_digest":"sha256:5916b56934b9f81d0d53609084fbd72ed179deb18560cbb4a44aefc0d6a19a05","observation_id":"8071cd29-f15f-4b4d-ab67-68d890790e7e","resolution":{"observed_at":"2026-08-07T00:25:04.948807Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-08-06T05:41:41.959671Z","title":"Payandeh, D","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2508.01539","last_updated":"2025-08-03T01:46:10Z","snapshot_observed_at":"2026-08-08T15:25:20.465265Z","submitted_at":"2025-08-03T01:46:10Z","title":"HALO: Human Preference Aligned Offline Reward Learning for Robot Navigation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T05:41:41.959671Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2508.01539"},"observation_digest":"sha256:036106e1d0ecff035dc136c0bc0581be1cb5326e0578effe1886e5f419bf8d9a","observation_id":"008c4d89-bd89-41cb-aa6b-87abc623945c","resolution":{"observed_at":"2026-08-06T05:41:41.959671Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-08-04T07:54:07.763791Z","title":"Social-llava: Enhancing robot navigation through human-language reasoning in social spaces,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.23509","last_updated":"2026-07-18T06:00:54Z","snapshot_observed_at":"2026-08-14T23:35:33.401673Z","submitted_at":"2025-10-27T16:47:15Z","title":"Logic-Guided Socially-aware Robot Navigation World Model","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-04T07:54:07.763791Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2510.23509"},"observation_digest":"sha256:25fb158544060d0ffd952a5e9529ff58c2c61c0da2dd5efdebb4bea4d1d62963","observation_id":"c14555a3-8daf-42db-838d-3ffdc0a58d06","resolution":{"observed_at":"2026-08-04T07:54:07.763791Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-08-03T13:51:32.307397Z","title":"Social-llava: Enhancing robot navigation through human-language reasoning in social spaces,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2512.22867","last_updated":"2026-07-03T08:07:47Z","snapshot_observed_at":"2026-08-14T08:24:23.911109Z","submitted_at":"2025-12-28T10:41:39Z","title":"MUSON: A Reasoning-oriented Multimodal Dataset for Socially Compliant Navigation in Urban Environments","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-03T13:51:32.307397Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2512.22867"},"observation_digest":"sha256:4e3ed9b905aebebb1415b00ceaac6e21e647945780b405d5ae0531a8f55a9fc6","observation_id":"e1cd809c-a776-4a25-8902-e60e3e343cdd","resolution":{"observed_at":"2026-08-03T13:51:32.307397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":"2501.09024","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-07-04T21:10:08.552391Z","title":"Social-llava: Enhancing robot navigation through human-language reasoning in social spaces","venue":null,"work_id":"8821b021-c012-4b4b-a6f3-a9004e6a5b56","year":2024},"citing_paper":{"arxiv_id":"2605.13321","last_updated":"2026-05-13T10:34:47Z","snapshot_observed_at":"2026-08-13T09:33:39.207244Z","submitted_at":"2026-05-13T10:34:47Z","title":"HCSG: Human-Centric Semantic-Geometric Reasoning for Vision-Language Navigation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-14T18:04:27.823075Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2605.13321"},"observation_digest":"sha256:4a91039014d84b4ee636c979662c3983dc2255d737c95f78e39c10b0e34722d2","observation_id":"2e49b047-0d85-4fa4-a951-0374467b9cf5","resolution":{"observed_at":"2026-05-14T18:07:33.852632Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":"2501.09024","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-07-04T21:10:08.552391Z","title":"Social-llava: Enhancing robot navigation through human-language reasoning in social spaces","venue":null,"work_id":"8821b021-c012-4b4b-a6f3-a9004e6a5b56","year":2024},"citing_paper":{"arxiv_id":"2606.10495","last_updated":"2026-06-09T07:18:01Z","snapshot_observed_at":"2026-08-13T05:21:01.334976Z","submitted_at":"2026-06-09T07:18:01Z","title":"Act on What You See: Unlocking Safe Social Navigation in Vision-Language-Action Models","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-27T12:52:38.764427Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2606.10495"},"observation_digest":"sha256:b37b7847d7cea139de38f17c2af299931e6cefff5af21b1dc03b564949513e89","observation_id":"3df5c082-6422-4754-994e-898505c50dd2","resolution":{"observed_at":"2026-07-03T06:07:41.692074Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":"2501.09024","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-07-04T21:10:08.552391Z","title":"Social-llava: Enhancing robot navigation through human-language reasoning in social spaces","venue":null,"work_id":"8821b021-c012-4b4b-a6f3-a9004e6a5b56","year":2024},"citing_paper":{"arxiv_id":"2606.26047","last_updated":"2026-06-24T17:26:17Z","snapshot_observed_at":"2026-08-10T17:39:48.018951Z","submitted_at":"2026-06-24T17:26:17Z","title":"Learning Robot Visual Navigation in Crowds via Intention-Aware Scene Representations","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-25T19:06:41.116675Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2606.26047"},"observation_digest":"sha256:ffdc0d63e7981b6caae8e7d95400adbac9d5176477513f010c79fcdd8f65a572","observation_id":"7e28c39f-1b1e-4f96-b5ad-e7672a97f5c4","resolution":{"observed_at":"2026-07-04T21:10:08.555837Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.09024","snapshot_observed_at":"2026-07-14T07:49:44.998466Z","title":"Payandeh, D","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.10991","last_updated":"2026-07-13T01:27:27Z","snapshot_observed_at":"2026-08-14T11:57:19.957207Z","submitted_at":"2026-07-13T01:27:27Z","title":"Think When It Matters: Conditional VLM Reasoning for Social Navigation with RL Policies","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-07-14T07:49:44.998466Z"},"links":{"cited_paper":"/paper/2501.09024","citing_paper":"/paper/2607.10991"},"observation_digest":"sha256:a4e955f0272b25992111249bc9eeabaee8fbb7e536bd387913eef23f5face76a","observation_id":"e2c62eb0-ee22-4109-8aba-be77a9913e6f","resolution":{"observed_at":"2026-07-14T07:49:44.998466Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2501.09024/citation-record","integrity":"/paper/2501.09024/integrity","json":"/paper/2501.09024/citation-record.json","paper":"/paper/2501.09024"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.185702Z","title":"Conflict avoidance in social navigation—a survey,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.185702Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:c23eb182ad69caaf8cdf732fc6831ade206aba61c9d12775ff1855df70eecf0b","observation_id":"a4726c3d-475b-4731-ba4f-605a126194e7","resolution":{"observed_at":"2026-08-10T23:01:08.185702Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.771970Z","title":"Principles and guidelines for evaluating social robot navigation algorithms,","venue":null,"work_id":"ba579319-7192-427c-b560-081e51936842","year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.194466Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:76b6b1a902e46e393bebdd6d83a927e2d66fac2175f47ac206de507c91a7fe57","observation_id":"eb20d4ee-6970-4a1f-9df9-9ee679e18f31","resolution":{"observed_at":"2026-08-10T23:01:09.778096Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.202669Z","title":"Core challenges of social robot navigation: A survey,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.202669Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:7ebc3e24181086bafb1e8f35b0a03cb5e5d1403eba830f36f8f6441d6555251a","observation_id":"ebcdaa8a-ade5-4efd-9ca4-9d6ee465bea2","resolution":{"observed_at":"2026-08-10T23:01:08.202669Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.209223Z","title":"Reciprocal n-body collision avoidance,","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.209223Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:f726b2d8bef98d80b0954bf47ffc15eee7ac68a1f7e5fb6e4017068f86ba85ef","observation_id":"24152565-5ae5-4dc9-88c3-f6c0a6320a42","resolution":{"observed_at":"2026-08-10T23:01:08.209223Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.217205Z","title":"Social force model for pedestrian dynam- ics,","venue":null,"work_id":null,"year":1995},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.217205Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:ef2f1fd49c77c27ce5ff82b40760e38e5f443114a8e20a8ff03d36c2971f2e85","observation_id":"48005019-b8f4-44d5-891d-8d9b0b8ce3f5","resolution":{"observed_at":"2026-08-10T23:01:08.217205Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.712061Z","title":"Social-aware robot navigation in urban environments,","venue":null,"work_id":"c701b3c0-6965-4c79-8277-a66f15e98eb6","year":2013},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.224865Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:58b2ec6772f090637d057c29a50a855811a0cc75bf8e66369eb58ba911979983","observation_id":"744bd63e-8133-4605-8ce8-5a075350ac3b","resolution":{"observed_at":"2026-08-10T23:01:09.718716Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.690912Z","title":"Socially compliant navigation dataset (scand): A large-scale dataset of demonstrations for social navigation,","venue":null,"work_id":"7eec7f82-ff3d-4f57-8bc3-8b7288fea7d9","year":2022},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.232397Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:570dab1049df6325a4b85efa54f5e25e8ef82652ded7ed9b35e42e1ba591e7d0","observation_id":"3ca4a8b0-44b3-4bb9-82f6-30c78c4c5ecb","resolution":{"observed_at":"2026-08-10T23:01:09.697360Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.668238Z","title":"Toward human-like social robot navigation: A large-scale, multi- modal, social human navigation dataset,","venue":null,"work_id":"2343723b-6349-4bf0-b091-f87b5163b5c2","year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.238569Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:0246a685b99c0466c61ac5230f6c9a32166ce6495bf965a26c558ebe2819eeeb","observation_id":"fa346a68-712e-4d0f-b1c0-4455c872309e","resolution":{"observed_at":"2026-08-10T23:01:09.675214Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.245238Z","title":"Sacson: Scalable autonomous control for social navigation,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.245238Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:660db3da3a0400ed3d35f1a92ed3787d1e1edf4dfb3ebba9568f1ae01b3f09d5","observation_id":"512e13d6-4e21-45e5-b633-fa532e986847","resolution":{"observed_at":"2026-08-10T23:01:08.245238Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.251225Z","title":"Visual instruction tuning,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.251225Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:7905331f6b46ba62b2fe8bc8eaf2e00ad036a856fd9eb8826292f299eefa5370","observation_id":"3bda7eea-b61a-4a5d-9de4-b1b6fdeda2fb","resolution":{"observed_at":"2026-08-10T23:01:08.251225Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.610452Z","title":"Vanp: Learning where to see for navigation with self-supervised vision-action pre-training,","venue":null,"work_id":"ab0770b7-8628-44ff-b538-cfb6b0ca2f1d","year":null},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.265221Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:a92fe27eb801783b37bc78ab1f2bbcdb6cbea70b8c7d6f2d1258211e409b628b","observation_id":"239329e5-3b5a-48d3-a97a-5918e6f680c4","resolution":{"observed_at":"2026-08-10T23:01:09.617980Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2204.05443","last_updated":"2022-04-30T02:50:01Z","snapshot_observed_at":"2026-08-13T16:04:05.315060Z","submitted_at":"2022-04-11T23:38:04Z","title":"A Protocol for Validating Social Navigation Policies","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2204.05443","snapshot_observed_at":"2026-08-10T23:01:08.281900Z","title":"A protocol for validating social navigation policies,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.281900Z"},"links":{"cited_paper":"/paper/2204.05443","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:73d93e3f27669531f9ba132bf722f718721d72a9174fb2122e6c525aa93ec939","observation_id":"3c95e634-110c-4e0c-aaa5-4f498170cd44","resolution":{"observed_at":"2026-08-10T23:01:08.281900Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.580060Z","title":"Robot companion: A social- force based approach with human awareness-navigation in crowded environments,","venue":null,"work_id":"6fdf97d0-3b3d-486a-a094-6c148d860614","year":2013},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.288918Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:95fb7aa3305ebf700f3641b65b71b31a3a258a7a819632661880e182e8ac2ff4","observation_id":"f5b59245-3717-4368-a346-563d42bf765b","resolution":{"observed_at":"2026-08-10T23:01:09.587588Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.295448Z","title":"Human-robot proxemics: physical and psychological distancing in human-robot interaction,","venue":null,"work_id":null,"year":2011},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.295448Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:2d1320d0e0251d349f1e80bd55af4dcfad0cf3cf6faf2af32314190f3697e9ff","observation_id":"d5d20c7f-5a56-4cd4-9412-055317b28543","resolution":{"observed_at":"2026-08-10T23:01:08.295448Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.547152Z","title":"Robot navigation in large-scale social maps: An action recognition approach,","venue":null,"work_id":"3496c873-f457-4577-afa3-b5a268aff264","year":2016},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.302398Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:08769131b4988f7cbd40ed9db16c7d16b3e002490935caaeff2eb8f398253230","observation_id":"9ddd1b75-2427-4f7f-9a78-d69e230d1e9e","resolution":{"observed_at":"2026-08-10T23:01:09.554168Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.528289Z","title":"Exploring reflective limitation of behavior cloning in autonomous vehicles,","venue":null,"work_id":"ac71378c-25d7-4f50-8925-3ec4bd15156f","year":2021},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.310808Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:de773904ffcf90d93a4afa8aeda67383bd0c27325354025ea30d350ffbc64ff6","observation_id":"b7f5e78c-8c89-4f36-b5ec-9cf001a0a247","resolution":{"observed_at":"2026-08-10T23:01:09.534511Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.508429Z","title":"Targeted learning: A hybrid approach to social robot navigation,","venue":null,"work_id":"88b35cd3-b57f-4fa6-bafc-e040a839b5f4","year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.319241Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:20e4c7e6c7ba126f7dc8c7c710f4a68b4e0b6661cec058306fd82cee6e778d79","observation_id":"95318e53-74da-4df5-9a91-8089a20fe3a0","resolution":{"observed_at":"2026-08-10T23:01:09.513928Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.489299Z","title":"Data-driven hri: Learn- ing social behaviors by example from human–human interaction,","venue":null,"work_id":"23a8a97b-eb91-42b1-ad86-736d09b3b23b","year":2016},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.330453Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:9494676dd853e6b151740bca9fa15ed4071a8c8df86c9e7007722f9de3d5c48f","observation_id":"ec5d2524-61e3-4aae-889d-c4021084716e","resolution":{"observed_at":"2026-08-10T23:01:09.495669Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.338366Z","title":"Appld: Adaptive planner parameter learning from demonstration,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.338366Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:8d15e649d39e7e56ad8c3e3fe67bb9cb5a61039e13d29bf234c599859597ac26","observation_id":"f400fbb3-ff21-40b0-b9ea-2e3296ac9358","resolution":{"observed_at":"2026-08-10T23:01:08.338366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.345895Z","title":"Learning model pre- dictive controllers with real-time attention for real-world navigation,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.345895Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:decab05bf268219e15904d6b8d0a81d1c572aaeac9caba37ff1612107dbd1804","observation_id":"f46ee27b-a1d3-47b7-8488-59fea0054375","resolution":{"observed_at":"2026-08-10T23:01:08.345895Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.426891Z","title":"Deep representation learning: Fundamentals, technologies, applications, and open challenges,","venue":null,"work_id":"f09fca12-a62c-402f-8356-5ee13c4a4e04","year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.352873Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:86f615f5ee13df0f9a72609a15be5380037c4344b1bce2a47d4c487e916c49af","observation_id":"a9a4bc09-8475-419c-ad0d-6f0c6f65ae05","resolution":{"observed_at":"2026-08-10T23:01:09.433151Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.06500","last_updated":"2023-06-15T08:00:18Z","snapshot_observed_at":"2026-08-13T18:58:34.541884Z","submitted_at":"2023-05-11T00:38:10Z","title":"InstructBLIP: Towards General-purpose Vision-Language Models with Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.06500","snapshot_observed_at":"2026-08-10T23:01:08.360509Z","title":"Instructblip: Towards general-purpose vision-language models with instruction tuning,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.360509Z"},"links":{"cited_paper":"/paper/2305.06500","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:1d0b59784346e3249642edfdaf71d55e147c9bb4956ccb90f9c8e6b543b931e9","observation_id":"ea6c641d-eed7-424a-92cf-079189d58809","resolution":{"observed_at":"2026-08-10T23:01:08.360509Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.06281","last_updated":"2024-08-20T03:56:03Z","snapshot_observed_at":"2026-07-06T15:53:19.485466Z","submitted_at":"2023-07-12T16:23:09Z","title":"MMBench: Is Your Multi-modal Model an All-around Player?","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.06281","snapshot_observed_at":"2026-08-10T23:01:08.368963Z","title":"Mmbench: Is your multi-modal model an all-around player?","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.368963Z"},"links":{"cited_paper":"/paper/2307.06281","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:8356c2daa324328ee068d280c2174880651e4a750a54fd18b8b3b7a19fa2c764","observation_id":"e5c43d32-e685-48be-bc9b-6f1404200c6a","resolution":{"observed_at":"2026-08-10T23:01:08.368963Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.04087","last_updated":"2023-12-28T16:01:36Z","snapshot_observed_at":"2026-08-13T10:59:39.237521Z","submitted_at":"2023-07-09T03:25:14Z","title":"SVIT: Scaling up Visual Instruction Tuning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.04087","snapshot_observed_at":"2026-08-10T23:01:08.375491Z","title":"Svit: Scaling up visual instruction tuning,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.375491Z"},"links":{"cited_paper":"/paper/2307.04087","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:27803cb6aff08234a15505a3638eecf988feeb0bd020cdd341e901150a9a621d","observation_id":"7bf96cce-6eb8-43fa-9f5c-324e5272c524","resolution":{"observed_at":"2026-08-10T23:01:08.375491Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.16602","last_updated":"2023-12-27T14:54:37Z","snapshot_observed_at":"2026-08-13T15:38:15.535882Z","submitted_at":"2023-12-27T14:54:37Z","title":"Visual Instruction Tuning towards General-Purpose Multimodal Model: A Survey","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.16602","snapshot_observed_at":"2026-08-10T23:01:08.385867Z","title":"Visual instruction tuning towards general-purpose multimodal model: A survey,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.385867Z"},"links":{"cited_paper":"/paper/2312.16602","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:47f6960d96465c48bf314ef39f3ecca876e8d9e7f3ca7dfa05c9d2ac87b79dc4","observation_id":"76c474ea-bf81-4d73-8188-d4bcdcaed8fa","resolution":{"observed_at":"2026-08-10T23:01:08.385867Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1405.0312","last_updated":"2015-02-21T01:48:49Z","snapshot_observed_at":"2026-07-06T03:42:37.507458Z","submitted_at":"2014-05-01T21:43:32Z","title":"Microsoft COCO: Common Objects in Context","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1405.0312","snapshot_observed_at":"2026-08-10T23:01:08.398839Z","title":"Microsoft coco: Common objects in context,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.398839Z"},"links":{"cited_paper":"/paper/1405.0312","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:5c3803e76261d000021fbf6d7c42022a14370341100b3e88a075c3806a698897","observation_id":"903c2df0-f736-44c0-a553-a9c8fb94cc6b","resolution":{"observed_at":"2026-08-10T23:01:08.398839Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.403755Z","title":"Lamm: Language-assisted multi-modal instruction-tuning dataset, framework, and benchmark,","venue":null,"work_id":"f818dd6b-5d80-48be-a482-7fafb68cde2c","year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.408998Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:69e078b6d658c69444b55fe75303fc5198ea17bee4b35d1bf8c7e09799dd0e23","observation_id":"10943af6-a93e-492c-a537-ab22892d774d","resolution":{"observed_at":"2026-08-10T23:01:09.412685Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11684","last_updated":"2024-06-17T07:55:59Z","snapshot_observed_at":"2026-07-06T17:31:53.996417Z","submitted_at":"2024-02-18T19:26:49Z","title":"ALLaVA: Harnessing GPT4V-Synthesized Data for Lite Vision-Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.11684","snapshot_observed_at":"2026-08-10T23:01:08.417685Z","title":"Allava: Harnessing gpt4v-synthesized data for lite vision-language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.417685Z"},"links":{"cited_paper":"/paper/2402.11684","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:3f1614952d4b39cdb1375ef57fe2b2289d3a39725bdf02322716c805b0b4d510","observation_id":"c338da44-49c9-4708-b092-ea59116be6b3","resolution":{"observed_at":"2026-08-10T23:01:08.417685Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.14150","last_updated":"2025-01-16T10:57:44Z","snapshot_observed_at":"2026-08-13T04:55:49.773157Z","submitted_at":"2023-12-21T18:59:12Z","title":"DriveLM: Driving with Graph Visual Question Answering","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.14150","snapshot_observed_at":"2026-08-10T23:01:08.423570Z","title":"Drivelm: Driving with graph visual question answering,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.423570Z"},"links":{"cited_paper":"/paper/2312.14150","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:d1713c9122b2a7c7fa406e6c0626f4a41e55b9c5c4d20cdbdcd6adfefaf9d77a","observation_id":"c2420a18-5892-49ce-9903-002b4a4e042d","resolution":{"observed_at":"2026-08-10T23:01:08.423570Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.14115","last_updated":"2024-09-26T15:30:00Z","snapshot_observed_at":"2026-08-13T04:55:52.245365Z","submitted_at":"2023-12-21T18:40:34Z","title":"LingoQA: Visual Question Answering for Autonomous Driving","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.14115","snapshot_observed_at":"2026-08-10T23:01:08.429869Z","title":"Lingoqa: Video question answering for autonomous driving,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.429869Z"},"links":{"cited_paper":"/paper/2312.14115","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:d88116a1927c9eaef72360279a6e37a0d279df9bcb120834cc41416468c723de","observation_id":"fdb258d4-928d-4931-92ba-c46911302914","resolution":{"observed_at":"2026-08-10T23:01:08.429869Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.11206","last_updated":"2023-05-18T17:45:22Z","snapshot_observed_at":"2026-08-08T19:18:06.171048Z","submitted_at":"2023-05-18T17:45:22Z","title":"LIMA: Less Is More for Alignment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.11206","snapshot_observed_at":"2026-08-10T23:01:08.436680Z","title":"Lima: Less is more for alignment,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.436680Z"},"links":{"cited_paper":"/paper/2305.11206","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:64491b05714ab7472b28a00e9b374338adc2140b10b48d8a73c2ae667172a4bc","observation_id":"da149138-46bd-40ef-9627-dddc905a9524","resolution":{"observed_at":"2026-08-10T23:01:08.436680Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.08701","last_updated":"2024-02-13T18:37:25Z","snapshot_observed_at":"2026-08-13T10:54:00.274565Z","submitted_at":"2023-07-17T17:59:40Z","title":"AlpaGasus: Training A Better Alpaca with Fewer Data","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.08701","snapshot_observed_at":"2026-08-10T23:01:08.443092Z","title":"Alpagasus: Training a better alpaca with fewer data,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.443092Z"},"links":{"cited_paper":"/paper/2307.08701","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:eb745d6f31c626ff6315a74b5b6ba05a427e70c4f5bb8777b8151431ae4bf0d3","observation_id":"8195f48a-3b06-46ea-8031-e6e4f5ad1915","resolution":{"observed_at":"2026-08-10T23:01:08.443092Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.06290","last_updated":"2024-07-26T18:09:11Z","snapshot_observed_at":"2026-08-13T10:57:21.545059Z","submitted_at":"2023-07-12T16:37:31Z","title":"Instruction Mining: Instruction Data Selection for Tuning Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.06290","snapshot_observed_at":"2026-08-10T23:01:08.449490Z","title":"Instruction mining: Instruction data selection for tuning large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.449490Z"},"links":{"cited_paper":"/paper/2307.06290","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:aedad027f0df78628d25aa87ec47442dd311a7603776f0bebab9079d10139fae","observation_id":"c6714b32-3947-4c83-b032-bb97638279d4","resolution":{"observed_at":"2026-08-10T23:01:08.449490Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.12067","last_updated":"2026-04-12T14:13:44Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-23T11:27:30Z","title":"MM-LIMA: Less Is More for Alignment in Multi-Modal Datasets","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.12067","snapshot_observed_at":"2026-08-10T23:01:08.455984Z","title":"Instructiongpt-4: A 200-instruction paradigm for fine-tuning minigpt-4,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.455984Z"},"links":{"cited_paper":"/paper/2308.12067","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:52d544dd94442a16f665d04a78d1e2d2547dd1f4655b76efac58ea441a3608d1","observation_id":"7aa22314-150e-4637-b044-0071fea85742","resolution":{"observed_at":"2026-08-10T23:01:08.455984Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2201.11903","last_updated":"2023-01-10T23:07:57Z","snapshot_observed_at":"2026-08-13T07:04:41.220509Z","submitted_at":"2022-01-28T02:33:07Z","title":"Chain-of-Thought Prompting Elicits Reasoning in Large Language Models","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2201.11903","snapshot_observed_at":"2026-08-10T23:01:08.462569Z","title":"Chain-of-thought prompting elicits reasoning in large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.462569Z"},"links":{"cited_paper":"/paper/2201.11903","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:dde03c2cd8eb066ed19bc7b13fece9a2487496c3d06ce00679230f227b727a5c","observation_id":"b8e7a55b-dda3-4166-9f4d-a94ba55b2935","resolution":{"observed_at":"2026-08-10T23:01:08.462569Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.381510Z","title":null,"venue":null,"work_id":"afd0b266-accd-4d40-9ec9-c77493ba0db3","year":2013},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.469560Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:f9960a7844a91c976d51269a02b7267b3f3ce3bbbc003b2e865b61f54a3103dd","observation_id":"b39d7d6c-dd34-4c75-95bb-5e92cf6fcebf","resolution":{"observed_at":"2026-08-10T23:01:09.388867Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.359601Z","title":"Marr, Vision : a computational investigation into the human representation and processing of visual information","venue":null,"work_id":"a8d6df62-357b-466e-93a5-7ee3a44a007e","year":2010},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.477208Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:b8c346364c63e9a8ae9acc005f9329efe7456af994ad3e0570de317b99832584","observation_id":"3ae28efb-4ca8-4f9c-80b9-6f14562b864d","resolution":{"observed_at":"2026-08-10T23:01:09.367646Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.328880Z","title":"Spatialvlm: Endowing vision-language models with spatial reasoning capabilities,","venue":null,"work_id":"a8ddb740-8b15-4857-bbb7-42020398368d","year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.484445Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:b80b8687299dfdd2d9b1dc89554ad7ae82cecc7ee857afb1a2328f385d6b98f9","observation_id":"3d43a7d1-e8fb-4156-b7dc-37296458d72b","resolution":{"observed_at":"2026-08-10T23:01:09.335847Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-10T23:01:08.491895Z","title":"Gpt-4 technical report,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.491895Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:616b9d71aa6bae024d8c7c763607a1dea4cda178b63169382e21de179d76fe0a","observation_id":"4d71acc0-5e54-4f84-968d-cf162780eb78","resolution":{"observed_at":"2026-08-10T23:01:08.491895Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-10T23:01:08.499496Z","title":"Gemini: a family of highly capable multimodal models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.499496Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:5311825d9c40141ae2747f165ab573923d877228671ec7bde8d7f72b4a0f827e","observation_id":"c7cfa1ed-d3bf-44fa-918e-eab9aacb3e6b","resolution":{"observed_at":"2026-08-10T23:01:08.499496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.508876Z","title":"Lora: Low-rank adaptation of large language models,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.508876Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:ee950b265f13683325dfc208b598be48a8577d8a2dcdb04395c947d960ded09f","observation_id":"846fb2f2-e1a8-436d-9b55-77606b0ea995","resolution":{"observed_at":"2026-08-10T23:01:08.508876Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01584","last_updated":"2024-10-15T01:16:20Z","snapshot_observed_at":"2026-08-14T09:35:24.079025Z","submitted_at":"2024-06-03T17:59:06Z","title":"SpatialRGPT: Grounded Spatial Reasoning in Vision Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01584","snapshot_observed_at":"2026-08-10T23:01:08.524496Z","title":"Spatialrgpt: Grounded spatial reasoning in vision language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.524496Z"},"links":{"cited_paper":"/paper/2406.01584","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:6b0f6aae7bcf1fd1229afdb90854eb30bd1d8bf1ddfd1d27951364460b2cf564","observation_id":"517a0395-b94f-455c-ac51-15f0b0de6eec","resolution":{"observed_at":"2026-08-10T23:01:08.524496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:09.290229Z","title":"Structured spatial reasoning with open vocabulary object detectors,","venue":null,"work_id":"5af92fc1-0551-4e0d-8c08-f2d28578585a","year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.532255Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:44106234cbd69ae50fd00c5ddde369e276a430fe10b75093467015418e744462","observation_id":"3b20cc39-d3fd-4efa-8abb-9158739bcd8f","resolution":{"observed_at":"2026-08-10T23:01:09.299091Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T23:01:08.538521Z","title":"Vision- and-language navigation: A survey of tasks, methods, and future directions,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.538521Z"},"links":{"citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:35b5d767082b80f32462cbe6dc2708f89112631d8342d83a79e56fedaf34dcb9","observation_id":"0468ba8d-c02f-4785-8ed1-3d6c64e7b565","resolution":{"observed_at":"2026-08-10T23:01:08.538521Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01415","last_updated":"2023-12-05T05:26:29Z","snapshot_observed_at":"2026-08-14T07:41:23.792152Z","submitted_at":"2023-10-02T17:59:57Z","title":"GPT-Driver: Learning to Drive with GPT","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.01415","snapshot_observed_at":"2026-08-10T23:01:08.544936Z","title":"Gpt-driver: Learning to drive with gpt,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.544936Z"},"links":{"cited_paper":"/paper/2310.01415","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:93cbd859962fc0fc7e7e991a7b446b9502733839e81b61ec15cb7578956dbea6","observation_id":"effb321a-65db-4508-b35d-580552d7f6a1","resolution":{"observed_at":"2026-08-10T23:01:08.544936Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11498","last_updated":"2025-03-30T03:37:48Z","snapshot_observed_at":"2026-08-13T04:16:48.542891Z","submitted_at":"2024-02-18T08:05:54Z","title":"Verifiably Following Complex Robot Instructions with Foundation Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.11498","snapshot_observed_at":"2026-08-10T23:01:08.550920Z","title":"Verifiably following complex robot instructions with foundation models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.550920Z"},"links":{"cited_paper":"/paper/2402.11498","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:d00d43254e74c0a3cda64c0f1c931fc796535f5272eeadd926ea24ddd72899ce","observation_id":"e6950e09-bbe7-4121-89a6-a8dc28bcf122","resolution":{"observed_at":"2026-08-10T23:01:08.550920Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.07035","last_updated":"2024-12-29T23:16:37Z","snapshot_observed_at":"2026-08-12T23:25:39.126520Z","submitted_at":"2024-07-09T16:53:36Z","title":"Vision-and-Language Navigation Today and Tomorrow: A Survey in the Era of Foundation Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.07035","snapshot_observed_at":"2026-08-10T23:01:08.557671Z","title":"Vision-and-language navigation today and tomorrow: A survey in the era of foundation models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.557671Z"},"links":{"cited_paper":"/paper/2407.07035","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:f3a205caa2a25cc3c9dfc404b20020b7194aa10cb7dafed425304c4c88c461a2","observation_id":"6277bf27-adcc-4719-acc0-937f968bba1c","resolution":{"observed_at":"2026-08-10T23:01:08.557671Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.09685","last_updated":"2021-10-16T18:40:34Z","snapshot_observed_at":"2026-08-11T08:20:29.798517Z","submitted_at":"2021-06-17T17:37:18Z","title":"LoRA: Low-Rank Adaptation of Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.09685","snapshot_observed_at":"2026-08-10T23:01:08.517015Z","title":"Available: https://arxiv.org/abs/2106.09685","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.517015Z"},"links":{"cited_paper":"/paper/2106.09685","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:3060042ac502362436bccb88de4ea37c519552812fb1d889bde401e7c7b8ff7f","observation_id":"f69939b2-5c5c-4889-bec5-beba66f11609","resolution":{"observed_at":"2026-08-10T23:01:08.517015Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08485","last_updated":"2023-12-11T17:46:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-04-17T17:59:25Z","title":"Visual Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08485","snapshot_observed_at":"2026-08-10T23:01:08.256568Z","title":"Available: https://arxiv.org/abs/2304.08485","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.256568Z"},"links":{"cited_paper":"/paper/2304.08485","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:05740c1a88eb9f3897d90a371ec278c1ce00ea0bdd7b161bf6e81b1173b9f8cf","observation_id":"4b91085b-47f2-4dca-95cf-3dd267d81efa","resolution":{"observed_at":"2026-08-10T23:01:08.256568Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.08109","last_updated":"2024-09-04T20:54:13Z","snapshot_observed_at":"2026-08-13T00:55:40.524750Z","submitted_at":"2024-03-12T22:33:08Z","title":"VANP: Learning Where to See for Navigation with Self-Supervised Vision-Action Pre-Training","version":3},"cited_work":{"arxiv_id":"2403.08109","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.08109","snapshot_observed_at":"2026-08-10T23:01:09.239877Z","title":"VANP: Learning Where to See for Navigation with Self-Supervised Vision-Action Pre-Training","venue":"cs.RO","work_id":"f127659b-0781-4ed5-a7ea-32fd32803d73","year":2024},"citing_paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-10T23:01:08.272974Z"},"links":{"cited_paper":"/paper/2403.08109","citing_paper":"/paper/2501.09024"},"observation_digest":"sha256:d95ee6e72e9eb07c109c4e459e5bd7ad20750245953c8f7b228f6008438a32c5","observation_id":"eb2bd44e-0beb-4e5d-aad2-578d1a875e12","resolution":{"observed_at":"2026-08-10T23:01:09.249181Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-14T06:32:32.682623+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2501.09024","last_updated":"2024-12-30T23:59:30Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-14T19:14:47.886608Z","submitted_at":"2024-12-30T23:59:30Z","title":"Social-LLaVA: Enhancing Robot Navigation through Human-Language Reasoning in Social Spaces"},"reference_resolution":{"displayed":50,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":34,"verified_exact":0,"verified_fuzzy":15},"total_outbound_references":50},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-14T06:32:32.682623+00:00","source":"crossref"},{"observed_at":"2026-08-14T06:32:18.44784+00:00","source":"retraction_watch"}],"thesis":"As of 15 August 2026, this Paper Citation Record lists 50 of 50 outbound references and 10 inbound Pith citation observations for arXiv:2501.09024."}