{"as_of":"2026-08-05T22:33:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d5c39cf6970cb3f266817c65d4e5a883800682d774caace1fbb3b269ffb0bfba","coverage":[{"denominator":38,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":38,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-27T19:54:58.135809Z","state":"measured"},{"denominator":39,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":39,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T00:48:36.239108Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-05T00:48:38.306341Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"cited_work":{"arxiv_id":"2606.08029","doi":null,"metadata_source":"pith","pith_arxiv_id":"2606.08029","snapshot_observed_at":"2026-08-05T00:48:38.306341Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","venue":"cs.RO","work_id":"e004c9dc-0c4a-41de-9f2e-8310ea676825","year":2026},"citing_paper":{"arxiv_id":"2608.00527","last_updated":"2026-08-01T08:40:04Z","snapshot_observed_at":"2026-08-05T22:16:25.724664Z","submitted_at":"2026-08-01T08:40:04Z","title":"SSTG-Nav: Metric-Grounded Spatial-Semantic Topological Graphs for Reusable Object Navigation","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-05T00:48:36.239108Z"},"links":{"cited_paper":"/paper/2606.08029","citing_paper":"/paper/2608.00527"},"observation_digest":"sha256:9ed47829c36d6b15b160ae8b6c70e900f451ce9ebc874599694de387ca8fead9","observation_id":"aced2130-fe82-4e97-a098-b757d133a386","resolution":{"observed_at":"2026-08-05T00:48:38.384115Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2606.08029/citation-record","integrity":"/paper/2606.08029/integrity","json":"/paper/2606.08029/citation-record.json","paper":"/paper/2606.08029"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2006.13171","last_updated":"2020-08-30T04:28:13Z","snapshot_observed_at":"2026-07-06T09:32:02.150002Z","submitted_at":"2020-06-23T17:18:54Z","title":"ObjectNav Revisited: On Evaluation of Embodied Agents Navigating to Objects","version":2},"cited_work":{"arxiv_id":"2006.13171","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2006.13171","snapshot_observed_at":"2026-07-04T19:20:06.888892Z","title":"Objectnav revisited: On evaluation of embodied agents navigating to objects","venue":null,"work_id":"2559f113-5e2a-419e-8aba-d15f8f50737c","year":2006},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2006.13171","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:10c1636f4eb952fb938daa61b6b6e1907bccf25e824bbf9e43de96f2f066f022","observation_id":"7dc7aaed-1eb4-4ba8-9886-c30b4e58e229","resolution":{"observed_at":"2026-07-02T21:07:24.185885Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1807.06757","last_updated":"2018-07-18T03:28:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2018-07-18T03:28:02Z","title":"On Evaluation of Embodied Navigation Agents","version":1},"cited_work":{"arxiv_id":"1807.06757","doi":null,"metadata_source":"pith","pith_arxiv_id":"1807.06757","snapshot_observed_at":"2026-07-04T09:09:43.494441Z","title":"On Evaluation of Embodied Navigation Agents","venue":"cs.AI","work_id":"3b074aa9-2ff9-4ad6-8796-6a25689ecfd3","year":2018},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/1807.06757","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:07e7a36d14b6c38c92f0c05885ca85df5a58772b8f48e940942339dc703b1532","observation_id":"c1ba877d-6a3d-4dcf-a56b-67e62c27f300","resolution":{"observed_at":"2026-07-02T21:07:24.190934Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:eb197ab4e3a85ef29457812c36c20223adcd0c04638a9af6c1468e37affe2fb9","observation_id":"3c39a9a6-4827-468a-bcc8-480fde3b1a22","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.13166","last_updated":"2023-07-06T06:25:33Z","snapshot_observed_at":"2026-07-06T14:46:15.598295Z","submitted_at":"2023-01-30T18:37:32Z","title":"ESC: Exploration with Soft Commonsense Constraints for Zero-shot Object Navigation","version":3},"cited_work":{"arxiv_id":"2301.13166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2301.13166","snapshot_observed_at":"2026-07-02T21:07:24.235438Z","title":"ESC: Exploration with Soft Commonsense Constraints for Zero-shot Object Navigation","venue":null,"work_id":"e886a807-56a7-4c46-8db3-18e379dfd9b2","year":2023},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2301.13166","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:ddb4b4b1a774f3ff4d1165ab9ac851c413cd6c595b6d800562027336e4cd4a7d","observation_id":"500e2cc0-9884-4dc0-a31a-193a7b3d5f1f","resolution":{"observed_at":"2026-07-02T21:07:24.237051Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":"Yokoyama, S","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:d5c78aab6e6b2d91ebc70d79292798e17e6bcb94cf227a290517e0bb4868d384","observation_id":"98a47746-041f-4d0b-b362-9f7b7a3b5a9a","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":"Kuang, H","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:d0674cbf1daeafe5b8cb06f0e26651bbc169ffbc7b36c477ce0392a9a3bbe873","observation_id":"adc9a3cb-8043-4f40-8fcc-4f9b58cad883","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:8af66351b8d709b3db193778660e6afb3379e849a84eb23fddc92e5c9564e90c","observation_id":"f14148c9-148b-494f-943e-b2e8d4f6d939","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.10439","last_updated":"2025-08-28T04:36:42Z","snapshot_observed_at":"2026-07-06T20:06:48.183096Z","submitted_at":"2024-12-11T09:50:35Z","title":"CogNav: Cognitive Process Modeling for Object Goal Navigation with LLMs","version":3},"cited_work":{"arxiv_id":"2412.10439","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.10439","snapshot_observed_at":"2026-07-04T19:20:06.905778Z","title":"Cognav: Cognitive process modeling for object goal navigation with llms","venue":null,"work_id":"14cb2472-93e5-43db-9226-eac825329385","year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2412.10439","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:0fde5c3601d07dcf598ea0deb231c778ab37059842f3f9d139953696046a3a39","observation_id":"0fafb9bc-b417-4a43-b76b-0450818abd2a","resolution":{"observed_at":"2026-07-02T21:07:24.245311Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.06729","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T21:07:24.238155Z","title":"STRIVE: Structured representation integrating VLM reasoning for efficient object navigation","venue":null,"work_id":"a0bbeb8d-d592-41aa-afdd-6701bea985fe","year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:f6a13110648fabf236547b9a9b9c9c6428dd44a7e336d5ad3ed1ba520138ecb0","observation_id":"78d49a6f-64a9-4a46-9ac1-296fb12e438b","resolution":{"observed_at":"2026-07-02T21:07:24.239522Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.17385","last_updated":"2025-04-17T17:59:27Z","snapshot_observed_at":"2026-07-06T19:38:04.948333Z","submitted_at":"2024-10-22T19:39:15Z","title":"Do Vision-Language Models Represent Space and How? Evaluating Spatial Frame of Reference Under Ambiguities","version":2},"cited_work":{"arxiv_id":"2410.17385","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2410.17385","snapshot_observed_at":"2026-07-02T21:07:24.216562Z","title":"Do vision-language models represent space and how? eval- uating spatial frame of reference under ambiguities","venue":null,"work_id":"f53c4559-e7da-40a9-935f-1e53b931b6be","year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2410.17385","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:8dcd8b4588bd38fb395cea9e9aed5b59adc389240beabb442a2ceeecb28c25c9","observation_id":"91c09a56-d392-490c-a84f-47e490910923","resolution":{"observed_at":"2026-07-02T21:07:24.218032Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:ce63906261b37318066f90b47815e1d2a0efa408c61107f2e72f184342e5e32c","observation_id":"945d4ca5-e165-4d44-a1c6-d39c50e6619f","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08669","last_updated":"2023-11-06T18:44:33Z","snapshot_observed_at":"2026-08-03T03:35:03.879513Z","submitted_at":"2023-10-12T19:01:06Z","title":"Multimodal Large Language Model for Visual Navigation","version":2},"cited_work":{"arxiv_id":"2310.08669","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2310.08669","snapshot_observed_at":"2026-07-02T21:07:24.219142Z","title":"Multimodal large lan- guage model for visual navigation","venue":null,"work_id":"42e9e228-026e-48df-b36c-2a837e87457e","year":2023},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2310.08669","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:ddc2bcf83cbebc2875564c479d23ba9efd03e04da2d97a2295a69da75a658fae","observation_id":"9cba033b-5cdb-45ae-9d23-71a860a3d69e","resolution":{"observed_at":"2026-07-02T21:07:24.220811Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.10154","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T21:07:24.224570Z","title":"arXiv preprint arXiv:2510.10154 , year=","venue":null,"work_id":"97449241-9641-4f0e-b859-2145e1ec88a9","year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:68cd00edd85916e2af1074bb22325c3db0a7f60de4b4b31ecae1224884521bbf","observation_id":"ff392ac0-ff1a-4e34-ab98-9f6b02e981b3","resolution":{"observed_at":"2026-07-02T21:07:24.226174Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2602.09972","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T21:07:24.227141Z","title":"Hydra-nav: Object navigation via adaptive dual-process reasoning","venue":null,"work_id":"3e84bd18-0c30-4737-9357-0f11751da02e","year":2026},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:5b1bcc652cb8bf34a910fea341d278e8831ed4d7cc8e9be09c71a1c86cf222de","observation_id":"90a81f87-9e3f-45b6-b8df-f44eb6b56ff4","resolution":{"observed_at":"2026-07-02T21:07:24.228801Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.15852","last_updated":"2024-06-30T11:14:13Z","snapshot_observed_at":"2026-07-06T17:34:57.509162Z","submitted_at":"2024-02-24T16:39:16Z","title":"NaVid: Video-based VLM Plans the Next Step for Vision-and-Language Navigation","version":7},"cited_work":{"arxiv_id":"2402.15852","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.15852","snapshot_observed_at":"2026-07-10T08:47:01.956969Z","title":"NaVid: Video-based VLM Plans the Next Step for Vision-and-Language Navigation","venue":"cs.CV","work_id":"a18e1283-0ce1-4c27-bf05-99efa4f917c4","year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2402.15852","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:96af326712a59fa26b500d2173074e386e8897d13a4122cc91d3e7ac3f1d1fd1","observation_id":"cd8255f0-b2ef-4c3a-bf67-2704f1ff50fc","resolution":{"observed_at":"2026-07-02T21:07:24.212081Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.06224","last_updated":"2025-02-06T10:14:36Z","snapshot_observed_at":"2026-07-06T20:03:42.903210Z","submitted_at":"2024-12-09T05:55:55Z","title":"Uni-NaVid: A Video-based Vision-Language-Action Model for Unifying Embodied Navigation Tasks","version":2},"cited_work":{"arxiv_id":"2412.06224","doi":null,"metadata_source":"pith","pith_arxiv_id":"2412.06224","snapshot_observed_at":"2026-07-11T02:27:49.380331Z","title":"Uni-NaVid: A Video-based Vision-Language-Action Model for Unifying Embodied Navigation Tasks","venue":"cs.RO","work_id":"13512be0-3aa0-406d-8602-a0edc0286028","year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2412.06224","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:8c340bb60bd09812d94f6ddfe20e8464d94d60b466182d7f6b3df132649367c8","observation_id":"0a4b5beb-100d-47b0-82a0-60d5147a01fd","resolution":{"observed_at":"2026-07-02T21:07:24.201546Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.25687","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T21:07:24.202581Z","title":"Omninav: A unified framework for prospective exploration and visual-language navigation","venue":null,"work_id":"3dab4811-a369-4aed-9612-71d164d55fb2","year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:453cb2ddfbb5ac6f2f4f9f4afb269f38a6be6c2a62f48d492aa0ed62b1be61f1","observation_id":"a7337eef-9711-4811-ad5a-9f609699fd64","resolution":{"observed_at":"2026-07-02T21:07:24.204309Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.06914","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T21:07:24.194844Z","title":"SysNav: Multi-level systematic cooperation enables real-world, cross-embodiment object navigation","venue":null,"work_id":"e13ea000-9db7-4d06-8384-747ce7841abf","year":2026},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:7d53cfeb9555f2f10a1627adf6d02bb370e6a239ea5a9fc7cb1725d7fe344c72","observation_id":"9c2fb86b-787f-4e4d-8a8b-f2c3c826ca71","resolution":{"observed_at":"2026-07-02T21:07:24.196507Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.02695","last_updated":"2024-02-06T05:15:20Z","snapshot_observed_at":"2026-07-06T17:11:55.112404Z","submitted_at":"2024-01-05T08:05:07Z","title":"VoroNav: Voronoi-based Zero-shot Object Navigation with Large Language Model","version":2},"cited_work":{"arxiv_id":"2401.02695","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2401.02695","snapshot_observed_at":"2026-07-04T16:59:58.766242Z","title":"V oronav: V oronoi-based zero-shot object navigation with large language model","venue":null,"work_id":"bcd8597a-edc9-4ff0-aba0-2914ab9c8dd8","year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2401.02695","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:07897d600cff5781ffe49ddfdef975f7a4e5f4385e4296cbd25eaa1d1ba23e70","observation_id":"d6e900c4-c6dc-4d37-9ce8-d1ca4c9be84a","resolution":{"observed_at":"2026-07-02T21:07:24.193682Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":"Majumdar, G","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:e8208d927e5a1f09efe9aaf2ad84b519be3db9248e2120bc8998f347d658e116","observation_id":"5e516736-83c2-438a-99c8-aa47448e26f9","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:725dc832547eac737c53ef5d54f44cab26164616ccb2f5659c8b2593b6c51b5b","observation_id":"d01aab07-5d66-4748-9480-b5acd9bceaf1","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":"Zhang, Y","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:46a72060d640bf4a8b6940443157f30301d3869c70e4780f165d2fea58c44569","observation_id":"4154abf5-ec77-427e-ab2b-3d049ae7d1a7","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.16425","last_updated":"2025-03-26T07:26:43Z","snapshot_observed_at":"2026-08-03T20:10:34.950874Z","submitted_at":"2024-11-25T14:27:55Z","title":"TopV-Nav: Unlocking the Top-View Spatial Reasoning Potential of MLLM for Zero-shot Object Navigation","version":2},"cited_work":{"arxiv_id":"2411.16425","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.16425","snapshot_observed_at":"2026-07-02T21:07:24.197654Z","title":"Topv-nav: Unlocking the top-view spatial reasoning potential of mllm for zero-shot object navigation","venue":null,"work_id":"40b4e195-ffca-4f72-8e61-9273664b717c","year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2411.16425","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:dbb959485f530cf44ac06e33049a14c3ae4751d47c10491f10f5ce19ee2a78c9","observation_id":"2709a626-ffe2-4fc9-9b7a-a056b8667dda","resolution":{"observed_at":"2026-07-02T21:07:24.199096Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":"Ramrakhya, E","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:ad16bc3faf828c955feea3492dc043b2ed732c42f156837788de7541ab05564f","observation_id":"a04c361a-fd8d-459e-b541-b8a338541447","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1709.06158","last_updated":"2017-09-18T20:34:48Z","snapshot_observed_at":"2026-08-02T16:47:23.623541Z","submitted_at":"2017-09-18T20:34:48Z","title":"Matterport3D: Learning from RGB-D Data in Indoor Environments","version":1},"cited_work":{"arxiv_id":"1709.06158","doi":"10.48550/arxiv.1709.06158","metadata_source":"pith","pith_arxiv_id":"1709.06158","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Matterport3D: Learning from RGB-D Data in Indoor Environments","venue":"cs.CV","work_id":"a6675134-1bd7-4d3f-9344-d7072e7449e9","year":2017},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/1709.06158","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:1bf9eb1dfd768fc872c025e974dca56368363c46a67e48f7f4cd19d2b1e7d9e4","observation_id":"3ab65a33-ee68-4f88-81a8-e94558110663","resolution":{"observed_at":"2026-07-02T21:07:24.206844Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13724","last_updated":"2023-10-19T17:29:17Z","snapshot_observed_at":"2026-08-02T17:55:38.209254Z","submitted_at":"2023-10-19T17:29:17Z","title":"Habitat 3.0: A Co-Habitat for Humans, Avatars and Robots","version":1},"cited_work":{"arxiv_id":"2310.13724","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.13724","snapshot_observed_at":"2026-07-08T02:44:27.650595Z","title":"Habitat 3.0: A co-habitat for humans, avatars and robots","venue":"cs.HC","work_id":"95bbd779-1839-47c2-a371-d4282687070f","year":2023},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2310.13724","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:70fd337ee147f2e7d8a8c6eb05279f371fc3fcae76fd6daeadd6efe9edf56b1d","observation_id":"571d7080-4527-4e3b-a73f-5319ab912db0","resolution":{"observed_at":"2026-07-02T21:07:24.209386Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.10479","last_updated":"2025-04-19T03:47:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-14T17:59:25Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","version":3},"cited_work":{"arxiv_id":"2504.10479","doi":"10.48550/arxiv.2504.10479","metadata_source":"pith","pith_arxiv_id":"2504.10479","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","venue":"cs.CV","work_id":"fe8637aa-12bc-4434-8d36-9f57b5eebcbe","year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2504.10479","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:730d973221a3d9f7a3b249f7bfb97f3f5c0993f74bde4d2b7a76da4e6372075a","observation_id":"f81f36a1-8aa7-4ed6-9505-6f89c7f02010","resolution":{"observed_at":"2026-07-02T21:07:24.214612Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-20T07:54:09.017512+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T07:54:09.017512+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":"Yadav, S","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:a4c80c0920ffcae2e8fdea5f4746e703b642f447b091fce7d7293701c1fe051c","observation_id":"76c122f2-9dca-4b21-8c73-d08e682f30cb","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":"Yadav, R","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:45c60465afa7468c557a76d16a9edc1adab19feb46ebe7d3b0579c725f49bcb7","observation_id":"c35edf42-86e5-4a40-a649-b3a7d0d4ed77","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15223","last_updated":"2024-03-22T14:15:27Z","snapshot_observed_at":"2026-07-06T17:49:03.650707Z","submitted_at":"2024-03-22T14:15:27Z","title":"TriHelper: Zero-Shot Object Navigation with Dynamic Assistance","version":1},"cited_work":{"arxiv_id":"2403.15223","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.15223","snapshot_observed_at":"2026-07-02T21:07:24.229956Z","title":"Zhang, Q","venue":null,"work_id":"f427e78b-5558-4ce0-a8f4-7a856e1577ef","year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2403.15223","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:231b2657ae3f7494e563294426579812cae184ce03f112442b9ecfd8ae10b420","observation_id":"d315154a-040c-43d3-805a-436325ae28a7","resolution":{"observed_at":"2026-07-02T21:07:24.231416Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.06487","last_updated":"2025-05-27T07:28:27Z","snapshot_observed_at":"2026-07-06T21:38:10.885548Z","submitted_at":"2025-05-27T07:28:27Z","title":"BeliefMapNav: 3D Voxel-Based Belief Map for Zero-Shot Object Navigation","version":1},"cited_work":{"arxiv_id":"2506.06487","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.06487","snapshot_observed_at":"2026-07-02T21:07:24.237572Z","title":"arXiv preprint arXiv:2506.06487 , year=","venue":null,"work_id":"d44f41e8-c179-4f82-8494-8497398d6daf","year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2506.06487","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:44dd0eaf6c96ddebca66e16b3b2d499deaeb2a9d23b8664ca31bd03bbb3db228","observation_id":"a1af2e3e-de6a-4b90-bbd4-0a3cc35aee4a","resolution":{"observed_at":"2026-07-02T21:07:24.238925Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.02247","last_updated":"2025-07-19T03:44:28Z","snapshot_observed_at":"2026-07-06T20:46:13.349012Z","submitted_at":"2025-03-04T03:51:36Z","title":"WMNav: Integrating Vision-Language Models into World Models for Object Goal Navigation","version":5},"cited_work":{"arxiv_id":"2503.02247","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2503.02247","snapshot_observed_at":"2026-07-04T10:49:46.489053Z","title":"arXiv preprint arXiv:2503.02247 , year=","venue":null,"work_id":"52a4f8d0-1c0f-42fd-9c8d-5b403cead7ec","year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2503.02247","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:3e8b821108fb5e4d5ec35b99dd7adbc3e52266529c859a98573e34a7da7523ab","observation_id":"2efc2e8c-f48a-4db4-9f7c-4eefa1c8668b","resolution":{"observed_at":"2026-07-02T21:07:24.226035Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.13207","last_updated":"2026-06-10T09:55:01Z","snapshot_observed_at":"2026-08-05T17:53:49.817441Z","submitted_at":"2025-11-17T10:19:13Z","title":"PIGEON: VLM-Driven Object Navigation via Points of Interest Selection","version":2},"cited_work":{"arxiv_id":"2511.13207","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.13207","snapshot_observed_at":"2026-07-02T21:07:24.240655Z","title":"Pigeon: Vlm-driven object navigation via points of interest selection","venue":"cs.RO","work_id":"92d7f288-90a1-420e-a5dc-c44b597d7dd7","year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/2511.13207","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:4c239a79e3a14bd913bbfd8914c5f9578324e7a8af12142e9c263fb28be07419","observation_id":"4b7ef735-5e44-4a6e-b727-27573a156c87","resolution":{"observed_at":"2026-07-02T21:07:24.242343Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1911.00357","last_updated":"2020-01-20T04:18:58Z","snapshot_observed_at":"2026-08-03T16:21:22.610677Z","submitted_at":"2019-11-01T13:07:37Z","title":"DD-PPO: Learning Near-Perfect PointGoal Navigators from 2.5 Billion Frames","version":2},"cited_work":{"arxiv_id":"1911.00357","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"1911.00357","snapshot_observed_at":"2026-07-04T16:59:58.769476Z","title":"Decentralized distributed PPO: solving pointgoal navigation","venue":null,"work_id":"7170e982-3143-4286-a431-71793e679d9a","year":1911},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"cited_paper":"/paper/1911.00357","citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:d669b8ea601c7171fda249462596f2534f29e75f6c1dc8949e1af78300c0c255","observation_id":"ce1fde40-4fad-4453-a94f-079add95459b","resolution":{"observed_at":"2026-07-02T21:07:24.188450Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:21d1089164142aaaddb0513857699be4fce0f287af08e10cd7e78a3b348c807f","observation_id":"b343d3ba-ee2e-4c0b-ba8b-ffdf4c8e4cb0","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:4d8f9140e9623ffc444a62d6a982959345ba9cfa88206c2c3eaa01d6af438586","observation_id":"fb0ce682-3948-4d03-abe8-81509c737dfc","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:c3fb3fb537406bfa1a1294e9110499ccac2fa42c57fefe596699e7fdbc2f125d","observation_id":"e58f00ab-7473-455a-a15b-dba59dacac7d","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-06-27T19:54:58.135809Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-06-27T19:54:58.135809Z"},"links":{"citing_paper":"/paper/2606.08029"},"observation_digest":"sha256:27a7c78731ca5a8cee601baf32b382f2c45e88dd9e6e37d430d0184d0e992b99","observation_id":"f81cfbdc-e1c6-4606-90e7-7efd6afa9c49","resolution":{"observed_at":"2026-06-27T19:54:58.135809Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2606.08029","last_updated":"2026-06-06T07:45:19Z","latest_version":1,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-02T17:57:16.436409Z","submitted_at":"2026-06-06T07:45:19Z","title":"IntentNav: Learning Spatial-Visual Object Navigation from Human Demonstrations"},"reference_resolution":{"displayed":38,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":14,"verified_exact":23,"verified_fuzzy":0},"total_outbound_references":38},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 38 of 38 outbound references and 1 inbound Pith citation observation for arXiv:2606.08029."}