{"as_of":"2026-08-17T02:31:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:b30b60d372d1838d021192649dc9ffadf28c1022129ace976f19b63938fef33a","coverage":[{"denominator":130,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T19:58:00.610186Z","state":"measured"},{"denominator":101,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":101,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-16T06:30:59.297886+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-10T08:17:17.920830Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-10T08:17:36.960909Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"cited_work":{"arxiv_id":"2506.14096","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.14096","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Im- age segmentation with large language models: A survey with perspectives for intelligent transportation systems","venue":null,"work_id":"d5a6dd22-d217-40cc-8177-19716e4328a9","year":2025},"citing_paper":{"arxiv_id":"2604.15946","last_updated":"2026-04-17T11:07:36Z","snapshot_observed_at":"2026-08-05T17:01:30.244509Z","submitted_at":"2026-04-17T11:07:36Z","title":"SENSE: Stereo OpEN Vocabulary SEmantic Segmentation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-10T08:17:17.920830Z"},"links":{"cited_paper":"/paper/2506.14096","citing_paper":"/paper/2604.15946"},"observation_digest":"sha256:ee54f7d5cea875c5cca00033f536068f09f5daa6f50ac3827848f3076381269e","observation_id":"b53b7687-c19d-417c-87ff-0882f86e0693","resolution":{"observed_at":"2026-05-10T08:17:36.962388Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.14096/citation-record","integrity":"/paper/2506.14096/integrity","json":"/paper/2506.14096/citation-record.json","paper":"/paper/2506.14096"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2310.14414","last_updated":"2024-06-20T19:36:38Z","snapshot_observed_at":"2026-08-16T14:49:58.499425Z","submitted_at":"2023-10-22T21:06:10Z","title":"Vision Language Models in Autonomous Driving: A Survey and Outlook","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.14414","snapshot_observed_at":"2026-08-15T19:58:00.150325Z","title":"Vision language models in autonomous driving: A survey and outlook,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.150325Z"},"links":{"cited_paper":"/paper/2310.14414","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:ccfefbcd5b79a7a2bf43db7cc3059fefe9f4efee35c0f8ef281807fca31b95f1","observation_id":"9024ef81-34ab-4c93-9f2a-dd7f993dcae3","resolution":{"observed_at":"2026-08-15T19:58:00.150325Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2202.02703","last_updated":"2024-12-16T23:42:56Z","snapshot_observed_at":"2026-08-16T22:45:04.602281Z","submitted_at":"2022-02-06T04:18:45Z","title":"Multi-modal Sensor Fusion for Auto Driving Perception: A Survey","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2202.02703","snapshot_observed_at":"2026-08-15T19:58:00.156301Z","title":"Multi-modal sensor fusion for auto driving perception: A survey,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.156301Z"},"links":{"cited_paper":"/paper/2202.02703","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:94c4542f86d054800db03ad2429d578f5d7e0936f4f7d6b6511ad6a432b3c55f","observation_id":"f89e7236-cbbf-46b5-85f4-fcdfc39e7047","resolution":{"observed_at":"2026-08-15T19:58:00.156301Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.161516Z","title":"Deeplab: Semantic image segmentation with deep convolutional nets, atrous convolution, and fully connected crfs,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.161516Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f787951bed1e50bf78f892712075da1482e15de449c743481b313e4ab9065e75","observation_id":"decd7360-2d83-4cc1-b2bc-5b6da9989b96","resolution":{"observed_at":"2026-08-15T19:58:00.161516Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.166227Z","title":"Mask r-cnn,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.166227Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:b703d705b40c690252a49c1dd28f665c8f3e1d311b52dc1a3e460dfd02476ce4","observation_id":"70a339fd-deb5-4b23-bbd7-958d665c97da","resolution":{"observed_at":"2026-08-15T19:58:00.166227Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.171524Z","title":"Swin transformer: Hierarchical vision transformer using shifted windows,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.171524Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:2e552c6e52e26cf1ad26b19d6da7270f90b68249e35d4c25fc992104eda15b1f","observation_id":"4a9fe05f-29a5-460f-aeb6-22c28c1385d1","resolution":{"observed_at":"2026-08-15T19:58:00.171524Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.183133Z","title":"Segmenter: Transformer for semantic segmentation,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.183133Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:31075eaffdedf4747b88316b640b17c59f35d002729ff441e9cb98f9413496a9","observation_id":"571d3f7b-ec25-4b38-8fcf-e6ea0530473b","resolution":{"observed_at":"2026-08-15T19:58:00.183133Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1604.01685","last_updated":"2016-04-07T15:39:22Z","snapshot_observed_at":"2026-08-14T22:02:49.221229Z","submitted_at":"2016-04-06T16:34:33Z","title":"The Cityscapes Dataset for Semantic Urban Scene Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1604.01685","snapshot_observed_at":"2026-08-15T19:58:00.188091Z","title":"The cityscapes dataset for semantic urban scene understanding,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.188091Z"},"links":{"cited_paper":"/paper/1604.01685","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:a2aff4d545c7f05cad3a0b3bc729f4667023a238149f20d1b30c306fe1962c31","observation_id":"647ab6cf-f494-45f1-9aee-e1f080633f89","resolution":{"observed_at":"2026-08-15T19:58:00.188091Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.192582Z","title":"Bdd100k: A diverse driving dataset for heterogeneous multitask learning,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.192582Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f1463eb40a53d9463bfd20b68109a34bd4d1122ee41d02b4846d34627dca3de8","observation_id":"33b271bd-5607-4b4e-8568-681f3074baec","resolution":{"observed_at":"2026-08-15T19:58:00.192582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.197371Z","title":"Encoder-decoder with atrous separable convolution for semantic image segmentation,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.197371Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:86e871fb51b36c196753e4a5b2dbf04fd9e683f9b09b1e1f44924a5e2c249b28","observation_id":"388dc0ce-c081-4721-9909-857a38b090a3","resolution":{"observed_at":"2026-08-15T19:58:00.197371Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.202613Z","title":"Panoptic segmenta- tion,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.202613Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:5a211ee37630b0e15e9d3eff4d8ad701fb98f7028c58b52fa97063b04d7496d8","observation_id":"2540af70-2b28-4472-8c2e-58345bfff7a3","resolution":{"observed_at":"2026-08-15T19:58:00.202613Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.207362Z","title":"Deep learning for 3d point clouds: A survey,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.207362Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:06afb0e598e51e03598b4272a90cc8ddd4ecd5c4a7baabd6d4cfb6857b8680ee","observation_id":"bcabac83-2532-4c8b-8a6b-5699c54911df","resolution":{"observed_at":"2026-08-15T19:58:00.207362Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-16T09:25:53.087782Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-15T19:58:00.211595Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.211595Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f699a0b4c550db961968aa00e0426e37c3078e3d8a21c55a9b0dbd7839834f4b","observation_id":"81ed69cc-4f7e-49a3-91cd-37edfbecde20","resolution":{"observed_at":"2026-08-15T19:58:00.211595Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.217808Z","title":"nuscenes: A multimodal dataset for autonomous driv- ing,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.217808Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:355929ef7362813bc2582a446209b1c46c3b44875b4dbe395cb17f0be976ac3a","observation_id":"bd2536d8-c59d-4519-bc33-9c9720253d08","resolution":{"observed_at":"2026-08-15T19:58:00.217808Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.221880Z","title":"The mapillary vistas dataset for semantic under- standing of street scenes,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.221880Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:5c6e89c66c508a1838a3376aa313f05580ed08a7604323e66285b31f0bc0fb0c","observation_id":"e2924f5b-bd53-4950-88ab-9b6b9c948f3e","resolution":{"observed_at":"2026-08-15T19:58:00.221880Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1810.04805","last_updated":"2019-05-24T20:37:26Z","snapshot_observed_at":"2026-08-14T18:16:28.847993Z","submitted_at":"2018-10-11T00:50:01Z","title":"BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.04805","snapshot_observed_at":"2026-08-15T19:58:00.226946Z","title":"Bert: Pre-training of deep bidirectional transformers for language understanding,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.226946Z"},"links":{"cited_paper":"/paper/1810.04805","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:843268ab0b6ef5b9e695921f3393bd9f51a303139177f02b65621f45aa6c73da","observation_id":"5cc553a6-ab59-42e5-901b-0be0a6a310f6","resolution":{"observed_at":"2026-08-15T19:58:00.226946Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2005.14165","last_updated":"2020-07-22T19:47:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-05-28T17:29:03Z","title":"Language Models are Few-Shot Learners","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.14165","snapshot_observed_at":"2026-08-15T19:58:00.233011Z","title":"Language models are few-shot learners,","venue":null,"work_id":null,"year":1901},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.233011Z"},"links":{"cited_paper":"/paper/2005.14165","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:e01ff7a137c43bb25cf1eccb4d219305f9ce05a5e6839cb1cacd02b564d74cb3","observation_id":"23df71c3-1892-4c57-b8dd-3b3c89d0b40b","resolution":{"observed_at":"2026-08-15T19:58:00.233011Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1910.10683","last_updated":"2023-09-19T15:14:48Z","snapshot_observed_at":"2026-08-14T16:16:21.567225Z","submitted_at":"2019-10-23T17:37:36Z","title":"Exploring the Limits of Transfer Learning with a Unified Text-to-Text Transformer","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1910.10683","snapshot_observed_at":"2026-08-15T19:58:00.238216Z","title":"Exploring the limits of transfer learning with a unified text-to-text transformer,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.238216Z"},"links":{"cited_paper":"/paper/1910.10683","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:ab26630a4640ee83d7c9dfe2f1eb49247d9ec3f6461c1f3c48597bfc649ec03c","observation_id":"9513ab42-57ef-404a-81da-e512852c46c3","resolution":{"observed_at":"2026-08-15T19:58:00.238216Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.241734Z","title":"Learning transferable visual models from natural language supervision,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.241734Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:dfe584a7a4d5ef2d250058c5cfab061219718aadf35a646a40e83b3d75050a1f","observation_id":"ff39f22e-f832-4b98-beb6-e96f8c9064a6","resolution":{"observed_at":"2026-08-15T19:58:00.241734Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.07193","last_updated":"2024-02-02T10:24:09Z","snapshot_observed_at":"2026-08-11T10:12:11.384939Z","submitted_at":"2023-04-14T15:12:19Z","title":"DINOv2: Learning Robust Visual Features without Supervision","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.07193","snapshot_observed_at":"2026-08-15T19:58:00.245896Z","title":"Dinov2: Learning robust visual features without supervision,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.245896Z"},"links":{"cited_paper":"/paper/2304.07193","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:adf2bb764e104c3b297c3f74b1bcdff41a755d3e3e266a846632d6088af0feb8","observation_id":"8af6236e-1345-42fc-8a1c-dc6f41a1c960","resolution":{"observed_at":"2026-08-15T19:58:00.245896Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.06718","last_updated":"2023-07-11T18:13:14Z","snapshot_observed_at":"2026-08-16T15:40:37.274062Z","submitted_at":"2023-04-13T17:59:40Z","title":"Segment Everything Everywhere All at Once","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.06718","snapshot_observed_at":"2026-08-15T19:58:00.257150Z","title":"Segment everything everywhere all at once,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.257150Z"},"links":{"cited_paper":"/paper/2304.06718","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f956b2a7cfce5e0af11ff89a28c5041edeadbe28171f07ea8d04c21e763dadd4","observation_id":"15273a73-0d31-4d96-b28b-4f34dbe8fb3e","resolution":{"observed_at":"2026-08-15T19:58:00.257150Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.12289","last_updated":"2024-06-25T17:55:35Z","snapshot_observed_at":"2026-08-13T08:07:31.942031Z","submitted_at":"2024-02-19T17:04:04Z","title":"DriveVLM: The Convergence of Autonomous Driving and Large Vision-Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.12289","snapshot_observed_at":"2026-08-15T19:58:00.269411Z","title":"Drivelm: Driving with language and vision models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.269411Z"},"links":{"cited_paper":"/paper/2402.12289","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:a93db3c201ecfb6fad51c6216abb6b4bc20ad4586f73f01d2d2c7f9a7c15c528","observation_id":"5c83ba3e-afba-491e-8b4e-1b52f9ac3ba7","resolution":{"observed_at":"2026-08-15T19:58:00.269411Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.272980Z","title":"Traffic scene perception via multimodal large language model with data augmentation and efficient training strategy,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.272980Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:60b3e58ab2b5189bc96a36b8f4d81d55307975a73fec18a9820c26024f157eb1","observation_id":"b5d7c779-4236-48f5-9dee-f3e246d03318","resolution":{"observed_at":"2026-08-15T19:58:00.272980Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.17608","last_updated":"2024-03-28T15:24:16Z","snapshot_observed_at":"2026-08-16T14:06:18.741060Z","submitted_at":"2024-03-26T11:39:00Z","title":"Fake or JPEG? Revealing Common Biases in Generated Image Detection Datasets","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.17608","snapshot_observed_at":"2026-08-15T19:58:00.276737Z","title":"Exploring the roles of large language models in reshaping transportation systems: A survey, framework, and roadmap,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.276737Z"},"links":{"cited_paper":"/paper/2403.17608","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:3f34d3878b25c01dc792b3c5302e645bbf0b008cfe18b9dcf956a091554abe8c","observation_id":"5a0271dc-a676-4a3c-9bb3-ebd2970deea4","resolution":{"observed_at":"2026-08-15T19:58:00.276737Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"abs/1023456","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:02.176383Z","title":"Traffic scene analysis using vision-language models,","venue":null,"work_id":"b24a043f-a35b-485d-801b-e7fbff678b75","year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.280826Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:1c184ae8e8269a781063344bd316dbcd333781daa3acf203566c3e4f1b9eecca","observation_id":"c14c471d-a991-4115-95bd-4600a6309f7b","resolution":{"observed_at":"2026-08-15T19:58:02.183394Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.284446Z","title":"Talk2car: Taking control of your self-driving car,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.284446Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:1684d4ad09fbbd6ac598cf7e78a1520bec482db3ae316abddafeb61d3391263d","observation_id":"251c3dc8-c253-4f01-8bd1-808e3c03c461","resolution":{"observed_at":"2026-08-15T19:58:00.284446Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.12345","last_updated":"2024-05-20T19:42:20Z","snapshot_observed_at":"2026-08-16T13:51:11.021853Z","submitted_at":"2024-05-20T19:42:20Z","title":"Existence and uniqueness of solutions in the Lipschitz space of a functional equation and its application to the behavior of the paradise fish","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.12345","snapshot_observed_at":"2026-08-15T19:58:00.288689Z","title":"Road-seg-vl: Vision-language dataset for road scene segmentation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.288689Z"},"links":{"cited_paper":"/paper/2405.12345","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:6da064fa990029fb5322856686c5a5fed89d44be1e22fa62219ea4959b8bc3ca","observation_id":"ef333248-23cc-464b-868b-ff7c75485572","resolution":{"observed_at":"2026-08-15T19:58:00.288689Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.293240Z","title":"Image segmentation using text and image prompts,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.293240Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f1210131c42697be6f2ba07d4104634b21bce3742df24b755401bddd43aa12db","observation_id":"aceacbfd-de73-4bfe-8489-8521ca8ebb69","resolution":{"observed_at":"2026-08-15T19:58:00.293240Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.297201Z","title":"Scaling open-vocabulary image segmentation with image-level labels,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.297201Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:62af11de791b5ffded3563d6db17f3be27caad7a82aba3c66ec7fa84a91b91ca","observation_id":"af70fe27-9875-43e5-88b8-0bc3e1f39d3b","resolution":{"observed_at":"2026-08-15T19:58:00.297201Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.02643","last_updated":"2023-04-05T17:59:46Z","snapshot_observed_at":"2026-08-08T05:14:59.435033Z","submitted_at":"2023-04-05T17:59:46Z","title":"Segment Anything","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.02643","snapshot_observed_at":"2026-08-15T19:58:00.302122Z","title":"Segment anything,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.302122Z"},"links":{"cited_paper":"/paper/2304.02643","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:b04283ba49fa8b9d3a468b379ae532227ddd1b034fe824523494242b1a3ece63","observation_id":"2c845b31-8d39-4af9-a18b-1d5a2e70d7cc","resolution":{"observed_at":"2026-08-15T19:58:00.302122Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.14289","last_updated":"2023-07-01T07:26:22Z","snapshot_observed_at":"2026-08-12T07:11:05.316372Z","submitted_at":"2023-06-25T16:37:25Z","title":"Faster Segment Anything: Towards Lightweight SAM for Mobile Applications","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.14289","snapshot_observed_at":"2026-08-15T19:58:00.306243Z","title":"Faster segment anything: Towards lightweight sam for mobile applications,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.306243Z"},"links":{"cited_paper":"/paper/2306.14289","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:35d4364b2ca20fc7485a7577598306aba522e95ea78665cf9331e552088d0bb8","observation_id":"edcffdb4-ca86-4466-99ac-1d075f5b6d08","resolution":{"observed_at":"2026-08-15T19:58:00.306243Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2205.03436","last_updated":"2022-07-21T21:54:49Z","snapshot_observed_at":"2026-08-16T17:02:01.668812Z","submitted_at":"2022-05-06T18:17:19Z","title":"EdgeViTs: Competing Light-weight CNNs on Mobile Devices with Vision Transformers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2205.03436","snapshot_observed_at":"2026-08-15T19:58:00.310282Z","title":"Edgevits: Competing light-weight cnns on mobile devices with vision transformers,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.310282Z"},"links":{"cited_paper":"/paper/2205.03436","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f508093733cd255ddda349e284f8d898e669a63e8586828d69ca8d19c5ef100a","observation_id":"36d2752c-e765-4263-b4e8-a1a24f05caab","resolution":{"observed_at":"2026-08-15T19:58:00.310282Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.315365Z","title":"Robust image classification with multi- modal large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.315365Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:0f04fdaac05771b83d106d1db676573ee544a0a0cdaf449d910ae18f557a425f","observation_id":"09603f13-1ec2-43cb-9bbb-d392a7225a02","resolution":{"observed_at":"2026-08-15T19:58:00.315365Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.328554Z","title":"Driving forward: Semantic segmenta- tion in autonomous vehicles,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.328554Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:5af9a386e6f304b7bae98f97a27f18bbeb1c53d65eb648ac2593914b8504e6f7","observation_id":"6a292cef-eeae-405b-9647-ffa49c9a484c","resolution":{"observed_at":"2026-08-15T19:58:00.328554Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.336564Z","title":"Real-time semantic segmentation for autonomous driving: A review of cnns, transformers, and beyond,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.336564Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:75999c2248fe2e3547ca93cedc08157d167041fa8bc24556d95b04d0b8ae6ca6","observation_id":"69f1325c-e9ef-4358-b8cd-5e6e6c6e64af","resolution":{"observed_at":"2026-08-15T19:58:00.336564Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.11234","last_updated":"2024-05-18T09:11:57Z","snapshot_observed_at":"2026-08-16T13:51:41.455824Z","submitted_at":"2024-05-18T09:11:57Z","title":"Peculiarities of the chemical enrichment of metal-poor Stars in the Milky Way Galaxy","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.11234","snapshot_observed_at":"2026-08-15T19:58:00.323962Z","title":"Available: https://arxiv.org/abs/2405.11234","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.323962Z"},"links":{"cited_paper":"/paper/2405.11234","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:1cd579b22e52f95fbb992fe31216a4c5e5855e7afac242a093e07e31dac9f6ea","observation_id":"7e32be33-ea56-40ca-b213-e0853e1778ed","resolution":{"observed_at":"2026-08-15T19:58:00.323962Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13457","last_updated":"2024-02-07T15:54:59Z","snapshot_observed_at":"2026-08-16T15:30:55.086046Z","submitted_at":"2023-05-22T19:57:23Z","title":"Invariant tori and boundedness of solutions of non-smooth oscillators with Lebesgue integrable forcing term","version":2},"cited_work":{"arxiv_id":"2305.13457","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.13457","snapshot_observed_at":"2026-08-15T19:58:02.009461Z","title":"Invariant tori and boundedness of solutions of non-smooth oscillators with Lebesgue integrable forcing term","venue":"math.DS","work_id":"444a2f0c-94bc-4797-ac4f-e2f73ab9b78e","year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.344986Z"},"links":{"cited_paper":"/paper/2305.13457","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:74e3315d669e1ed7ce2bfd67a4586a8341f48e2e9a32b265934b624bdc7f8dcc","observation_id":"e6dc2a03-451a-4f51-b52c-b8cf2d16ca82","resolution":{"observed_at":"2026-08-15T19:58:02.014780Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.332582Z","title":"Available: https://www.keylabs.ai/blog/ driving-forward-semantic-segmentation-in-autonomous-vehicles/","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.332582Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:853ff4857182ad77392b1bd51db3dd5c2b8153385374a4eec0b3d967aac412cb","observation_id":"62029bee-d5da-4d03-9b51-aaf310f1f310","resolution":{"observed_at":"2026-08-15T19:58:00.332582Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.353143Z","title":"Oneformer: One transformer to rule them all for universal image segmentation,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.353143Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:891b2080d03fa973ea6de982bc9d24d7f109163d55725c12dffd4a93ca642f2d","observation_id":"91a0eeec-eaf4-493d-a13b-41ed62f94829","resolution":{"observed_at":"2026-08-15T19:58:00.353143Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1506.01497","last_updated":"2016-01-06T06:30:17Z","snapshot_observed_at":"2026-08-14T22:45:50.420062Z","submitted_at":"2015-06-04T07:58:34Z","title":"Faster R-CNN: Towards Real-Time Object Detection with Region Proposal Networks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1506.01497","snapshot_observed_at":"2026-08-15T19:58:00.340681Z","title":"Faster r-cnn: Towards real-time object detection with region proposal networks,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.340681Z"},"links":{"cited_paper":"/paper/1506.01497","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:523a467afdce369edff5266c0731b00c46148b8395b8e9e81a961ca3a9a46ae7","observation_id":"350ca310-9e09-442d-814f-ffae257b6413","resolution":{"observed_at":"2026-08-15T19:58:00.340681Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.360510Z","title":"Fully convolutional networks for semantic segmentation,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.360510Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f0b87a1a6e13965d854702a96d1703d2bb9df18be84814304b7ccadaa33891ee","observation_id":"75d79032-23cf-449d-9a86-38c4a27a29bf","resolution":{"observed_at":"2026-08-15T19:58:00.360510Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.349724Z","title":"Masked-attention mask transformer for universal im- age segmentation,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.349724Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:50ffed7f68d11dce787cdb9b4e3fcb18f5fe1eecab365160b3463ae37256a1a8","observation_id":"3adeb662-cd88-46af-bc7e-d32118e083fd","resolution":{"observed_at":"2026-08-15T19:58:00.349724Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.367668Z","title":"U-net: Convolutional net- works for biomedical image segmentation,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.367668Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:07b781a028523937cbf99ef08d68d4e5db71e046e888679ad4f573a37fdd732e","observation_id":"cff49262-f75a-4b6d-ac54-925f5de81230","resolution":{"observed_at":"2026-08-15T19:58:00.367668Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.356499Z","title":"Normalized cuts and image segmentation,","venue":null,"work_id":null,"year":2000},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.356499Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:300666578ea20cb86d9fb6dae2c56c31938cfcd23875f3c7a538cc2bd4eb5572","observation_id":"4490283d-8630-41e5-b47a-e7d5130cae8b","resolution":{"observed_at":"2026-08-15T19:58:00.356499Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.374463Z","title":"Cross-modal self-attention network for referring image segmentation,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.374463Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:05df6064abe29deb7968e2f24808511ed955f675a61fc86807ee07edde64d1ad","observation_id":"e04d352a-416e-4996-8709-80595dafadde","resolution":{"observed_at":"2026-08-15T19:58:00.374463Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.364133Z","title":"Segnet: A deep con- volutional encoder-decoder architecture for image segmentation,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.364133Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:6eae4b3f5e29e4b399c9f8a75a005bdc307fb7b4e218a9d7997c30778228b85b","observation_id":"06109288-1f65-4725-b2ed-7214a0acb0bf","resolution":{"observed_at":"2026-08-15T19:58:00.364133Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.381843Z","title":"Phrasecut: Language-based image segmentation in the wild,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.381843Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:b09b75c1562439d86ec79b07d4ded2dc7c616b67e3e4af98690e7693a4a9461d","observation_id":"2b306609-5358-43cb-87a1-adb64d6b1592","resolution":{"observed_at":"2026-08-15T19:58:00.381843Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.371100Z","title":"Segformer: Simple and efficient design for semantic segmentation with transformers,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.371100Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:a51b5296b6cdfb7e1dca6a7a932a44487b840222b320a798ec2a9c57fed9eb4b","observation_id":"3d48b828-9da1-4ce8-9f47-f0854cbcb23f","resolution":{"observed_at":"2026-08-15T19:58:00.371100Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.389543Z","title":"Vilbert: Pretraining task- agnostic visiolinguistic representations for vision-and-language tasks,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.389543Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:47fb5762a49409e3d2859a57326e4bf49e86f414223486038152fcf50da668d6","observation_id":"845feecc-aada-4115-851d-826350490344","resolution":{"observed_at":"2026-08-15T19:58:00.389543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.378346Z","title":"Bi-directional re- lationship inferring network for referring image segmentation,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.378346Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:da0a4cb1a105b5df605cce20ad3ad904be93a439e95d52774cb153330a37707f","observation_id":"b60d39ab-5105-4f9a-8051-ead25e1ee50b","resolution":{"observed_at":"2026-08-15T19:58:00.378346Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.396379Z","title":"Imagenet classification with deep convolutional neural networks,","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.396379Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:77d25bac49c4abff0bf85a38f306595f9fdece7b078e73220ce7acc2c27f2949","observation_id":"73241321-e5a6-4e07-abfe-bb75f5882e26","resolution":{"observed_at":"2026-08-15T19:58:00.396379Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.385541Z","title":"Refvos: A closer look at referring expressions for video object segmentation,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.385541Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:c07fa68cf484213710f984a350f912acb454a20f514912ed532c0b5cb2d146f6","observation_id":"4958ea0c-5684-4a7f-a240-f252f6a57d3e","resolution":{"observed_at":"2026-08-15T19:58:00.385541Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.404520Z","title":"Attention is all you need,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.404520Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:933bb3569f0edfb42c604b18810a954957ff6d1604da42b95a13d9a1ac4f0da0","observation_id":"d6e573bc-74c0-45ee-a36d-d3fb5b3cc343","resolution":{"observed_at":"2026-08-15T19:58:00.404520Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.392927Z","title":"Lxmert: Learning cross-modality encoder rep- resentations from transformers,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.392927Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:8fd8073df46da516ba8adb430fc7de3572d98148617ca72b0828f2f4bd8d1cdc","observation_id":"8d43780a-2e77-4301-b798-b1e969882036","resolution":{"observed_at":"2026-08-15T19:58:00.392927Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.08767","last_updated":"2024-04-12T18:45:51Z","snapshot_observed_at":"2026-08-16T14:01:20.100950Z","submitted_at":"2024-04-12T18:45:51Z","title":"LLM-Seg: Bridging Image Segmentation and Large Language Model Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.08767","snapshot_observed_at":"2026-08-15T19:58:00.416443Z","title":"Llm-seg: Bridging image segmentation with large language models reasoning,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.416443Z"},"links":{"cited_paper":"/paper/2404.08767","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:9822dc9965c81620f502098fc3fff468945da5dfe2d9a02b715d9be6fa92d173","observation_id":"e0eb2a2a-f75a-48e2-9b50-68f954a35086","resolution":{"observed_at":"2026-08-15T19:58:00.416443Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1409.1556","last_updated":"2015-04-10T16:25:04Z","snapshot_observed_at":"2026-08-14T23:20:42.336514Z","submitted_at":"2014-09-04T19:48:04Z","title":"Very Deep Convolutional Networks for Large-Scale Image Recognition","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1409.1556","snapshot_observed_at":"2026-08-15T19:58:00.400744Z","title":"Very deep convolutional networks for large-scale image recognition,","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.400744Z"},"links":{"cited_paper":"/paper/1409.1556","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:ef5bb3dc637c7823262d246cca1c676e91322ea8e637583fa1cdc6c313792e61","observation_id":"ac976d36-1439-4902-87fe-e872cfdd73d2","resolution":{"observed_at":"2026-08-15T19:58:00.400744Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.12345","last_updated":"2023-05-21T04:43:55Z","snapshot_observed_at":"2026-08-16T15:31:25.018435Z","submitted_at":"2023-05-21T04:43:55Z","title":"Overspinning a rotating black hole in semiclassical gravity with type-A trace anomaly","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.12345","snapshot_observed_at":"2026-08-15T19:58:00.426013Z","title":"A survey on vision-language foundation models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.426013Z"},"links":{"cited_paper":"/paper/2305.12345","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:bf38bab9d1a579653110f75285e8bd858baa39cfdb637384fe5899b95b9677d4","observation_id":"54cbab60-a88e-4aed-8fb7-545a0c911fdf","resolution":{"observed_at":"2026-08-15T19:58:00.426013Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1706.03762","last_updated":"2023-08-02T00:41:18Z","snapshot_observed_at":"2026-08-17T01:19:18.409791Z","submitted_at":"2017-06-12T17:57:34Z","title":"Attention Is All You Need","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1706.03762","snapshot_observed_at":"2026-08-15T19:58:00.408725Z","title":"Available: https://arxiv.org/abs/1706.03762","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.408725Z"},"links":{"cited_paper":"/paper/1706.03762","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:1eca0c657d39d732d93a84fe06bc3f61f0944caf678ad1db7a20d1d798e5aa8c","observation_id":"e2953afd-f3b6-477a-bcfa-7b0d7b045835","resolution":{"observed_at":"2026-08-15T19:58:00.408725Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.412511Z","title":"Semantic segmentation datasets for autonomous driving,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.412511Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:e76172c47cf88a5b3e9550696fa45fb1cb9108f6b2b6706eda0c359bdaef6abd","observation_id":"6684055d-7d0f-41ae-ad8a-cdd2db99af1b","resolution":{"observed_at":"2026-08-15T19:58:00.412511Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.17890","last_updated":"2024-07-19T08:12:24Z","snapshot_observed_at":"2026-08-16T13:57:20.964640Z","submitted_at":"2024-04-27T12:55:13Z","title":"DPER: Diffusion Prior Driven Neural Representation for Limited Angle and Sparse View CT Reconstruction","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.17890","snapshot_observed_at":"2026-08-15T19:58:00.439085Z","title":"Open-vocabulary vision-language segmentation for autonomous driving,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.439085Z"},"links":{"cited_paper":"/paper/2404.17890","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:113de1cbfdff2ad8d6947aff911fd23e6c9b7ed6b883c1e946b82342d3205ea3","observation_id":"9885fcc2-eb39-40a3-bc44-d7844e31e5c4","resolution":{"observed_at":"2026-08-15T19:58:00.439085Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12320","last_updated":"2023-11-21T03:32:01Z","snapshot_observed_at":"2026-08-16T14:41:38.626307Z","submitted_at":"2023-11-21T03:32:01Z","title":"A Survey on Multimodal Large Language Models for Autonomous Driving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12320","snapshot_observed_at":"2026-08-15T19:58:00.421328Z","title":"A survey on multimodal large language models for autonomous driving,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.421328Z"},"links":{"cited_paper":"/paper/2311.12320","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:c63666ce89958677eab8b16fd2d7079c96df18b1d99839b3f9628fcfd35ec299","observation_id":"254657f6-bd64-4f14-94d3-bf0f2c727af9","resolution":{"observed_at":"2026-08-15T19:58:00.421328Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.446800Z","title":"Are we ready for autonomous driving? the kitti vision benchmark suite,","venue":null,"work_id":null,"year":2012},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.446800Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:d3a3c3ee3b27415d41d4029550652b586433ef81e6e72a4a6ccd9a574807030d","observation_id":"bfd87f47-d46b-41e4-8a7a-f2d0ad3e11e2","resolution":{"observed_at":"2026-08-15T19:58:00.446800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.09812","last_updated":"2024-05-16T04:54:41Z","snapshot_observed_at":"2026-08-16T13:52:18.003333Z","submitted_at":"2024-05-16T04:54:41Z","title":"Mean-field and cumulant approaches to modelling organic polariton physics","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.09812","snapshot_observed_at":"2026-08-15T19:58:00.429947Z","title":"Architectures for vision-language segmentation: A survey,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.429947Z"},"links":{"cited_paper":"/paper/2405.09812","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:7eeee058b33a50eab75ab3c594b1bd8dc54e52ffb30650c98af959ca5176b2c4","observation_id":"3c9e90d3-4a07-44b4-9a1b-20674f6aa6a0","resolution":{"observed_at":"2026-08-15T19:58:00.429947Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.07115","last_updated":"2022-07-18T17:56:53Z","snapshot_observed_at":"2026-08-16T19:42:46.594723Z","submitted_at":"2022-07-14T17:59:37Z","title":"XMem: Long-Term Video Object Segmentation with an Atkinson-Shiffrin Memory Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2207.07115","snapshot_observed_at":"2026-08-15T19:58:00.433822Z","title":"Xmem: Long-term video object segmentation with an atkinson-shiffrin memory model,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.433822Z"},"links":{"cited_paper":"/paper/2207.07115","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:73d66f7f1c29d7037d1c6f909c144c8497548925f5de19059ca7bb80024824d6","observation_id":"1e2acdb4-0269-44ab-889c-c05328cfb671","resolution":{"observed_at":"2026-08-15T19:58:00.433822Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.08299","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:01.859771Z","title":"Efficient unstructured pruning of mamba state-space models for resource-constrained environments,","venue":null,"work_id":"6a7df7d9-c48e-479b-a74f-ca8cab30db1a","year":2025},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.460168Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:1aa960e70c3df2bac8e8baa92a2667a8f804c79d91033b47c74eed1a785b7952","observation_id":"cac7fad4-be08-47b2-9ad5-30d6b9f0b730","resolution":{"observed_at":"2026-08-15T19:58:01.866581Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.14567","last_updated":"2024-04-22T20:29:58Z","snapshot_observed_at":"2026-08-16T13:58:44.706224Z","submitted_at":"2024-04-22T20:29:58Z","title":"WangLab at MEDIQA-M3G 2024: Multimodal Medical Answer Generation using Large Language Models","version":1},"cited_work":{"arxiv_id":"2404.14567","doi":null,"metadata_source":"pith","pith_arxiv_id":"2404.14567","snapshot_observed_at":"2026-08-15T19:58:01.879198Z","title":"WangLab at MEDIQA-M3G 2024: Multimodal Medical Answer Generation using Large Language Models","venue":"cs.CL","work_id":"0f6d97ee-d110-4f33-85dc-5f44e8c6b62e","year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.443101Z"},"links":{"cited_paper":"/paper/2404.14567","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f137bc8eae71110dbf59b3025bf7d585218952370f13488a35d34e10d74e5a21","observation_id":"f6b1fb5a-955f-41dd-a0fb-bcdb5043e826","resolution":{"observed_at":"2026-08-15T19:58:01.883605Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2103.10385","last_updated":"2023-10-25T06:15:58Z","snapshot_observed_at":"2026-08-16T18:37:56.706777Z","submitted_at":"2021-03-18T17:13:50Z","title":"GPT Understands, Too","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2103.10385","snapshot_observed_at":"2026-08-15T19:58:00.468134Z","title":"P-tuning: Prompt tuning can be as good as fine-tuning on language models,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.468134Z"},"links":{"cited_paper":"/paper/2103.10385","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:bf1f2bd021da728a3ce85596693d4a7eec32e3f9ec245dc1ae354c0d6e134395","observation_id":"7c30b918-481c-45d1-bfa3-525b5b07766d","resolution":{"observed_at":"2026-08-15T19:58:00.468134Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.450695Z","title":"Deep residual learning for image recognition,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.450695Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:844a1b9b28f1b71592ce1fe4dcaa28a29deea5b58c62dc7cd87f86d6b5b877ce","observation_id":"e9c754aa-97bf-43c3-b9da-318f88211f36","resolution":{"observed_at":"2026-08-15T19:58:00.450695Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.455566Z","title":"Efficientvit: Memory efficient vision transformer with cascaded group attention,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.455566Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:c24c42ad8f92fde4b7c82815459d644ce90f0c9e1a5537871607586c7e92503f","observation_id":"3f4d3675-8176-4d76-9df9-7a809b33d044","resolution":{"observed_at":"2026-08-15T19:58:00.455566Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.479020Z","title":"Multimodal compact bilinear pooling for visual ques- tion answering and visual grounding,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.479020Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:e49d1f69134c36a029ea3cd53ee8054dd72a466fe09a4e176e7570dd7645cbb4","observation_id":"fc964efb-ffd4-4742-8326-6e62fa99ccc7","resolution":{"observed_at":"2026-08-15T19:58:00.479020Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.08691","last_updated":"2021-09-02T17:34:41Z","snapshot_observed_at":"2026-08-16T20:01:36.160048Z","submitted_at":"2021-04-18T03:19:26Z","title":"The Power of Scale for Parameter-Efficient Prompt Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.08691","snapshot_observed_at":"2026-08-15T19:58:00.464208Z","title":"The power of scale for parameter-efficient prompt tuning,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.464208Z"},"links":{"cited_paper":"/paper/2104.08691","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:de7da8b551eea16b0267af706057fd730e71cf85a60608e4a36b9782b21895a9","observation_id":"313c5e9f-78f6-409b-95d2-aade2de02de8","resolution":{"observed_at":"2026-08-15T19:58:00.464208Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2102.02779","last_updated":"2021-05-23T23:12:46Z","snapshot_observed_at":"2026-08-17T00:24:12.793055Z","submitted_at":"2021-02-04T17:59:30Z","title":"Unifying Vision-and-Language Tasks via Text Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2102.02779","snapshot_observed_at":"2026-08-15T19:58:00.486496Z","title":"Unifying vision-and-language tasks via text genera- tion,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.486496Z"},"links":{"cited_paper":"/paper/2102.02779","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:9189e4e6cef9ab93c0f8a6956213457f2d8fead95ab426b7fb253ee24ff1bea6","observation_id":"4701650d-827b-4e98-8b9d-2bb1acff2070","resolution":{"observed_at":"2026-08-15T19:58:00.486496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.471775Z","title":"Unify, align and refine: A unified framework for vision-and-language pre-training,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.471775Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:23854ea913b2da5e8ee4ca20c00c50744b1a7365405795c308726c153ed9a4f5","observation_id":"453c9653-25de-4a8b-a3a6-b05e1c1dd0c6","resolution":{"observed_at":"2026-08-15T19:58:00.471775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.12597","last_updated":"2023-06-15T07:57:29Z","snapshot_observed_at":"2026-08-12T12:01:54.105712Z","submitted_at":"2023-01-30T00:56:51Z","title":"BLIP-2: Bootstrapping Language-Image Pre-training with Frozen Image Encoders and Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.12597","snapshot_observed_at":"2026-08-15T19:58:00.475144Z","title":"Blip-2: Bootstrapping language- image pre-training with frozen image encoders and large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.475144Z"},"links":{"cited_paper":"/paper/2301.12597","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:2ef0511d3fd019cc9350fbdc6a09cef848d05706599ac7f700ac708852773c9f","observation_id":"93b2528d-53cc-401b-983c-3b008f100bc6","resolution":{"observed_at":"2026-08-15T19:58:00.475144Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14876","last_updated":"2025-01-23T01:08:19Z","snapshot_observed_at":"2026-08-17T00:25:19.104335Z","submitted_at":"2024-04-02T01:42:32Z","title":"Precise and Robust Sidewalk Detection: Leveraging Ensemble Learning to Surpass LLM Limitations in Urban Environments","version":2},"cited_work":{"arxiv_id":"2405.14876","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.14876","snapshot_observed_at":"2026-08-15T19:58:01.701667Z","title":"Precise and Robust Sidewalk Detection: Leveraging Ensemble Learning to Surpass LLM Limitations in Urban Environments","venue":"cs.CV","work_id":"62292b60-7295-4435-920a-e7b8fe929193","year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.497982Z"},"links":{"cited_paper":"/paper/2405.14876","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:8ce6c1e02520567061b0bdc346552c359fbeaac2ea4aa1d7f1adb69d84bcf059","observation_id":"fd92956f-e1d5-4be3-8dbc-831c488787e0","resolution":{"observed_at":"2026-08-15T19:58:01.706727Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.482776Z","title":"Bilinear attention net- works,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.482776Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:e7ab1f69bdd3e86ee0e0595945dcb4ea7dbd4662d511252c6bb40ce30c1d6293","observation_id":"9709f0da-de33-451f-9fe9-94b96b7e65f9","resolution":{"observed_at":"2026-08-15T19:58:00.482776Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.505857Z","title":"Crash time matters: Hybridmamba for fine-grained temporal localization in traffic surveillance footage,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.505857Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:8eb7ac8f44e36772a364b7b90678fed2dc6074336b4ebeca72d0716e5452e3ca","observation_id":"a35832fb-3e71-4604-a6d0-cd21b98f960b","resolution":{"observed_at":"2026-08-15T19:58:00.505857Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03744","last_updated":"2024-05-15T19:22:44Z","snapshot_observed_at":"2026-08-15T02:35:59.111911Z","submitted_at":"2023-10-05T17:59:56Z","title":"Improved Baselines with Visual Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03744","snapshot_observed_at":"2026-08-15T19:58:00.490113Z","title":"Llava-1.5: Improved baselines for visual instruction tuning,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.490113Z"},"links":{"cited_paper":"/paper/2310.03744","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:13d9125a69bcbbb67ceac2f2ba41d14d6020061da737d459502cc6a886ff22bf","observation_id":"fe95bcdc-6188-4d4a-94c6-7ae2f67b6949","resolution":{"observed_at":"2026-08-15T19:58:00.490113Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.00692","last_updated":"2024-05-01T05:10:13Z","snapshot_observed_at":"2026-08-16T15:11:32.772599Z","submitted_at":"2023-08-01T17:50:17Z","title":"LISA: Reasoning Segmentation via Large Language Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.00692","snapshot_observed_at":"2026-08-15T19:58:00.493813Z","title":"Lisa: Reasoning segmentation via large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.493813Z"},"links":{"cited_paper":"/paper/2308.00692","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:027e8062db226453ac76f7e21f6fe142972650081898f10bf77d3fbddf4dc0ec","observation_id":"bc203a75-bf37-407e-95a1-63c658c2f0dc","resolution":{"observed_at":"2026-08-15T19:58:00.493813Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.00194","last_updated":"2023-07-31T23:15:45Z","snapshot_observed_at":"2026-08-16T15:11:47.908022Z","submitted_at":"2023-07-31T23:15:45Z","title":"Optimal Qubit Reuse for Near-Term Quantum Computers","version":1},"cited_work":{"arxiv_id":"2308.00194","doi":null,"metadata_source":"pith","pith_arxiv_id":"2308.00194","snapshot_observed_at":"2026-08-15T19:58:01.522892Z","title":"Optimal Qubit Reuse for Near-Term Quantum Computers","venue":"quant-ph","work_id":"435a8932-d3b6-4273-8cb1-a55466e5b2b5","year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.517806Z"},"links":{"cited_paper":"/paper/2308.00194","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:6c52a21e331c90c70f0a257748978e0a4965f6cea3ae01a5fc4e514f0b78eb8c","observation_id":"937822cb-63f4-401c-97df-a274ed196983","resolution":{"observed_at":"2026-08-15T19:58:01.530270Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2504.19684","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:01.677795Z","title":"Clearvision: Leveraging cyclegan and siglip-2 for robust all- weather classification in traffic camera imagery,","venue":null,"work_id":"1a81f2ba-c83d-4c74-b15b-d62bfff9d023","year":2025},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.501825Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:47eca61d5752febe83bf37adf6260fc01ae1900ab5494ff27a6ce5f62921c395","observation_id":"d6c172a5-485e-484c-92ea-f33e9f928537","resolution":{"observed_at":"2026-08-15T19:58:01.688409Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2212.08411","last_updated":"2022-12-16T11:25:01Z","snapshot_observed_at":"2026-08-16T16:08:09.865445Z","submitted_at":"2022-12-16T11:25:01Z","title":"Indiscernibles and satisfaction classes in arithmetic","version":1},"cited_work":{"arxiv_id":"2212.08411","doi":null,"metadata_source":"pith","pith_arxiv_id":"2212.08411","snapshot_observed_at":"2026-08-15T19:58:01.498605Z","title":"Indiscernibles and satisfaction classes in arithmetic","venue":"math.LO","work_id":"f47c7614-fc05-423b-8745-e2e77622b56c","year":2022},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.527503Z"},"links":{"cited_paper":"/paper/2212.08411","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:537c902a7e70b01d895b879c14a0cdb82d7b252fb5f191a47a621108264c03af","observation_id":"63737ba6-55a8-4212-b1b5-55039a8e4b78","resolution":{"observed_at":"2026-08-15T19:58:01.503885Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.510086Z","title":"Pointnet: Deep learning on point sets for 3d classification and segmentation,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.510086Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:7cf18fc011ae0ec64366fccd355c079315465a05b625b5542a867409002c84d1","observation_id":"35631a85-ca93-4d54-b6e3-6c38204dcdc1","resolution":{"observed_at":"2026-08-15T19:58:00.510086Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.513878Z","title":"Pointnet++: Deep hierarchical feature learning on point sets in a metric space,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.513878Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:9d43fe1128c3f8d7a71bc2fc55b6b6a42ef1687b72ecfe8d7a99b58197766903","observation_id":"08af4a26-a045-41e7-815f-37b3c45d3576","resolution":{"observed_at":"2026-08-15T19:58:00.513878Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.539234Z","title":"Communication-efficient learning of deep networks from decentral- ized data,","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.539234Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:9d4f8ad863046da7ced5140712e95108659960c2203e7ba8e03c1bb34a08db17","observation_id":"bb5bfc43-fed8-4fc9-a1dc-3f46bb210ee4","resolution":{"observed_at":"2026-08-15T19:58:00.539234Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.522947Z","title":"Clip2scene: Towards label-efficient 3d scene understanding by clip,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.522947Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:5134f69ded87ecb010b7da2635423461b97b33110d77aa0328abf73ab4aa806e","observation_id":"94960be8-8710-40b6-a687-c7f5fb99288b","resolution":{"observed_at":"2026-08-15T19:58:00.522947Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.546211Z","title":"Federated learning: Challenges, methods, and future directions,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.546211Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:2c1fbf7dd4edc812d8772db0cd558edfd37b63c16456f76876921e798d3844a2","observation_id":"dd2c064d-8450-4298-9b70-1e19fd8449c4","resolution":{"observed_at":"2026-08-15T19:58:00.546211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.18115","last_updated":"2024-06-10T10:52:54Z","snapshot_observed_at":"2026-08-16T14:14:38.861826Z","submitted_at":"2024-02-28T07:05:27Z","title":"UniVS: Unified and Universal Video Segmentation with Prompts as Queries","version":2},"cited_work":{"arxiv_id":"2402.18115","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.18115","snapshot_observed_at":"2026-08-15T19:58:01.479137Z","title":"UniVS: Unified and Universal Video Segmentation with Prompts as Queries","venue":"cs.CV","work_id":"14f174ec-b747-42e2-bec4-d2f8c2ea6cb2","year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.531795Z"},"links":{"cited_paper":"/paper/2402.18115","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:dc162b97232402c8a31d6fe890e2c7a39464c0200e156ea625fbea09f97462ef","observation_id":"252ce833-a51e-4bb0-bcd9-c49496be9667","resolution":{"observed_at":"2026-08-15T19:58:01.484159Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.535726Z","title":"Video object segmentation: A survey,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.535726Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:1fc87141534e45c4ce664fe720c2f2348bc9353f4c99027a41af74890eb87f94","observation_id":"f439cf81-3bb4-47ca-a43c-43af95a9b196","resolution":{"observed_at":"2026-08-15T19:58:00.535726Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.557443Z","title":"Towards explainable traffic flow prediction with large language models,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.557443Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f0f2461c44bb250f69b7b878685cf0d960258ba0de773b6dc60edef86897a530","observation_id":"f093e63f-f03e-4fb6-92f8-4d3c32d97d8c","resolution":{"observed_at":"2026-08-15T19:58:00.557443Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.542692Z","title":"Advances and open problems in federated learning,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.542692Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:6e1b8b5e3b30ab3a30af3d4399b7e29975ad408b28c0daada441cbe51847d237","observation_id":"7eec3855-938b-4d41-b691-4a51b50e53b6","resolution":{"observed_at":"2026-08-15T19:58:00.542692Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:02.508893Z","title":"A review of deep learning- based methods for pavement defect detection,","venue":null,"work_id":"54ee7174-acb5-4c9d-9c51-713fb718a386","year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.566214Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:ab95092e4aed649a289e02a27d68f54554a439174bc222900c0247b35d902612","observation_id":"b4b3495f-b751-48e2-8c80-0686d580c2a9","resolution":{"observed_at":"2026-08-15T19:58:02.512999Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1904.09333","last_updated":"2020-07-15T10:19:00Z","snapshot_observed_at":"2026-08-14T16:43:51.144969Z","submitted_at":"2019-04-19T20:26:26Z","title":"General derivative Thomae formula for singular half-periods","version":3},"cited_work":{"arxiv_id":"1904.09333","doi":null,"metadata_source":"pith","pith_arxiv_id":"1904.09333","snapshot_observed_at":"2026-08-15T19:58:01.458389Z","title":"General derivative Thomae formula for singular half-periods","venue":"math.AG","work_id":"a3214dba-0f4c-40c4-b794-ad33ae5108fd","year":2019},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.550054Z"},"links":{"cited_paper":"/paper/1904.09333","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:24c5560466da90bc03f8cc4187bdeb8212188dea65cc6f47fc0faeda347b81f8","observation_id":"7c7e5b9b-387b-4c5c-8125-95f033617598","resolution":{"observed_at":"2026-08-15T19:58:01.464460Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.553880Z","title":"Cooper: A query-based collaborative perception framework for 3d object detection,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":96,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.553880Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:04f8ad9e963b0155b672c9166cad07115fc71a98ca9d72d8a0eee3b12d3e3b2d","observation_id":"22ef35fa-78d8-4675-9e7b-b1b6c0ca6cd8","resolution":{"observed_at":"2026-08-15T19:58:00.553880Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.09935","last_updated":"2024-10-23T19:59:33Z","snapshot_observed_at":"2026-08-16T14:09:38.287623Z","submitted_at":"2024-03-15T00:35:08Z","title":"Topological frequency conversion in rhombohedral multilayer graphene","version":2},"cited_work":{"arxiv_id":"2403.09935","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.09935","snapshot_observed_at":"2026-08-15T19:58:01.357490Z","title":"Topological frequency conversion in rhombohedral multilayer graphene","venue":"cond-mat.mes-hall","work_id":"48a08c6c-0f5b-4bcc-8c9c-e25915a8d659","year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":97,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.579441Z"},"links":{"cited_paper":"/paper/2403.09935","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:17be01e8e26ecb22e95520e8d0900a6e828403f883fcaa6795d2034ae2371fc8","observation_id":"1ce0542f-456c-4e8b-83b5-2b32e23f2bf7","resolution":{"observed_at":"2026-08-15T19:58:01.362235Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:00.562218Z","title":"Robust and precise sidewalk detection with ensemble learning: Enhancing road safety and facilitating curb space management,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.562218Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:1d1fd18b9e1ef86d7d330a4282f140ee379ed68c4a3ca6edd9ed7b5c20201004","observation_id":"dfc44bcc-ce4d-435b-b010-58eb2323482e","resolution":{"observed_at":"2026-08-15T19:58:00.562218Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.08226","last_updated":"2023-10-12T11:17:54Z","snapshot_observed_at":"2026-08-16T14:52:44.871118Z","submitted_at":"2023-10-12T11:17:54Z","title":"Evaluate PAC codes via Efficient Estimation on Weight Distribution","version":1},"cited_work":{"arxiv_id":"2310.08226","doi":null,"metadata_source":"pith","pith_arxiv_id":"2310.08226","snapshot_observed_at":"2026-08-15T19:58:01.323904Z","title":"Evaluate PAC codes via Efficient Estimation on Weight Distribution","venue":"cs.IT","work_id":"39373a42-9198-4067-b67a-8d2546e3dca5","year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":99,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.591071Z"},"links":{"cited_paper":"/paper/2310.08226","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:ca2caca762f15d6a3fde788373dd7fe31f85b9d8e6f399b7b1f4cb5d7337563d","observation_id":"b11aa495-8c5c-49e0-b22a-e201d0125a27","resolution":{"observed_at":"2026-08-15T19:58:01.329276Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:02.496077Z","title":"Road pothole detection and classification using deep convolutional neu- ral networks,","venue":null,"work_id":"181ccd5d-d37f-43c2-bec7-3271a2e5b487","year":2018},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":100,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.570631Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:66d1d5b4f56c5db749a212c9e7fa9f15738747b2f197c941ed6d1ad70ba47193","observation_id":"28207b59-6aab-4c8b-94e5-76f081eef182","resolution":{"observed_at":"2026-08-15T19:58:02.500592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:02.481653Z","title":"Lingo-1: A foundation model for language-driven autonomous vehicles,","venue":null,"work_id":"4df3f513-2bd7-413f-93ff-118c8982ef87","year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":101,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.574745Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:504c28b72f4cc25f9e09b8f5fa30be21597ce4703ff11ccf097e21de460896ff","observation_id":"1d5db5cb-ccb8-40ef-a346-232fdf2d0157","resolution":{"observed_at":"2026-08-15T19:58:02.486793Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2403.11545","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T19:58:01.258423Z","title":"Talk2bev: Language- grounded bird’s-eye-view for autonomous driving,","venue":null,"work_id":"44377ca0-8a8f-410b-8f7e-abb53da56f06","year":2024},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":102,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.610186Z"},"links":{"citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:6afab9063ec3706bb2f8eb565f21bfc991dee5c43f95165585eea63eebc757d3","observation_id":"c5f67be0-bad6-48d9-86b9-a8e8e084c486","resolution":{"observed_at":"2026-08-15T19:58:01.265524Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-16T06:30:59.297886+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.07488","last_updated":"2023-12-21T05:37:58Z","snapshot_observed_at":"2026-08-16T14:35:23.450080Z","submitted_at":"2023-12-12T18:24:15Z","title":"LMDrive: Closed-Loop End-to-End Driving with Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.07488","snapshot_observed_at":"2026-08-15T19:58:00.586091Z","title":"Lmdrive: Closed-loop end-to-end driving with large language models,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems","version":2},"reference_index":103,"source":"pdf_text","source_observed_at":"2026-08-15T19:58:00.586091Z"},"links":{"cited_paper":"/paper/2312.07488","citing_paper":"/paper/2506.14096"},"observation_digest":"sha256:f2c2172402f3cc0ae0be964e51ab881e6066221e185af900009de2567ed3202f","observation_id":"23935f2c-e0a4-4d53-b27e-19e63b15e906","resolution":{"observed_at":"2026-08-15T19:58:00.586091Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.14096","last_updated":"2025-09-05T22:03:39Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-17T00:24:39.245252Z","submitted_at":"2025-06-17T01:20:50Z","title":"Image Segmentation with Large Language Models: A Survey with Perspectives for Intelligent Transportation Systems"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":4,"parse_uncertain":0,"unresolved":84,"verified_exact":9,"verified_fuzzy":3},"total_outbound_references":130},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-16T06:30:59.297886+00:00","source":"crossref"},{"observed_at":"2026-08-16T06:30:54.164669+00:00","source":"retraction_watch"}],"thesis":"As of 17 August 2026, this Paper Citation Record lists 100 of 130 outbound references and 1 inbound Pith citation observation for arXiv:2506.14096."}