{"as_of":"2026-08-12T23:26:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:e41be3afffbb081f1ee3b428ad156a492ccc15daca7c2aa2a22353a18009dd73","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":10,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":10,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":10,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":10,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T12:24:56.480508Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T23:17:28.876548Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-08-12T12:24:56.480508Z","title":"Contextual non-local alignment over full-scale repre- sentation for text-based person search","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.17776","last_updated":"2025-07-27T12:25:05Z","snapshot_observed_at":"2026-08-12T12:16:20.057740Z","submitted_at":"2024-11-26T09:50:15Z","title":"Beyond Walking: A Large-Scale Image-Text Benchmark for Text-based Person Anomaly Search","version":3},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T12:24:56.480508Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2411.17776"},"observation_digest":"sha256:36f2b32afc1a8eeec586e068b5ebb8445e98ffedf6e4c34056461152255f06ba","observation_id":"3dc1fc50-2577-4161-b886-1d75c94151af","resolution":{"observed_at":"2026-08-12T12:24:56.480508Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-08-10T23:20:13.344579Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.20646","last_updated":"2024-12-30T01:38:14Z","snapshot_observed_at":"2026-08-11T22:46:51.834751Z","submitted_at":"2024-12-30T01:38:14Z","title":"Enhancing Visual Representation for Text-based Person Searching","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-10T23:20:13.344579Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2412.20646"},"observation_digest":"sha256:47d148874d5cdd07dee45e397d78f98f1afee21fc559f065f55dab00db958ca9","observation_id":"f540f493-9fb9-4e7c-8a9b-e4985dc8957f","resolution":{"observed_at":"2026-08-10T23:20:13.344579Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-08-10T22:58:11.484733Z","title":"arXiv preprint arXiv:2101.03036 (2021) 1, 4, 5, 6, 9","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.00318","last_updated":"2024-12-31T07:29:50Z","snapshot_observed_at":"2026-08-10T23:29:15.415229Z","submitted_at":"2024-12-31T07:29:50Z","title":"Improving Text-based Person Search via Part-level Cross-modal Correspondence","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T22:58:11.484733Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2501.00318"},"observation_digest":"sha256:8b4643ba4a3217267cae9a85bf48aadae52f52efb288d5d1d4215d660856837b","observation_id":"728d96e7-4f58-4c87-a710-84fccca0ecb7","resolution":{"observed_at":"2026-08-10T22:58:11.484733Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":"2101.03036","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-07-02T23:17:28.876548Z","title":"arXiv preprint arXiv:2101.03036 (2021) 3, 15","venue":null,"work_id":"6ef6972e-c847-4404-accc-14a0af46c834","year":2021},"citing_paper":{"arxiv_id":"2604.27122","last_updated":"2026-06-27T17:42:07Z","snapshot_observed_at":"2026-08-11T14:22:00.391483Z","submitted_at":"2026-04-29T19:18:36Z","title":"InterPartAbility: Phrase-Region Grounding for Interpretable Text-to-Image Person Re-Identification","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-07T09:29:09.212522Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2604.27122"},"observation_digest":"sha256:29f09f3e921d98cce99bc91a6f77ae2be3f9ffafde167227a996814400a71dec","observation_id":"5a2cbe8c-cce3-41b2-bbcf-70074f481da1","resolution":{"observed_at":"2026-05-12T09:46:26.295511Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":"2101.03036","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-07-02T23:17:28.876548Z","title":"arXiv preprint arXiv:2101.03036 (2021) 3, 15","venue":null,"work_id":"6ef6972e-c847-4404-accc-14a0af46c834","year":2021},"citing_paper":{"arxiv_id":"2604.27122","last_updated":"2026-06-27T17:42:07Z","snapshot_observed_at":"2026-08-11T14:22:00.391483Z","submitted_at":"2026-04-29T19:18:36Z","title":"InterPartAbility: Phrase-Region Grounding for Interpretable Text-to-Image Person Re-Identification","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-01T08:23:22.559471Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2604.27122"},"observation_digest":"sha256:1883efd1a9ac581fedf34a8c137db83fe0117615d1ad0e302f287e2cc5864fb9","observation_id":"252fe246-06c6-4035-8234-5e237682f7d1","resolution":{"observed_at":"2026-07-01T08:25:33.090613Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":"2101.03036","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-07-02T23:17:28.876548Z","title":"arXiv preprint arXiv:2101.03036 (2021) 3, 15","venue":null,"work_id":"6ef6972e-c847-4404-accc-14a0af46c834","year":2021},"citing_paper":{"arxiv_id":"2606.01825","last_updated":"2026-07-01T05:50:06Z","snapshot_observed_at":"2026-08-05T10:25:09.105468Z","submitted_at":"2026-06-01T07:41:44Z","title":"ROGLE: Robust Global-Local Alignment with Automated Region Supervision for Text-Based Person Search","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-06-28T15:18:17.427707Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2606.01825"},"observation_digest":"sha256:76995defaba6993fb7ddec2f4ddd0abc355ed89e15e0a577737e335466604003","observation_id":"f0bd8b06-c1c0-4156-888c-17bfbe9df29f","resolution":{"observed_at":"2026-07-01T22:36:16.994710Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":"2101.03036","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-07-02T23:17:28.876548Z","title":"arXiv preprint arXiv:2101.03036 (2021) 3, 15","venue":null,"work_id":"6ef6972e-c847-4404-accc-14a0af46c834","year":2021},"citing_paper":{"arxiv_id":"2606.01825","last_updated":"2026-07-01T05:50:06Z","snapshot_observed_at":"2026-08-05T10:25:09.105468Z","submitted_at":"2026-06-01T07:41:44Z","title":"ROGLE: Robust Global-Local Alignment with Automated Region Supervision for Text-Based Person Search","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-07-02T23:17:03.456746Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2606.01825"},"observation_digest":"sha256:ed7d19cd5ae897235a67d49f9510dff9bf172dffe5a0a595418443ec0a351a3d","observation_id":"777549a0-2cbe-4680-b9ba-4ed0f4a1727c","resolution":{"observed_at":"2026-07-02T23:17:28.877978Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":"2101.03036","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-07-02T23:17:28.876548Z","title":"arXiv preprint arXiv:2101.03036 (2021) 3, 15","venue":null,"work_id":"6ef6972e-c847-4404-accc-14a0af46c834","year":2021},"citing_paper":{"arxiv_id":"2606.30458","last_updated":"2026-06-29T15:24:03Z","snapshot_observed_at":"2026-08-07T19:25:03.772420Z","submitted_at":"2026-06-29T15:24:03Z","title":"Cross-Resolution Semantic Transfer for Robust Text-to-Image Retrieval in Low-Resolution Surveillance","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-06-30T06:52:54.640706Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2606.30458"},"observation_digest":"sha256:c28fcadd6eafbc466a1986224a0332ef8a6ca9fc02111c347acb44cb9cb4d396","observation_id":"2b11fd3c-fc68-4687-b332-1c4e8605f927","resolution":{"observed_at":"2026-06-30T06:54:20.324306Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-08-02T01:01:00.886331Z","title":"arXiv preprint arXiv:2101.03036 (2021) 52","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.14821","last_updated":"2026-07-16T10:39:22Z","snapshot_observed_at":"2026-08-07T23:46:00.212599Z","submitted_at":"2026-07-16T10:39:22Z","title":"Blurring Modal Boundaries: A Unified Survey from Single- to Multi-Modal Person Re-ldentification","version":1},"reference_index":141,"source":"pdf_text","source_observed_at":"2026-08-02T01:01:00.886331Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2607.14821"},"observation_digest":"sha256:63d276d5e3de40ff901198323a0809e20dbe5a373e351a68867f387e3f59be92","observation_id":"21e104ce-2347-404d-9ffc-4a3091f96528","resolution":{"observed_at":"2026-08-02T01:01:00.886331Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2101.03036","snapshot_observed_at":"2026-08-11T22:37:00.006591Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2608.09152","last_updated":"2026-08-10T05:59:07Z","snapshot_observed_at":"2026-08-12T23:20:35.322407Z","submitted_at":"2026-08-10T05:59:07Z","title":"LightAIR: Lightweight Action Inversion and Riemannian Rectification for Text-based Person Anomaly Search","version":1},"reference_index":152,"source":"pdf_text","source_observed_at":"2026-08-11T22:37:00.006591Z"},"links":{"cited_paper":"/paper/2101.03036","citing_paper":"/paper/2608.09152"},"observation_digest":"sha256:eff297bf4f2a36e4b425a5b2e618457cf56131829d41885ae6a7bab544714999","observation_id":"1b1096cb-b903-4aa9-9c0a-8adb0ae81ea5","resolution":{"observed_at":"2026-08-11T22:37:00.006591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2101.03036/citation-record","integrity":"/paper/2101.03036/integrity","json":"/paper/2101.03036/citation-record.json","paper":"/paper/2101.03036"},"outbound":[],"paper":{"arxiv_id":"2101.03036","last_updated":"2021-01-08T14:30:07Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-09T16:35:45.671134Z","submitted_at":"2021-01-08T14:30:07Z","title":"Contextual Non-Local Alignment over Full-Scale Representation for Text-Based Person Search"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 10 inbound Pith citation observations for arXiv:2101.03036."}