{"as_of":"2026-08-23T21:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8224bb2232b694bcda5bdcaf0f52e1e1d9e4c4155097d871ed42cbce09fe07ec","coverage":[{"denominator":43,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":43,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T11:40:36.924564Z","state":"measured"},{"denominator":44,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":44,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-23T06:30:58.430688+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-15T07:16:05.225105Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-15T07:19:49.927864Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"cited_work":{"arxiv_id":"2412.15106","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2412.15106","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Knowing where to focus: Attention- guided alignment for text-based person search","venue":null,"work_id":"24e30eba-f65b-4ec2-b27a-53cb66356752","year":2024},"citing_paper":{"arxiv_id":"2605.01725","last_updated":"2026-05-14T10:16:56Z","snapshot_observed_at":"2026-08-16T23:20:27.163699Z","submitted_at":"2026-05-03T05:49:27Z","title":"Motion-Aware Caching for Efficient Autoregressive Video Generation","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-05-15T07:16:05.225105Z"},"links":{"cited_paper":"/paper/2412.15106","citing_paper":"/paper/2605.01725"},"observation_digest":"sha256:09644043e3d25755e1be0eea5c5cf1b7a2f8fe16a89e99beaed453be55e4b2a7","observation_id":"ebff62f9-0c04-4fb8-93fa-d0b166fccf48","resolution":{"observed_at":"2026-05-15T07:19:49.931364Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2412.15106/citation-record","integrity":"/paper/2412.15106/integrity","json":"/paper/2412.15106/citation-record.json","paper":"/paper/2412.15106"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:36.754218Z","title":"Tran- sreid: Transformer-based object re-identification,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.754218Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:f5cfb162c837c84df055decc4e5964f431f7ef269f55cd2324d373d8af0e3cd2","observation_id":"167144a4-d441-459d-84c6-f8c52a9bca8a","resolution":{"observed_at":"2026-08-11T11:40:36.754218Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.16684","last_updated":"2024-08-29T16:31:05Z","snapshot_observed_at":"2026-08-16T13:22:51.493742Z","submitted_at":"2024-08-29T16:31:05Z","title":"PartFormer: Awakening Latent Diverse Representation from Vision Transformer for Object Re-Identification","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.16684","snapshot_observed_at":"2026-08-11T11:40:36.759255Z","title":"Partformer: Awakening latent diverse representation from vision transformer for object re-identification,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.759255Z"},"links":{"cited_paper":"/paper/2408.16684","citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:bdaeab5ebbd4e79f3ba153a9ab4aea99abc63d635ee3d17056340f0bc3b1d38b","observation_id":"b1e3d541-5f07-418d-9d14-566c3d50d2ad","resolution":{"observed_at":"2026-08-11T11:40:36.759255Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.457271Z","title":"Occluded person re-identification via saliency-guided patch transfer,","venue":null,"work_id":"9402d30f-db11-4ef8-b30e-960dc20c3c2a","year":2024},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.764032Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:9fe98ec0420ab7563195a4f6a5789211f652855d636003787cfbb119f1c34c43","observation_id":"5bee19bb-63ae-4869-b4f7-bc33e319a731","resolution":{"observed_at":"2026-08-11T11:40:37.461173Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.445025Z","title":"Body part-level domain alignment for domain-adaptive person re-identification with transformer framework,","venue":null,"work_id":"1ae31369-939e-4eb3-b030-ce7bf9d0d3d2","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.768773Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:8a63147da5fa1cc69bea81d33528e63d1108e27d42cdf75d09cb2433ea6e1b6f","observation_id":"7540cced-cccb-4b81-bfb8-cd636a21fcf9","resolution":{"observed_at":"2026-08-11T11:40:37.449640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.433736Z","title":"Dynamic prototype mask for occluded person re-identification,","venue":null,"work_id":"fa0a098e-dd1c-4f68-ba8a-6c3ee8367008","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.772928Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:e0f013d5f04d472e4389b159ad4ed09111c54a9a7d3eb26fdf21c7e6d22bcd6f","observation_id":"29d15733-ecf3-4488-a9a5-86aacc2f127e","resolution":{"observed_at":"2026-08-11T11:40:37.437642Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.422803Z","title":"Person re-identification with hierarchical discriminative spatial aggregation,","venue":null,"work_id":"b5dfdfb1-e094-45de-8f06-3078a45e5878","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.777054Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:513ebef2a4a1aad1e1a541e2fb7de6b7d163d7b5e4ce4ef810f7f36add2909c8","observation_id":"3158068d-ccd7-4dfa-ab94-00cc371008a4","resolution":{"observed_at":"2026-08-11T11:40:37.426346Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.411281Z","title":"Dual-adversarial representation disentanglement for visible infrared person re-identification,","venue":null,"work_id":"cb5804a1-d61c-47d3-993d-da63abe84625","year":2023},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.781342Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:3d4dfc2be03b1952c4a25728e640a8eab85c53175008db34f4d968fd1466ce65","observation_id":"5f0a6a0f-ffb9-4474-b361-cf8526515808","resolution":{"observed_at":"2026-08-11T11:40:37.415614Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.398402Z","title":"Multi-view evolutionary training for unsupervised domain adaptive person re-identification,","venue":null,"work_id":"cd5ca5a3-08df-4719-ba40-3bd9cdee1af3","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.785075Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:3dcacde08dd684b7ed4b803d26380322a2a4e8d048b3454638983c9c9ca14b7c","observation_id":"5fac8d44-f8b6-4c42-8abd-ae18e66d972a","resolution":{"observed_at":"2026-08-11T11:40:37.403015Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.385762Z","title":"Self-supervised recovery and guide for low-resolution person re-identification,","venue":null,"work_id":"6296059b-3409-4164-a0ee-e234414a394f","year":2024},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.788711Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:9e5576bf269bf27f93cfc8a1c3a67b2d57a62c77323fe2f7a3254529b4843ac2","observation_id":"d57b98fb-2945-4849-a409-1a02dee824c8","resolution":{"observed_at":"2026-08-11T11:40:37.389832Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.372849Z","title":"Region generation and assessment network for occluded person re- identification,","venue":null,"work_id":"8a24feb0-17a3-4237-a13c-14e02fcff546","year":2023},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.792533Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:b0c0969dc6c64320ae43c76faf2e3d06d62fe037f446f8233a92c59c45cb939f","observation_id":"f5f5a640-2a1d-42db-be5c-640fe2c47c25","resolution":{"observed_at":"2026-08-11T11:40:37.377393Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.359814Z","title":"Text-based person search with limited data,","venue":null,"work_id":"1a55a43d-b1a1-44b5-b545-b12ca1c82aa1","year":2021},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.796361Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:4b138f5bdbdb9125c2479cb0349dcce01b8c87f4d81b12d802df9e9da6faea17","observation_id":"49cb6300-2a5b-4956-885e-860b16b0a6a6","resolution":{"observed_at":"2026-08-11T11:40:37.364302Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.347187Z","title":"Cross-modal implicit relation reasoning and align- ing for text-to-image person retrieval,","venue":null,"work_id":"5cfbd0ad-a6a2-453c-8f8e-d1b8b1d6a4e3","year":2023},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.800359Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:e19ba06814398813227ba140ed9cd28b68db6d8e2d3d9c9101372935c9a8cbc6","observation_id":"6eedd4d4-fa2b-458a-9f0c-d3b27bd92842","resolution":{"observed_at":"2026-08-11T11:40:37.351632Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.333772Z","title":"Rasa: relation and sensitivity aware representation learning for text- based person search,","venue":null,"work_id":"95e2f5aa-8e39-4c7b-8113-9c03be5f2ea2","year":2023},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.804510Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:64d53c16d39511b4f96c1311425480cfb628db055d126203c1ae64dc7bd2bebf","observation_id":"e3b4f96d-5c12-430d-a54f-2f92bd0e4e57","resolution":{"observed_at":"2026-08-11T11:40:37.338611Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:36.808120Z","title":"Learning transferable visual models from natural language supervision,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.808120Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:31cad3a353f722fd4a49245db6df95eb526667b38c71a61ec79d4913a250933d","observation_id":"950d5460-9eb0-4ad9-a470-a4685bf2ad70","resolution":{"observed_at":"2026-08-11T11:40:36.808120Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:36.812322Z","title":"Align before fuse: Vision and language representation learning with momentum distillation,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.812322Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:4f04656d1aba3f8066bc72f302a7c4f10391b19e7ee22e009e50dd26daa43a37","observation_id":"518b52cf-6af6-4549-894b-35df2012fc05","resolution":{"observed_at":"2026-08-11T11:40:36.812322Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.306316Z","title":"Towards unified text-based person retrieval: A large-scale multi-attribute and lan- guage search benchmark,","venue":null,"work_id":"712d7e8d-891b-44ca-a81a-fdf017093a04","year":2023},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.816144Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:c298ab751123f1d81d9e3fa30e3ceb0890b2b63a91f5b5cfdd9ecda42f042724","observation_id":"f10edc94-8380-4c58-aac7-7ed1619b046d","resolution":{"observed_at":"2026-08-11T11:40:37.310409Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.294170Z","title":"Person search with natural language description,","venue":null,"work_id":"4d14fc67-e0e8-4796-a3ff-f0bd0a8a81de","year":2017},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.820116Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:e4ffe58a6911c1c8f297d17f105a3a167592552d93ab142d346f00cfb59e0a7c","observation_id":"b2f0d1de-0e1d-4fad-b6e5-637fe52970fe","resolution":{"observed_at":"2026-08-11T11:40:37.298627Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.281688Z","title":"Deep cross-modal projection learning for image- text matching,","venue":null,"work_id":"f950dcfd-4555-4626-ada7-43a2dde96471","year":2018},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.824016Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:caa619d15ebcbe7c17966ab816e7339a37992551732f4f3b0c8802bbce9d459d","observation_id":"17f065f5-5145-4672-b9c0-a527c2b045a0","resolution":{"observed_at":"2026-08-11T11:40:37.286162Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2107.12666","last_updated":"2021-08-09T02:21:14Z","snapshot_observed_at":"2026-08-16T18:07:16.925424Z","submitted_at":"2021-07-27T08:26:47Z","title":"Semantically Self-Aligned Network for Text-to-Image Part-aware Person Re-identification","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2107.12666","snapshot_observed_at":"2026-08-11T11:40:36.827885Z","title":"Semantically self-aligned network for text-to-image part-aware person re-identification,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.827885Z"},"links":{"cited_paper":"/paper/2107.12666","citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:30e7233c9aee8736672f615cd9738cfe6a294e37bc96b1b2a0f014b07296cec3","observation_id":"021bc82b-c750-46d9-8099-3c557305c9e3","resolution":{"observed_at":"2026-08-11T11:40:36.827885Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.269618Z","title":null,"venue":null,"work_id":"3c867709-3a26-40b0-bc8d-cc5373a0c669","year":2016},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.831802Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:2ce6b832421e228b8e09cf40c9d23e6182b37f17a90e666b56f99420a23725de","observation_id":"b2b0de81-accc-430c-b4a6-fc5802f5b6d2","resolution":{"observed_at":"2026-08-11T11:40:37.273429Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1810.04805","last_updated":"2019-05-24T20:37:26Z","snapshot_observed_at":"2026-08-14T18:16:28.847993Z","submitted_at":"2018-10-11T00:50:01Z","title":"BERT: Pre-training of Deep Bidirectional Transformers for Language Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1810.04805","snapshot_observed_at":"2026-08-11T11:40:36.835967Z","title":"Bert: Pre-training of deep bidirectional transformers for language understanding,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.835967Z"},"links":{"cited_paper":"/paper/1810.04805","citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:2ab247765aef77433040cede8d241209a4fa7a42413f526a120fb763ce1c635c","observation_id":"b001d6c2-f58b-4f81-9cf8-a5c064edcd77","resolution":{"observed_at":"2026-08-11T11:40:36.835967Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.257355Z","title":"Asymmetric cross- scale alignment for text-based person search,","venue":null,"work_id":"4ea00a16-41c2-4973-8b53-301c9bd7cf94","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.839939Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:09867c1d3e91fe787ecf9da5d020e0185c4f9035eb9fecec6adc58d968833eae","observation_id":"c78e091f-dd7f-4108-8d75-32ca16b7f9a7","resolution":{"observed_at":"2026-08-11T11:40:37.261639Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.245022Z","title":"Clip-driven fine-grained text- image person re-identification,","venue":null,"work_id":"39750020-08a1-498f-a081-35281708bd9c","year":2023},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.843673Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:19af024b654d80f194a0e6dc9a55cf63bbe908e55d08a44312d7934ed148e0eb","observation_id":"0617108e-b078-486e-b7d3-6035339aa99b","resolution":{"observed_at":"2026-08-11T11:40:37.249530Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.231237Z","title":"Bilma: Bidirectional local-matching for text-based person re-identification,","venue":null,"work_id":"40b712d5-75bc-401d-ac35-3e980f42979e","year":2023},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.847799Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:d30f8c1c277ffa269f697df68922fcb8f966ab96ae03050e1d08e790059c057c","observation_id":"ceeca4b8-f22c-4ee4-9c8f-510f77a062d1","resolution":{"observed_at":"2026-08-11T11:40:37.235478Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.219021Z","title":"Harnessing the power of mllms for transferable text-to-image person reid,","venue":null,"work_id":"4c6cf423-811b-40d6-8493-ebfa67df9802","year":2024},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.852085Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:e3a1090d2a68d029b4115234882311f44433b067768cf224d7f73b64b76121c5","observation_id":"58e60293-2a15-4db5-9eb6-eb9bd8ea1bd8","resolution":{"observed_at":"2026-08-11T11:40:37.223293Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.205853Z","title":"Plot: Text-based person search with part slot attention for corresponding part discovery,","venue":null,"work_id":"a2d33881-558e-41aa-a178-a608baf7b4a0","year":2024},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.856279Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:1ff1df1bbf2738b31c4bea1b3427367a19877f74bc8a5506702cedbe67a590e1","observation_id":"c8881539-a0da-41c2-8edf-ca513983ebcc","resolution":{"observed_at":"2026-08-11T11:40:37.210160Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.193524Z","title":"“cloze procedure","venue":null,"work_id":"f711584c-df7c-4d11-9852-8949d625751d","year":1953},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.860190Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:3843600ecced3a16d3446086028a1828b8bf7e845ebaabb3848b1745769635ea","observation_id":"18163ba2-c622-4cc3-b900-cac9f41ed4de","resolution":{"observed_at":"2026-08-11T11:40:37.197493Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.181625Z","title":"Dual-path convolutional image-text embeddings with instance loss,","venue":null,"work_id":"efe61501-9125-40a5-bb79-b01327ca8116","year":2020},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.863992Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:10c2f668713590f92fe0e300417f03d2f2dc2d2f69751b0366f74e50fd0c0ad9","observation_id":"060a8ae8-3c25-4cb0-9afb-0204e1a2e2af","resolution":{"observed_at":"2026-08-11T11:40:37.185945Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.170021Z","title":"Vitaa: Visual-textual attributes alignment in person search by natural language,","venue":null,"work_id":"59ea97a1-601c-4265-828a-2aa137a9e8ef","year":2020},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.871539Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:e072b2a92557faf3cb293694d45fe580646060a61ed7fc98f82a598a20f41fb2","observation_id":"19d893b8-e008-4bb2-90f7-bcca17785fe6","resolution":{"observed_at":"2026-08-11T11:40:37.174029Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.158381Z","title":"Dssl: Deep surroundings-person separation learning for text- based person retrieval,","venue":null,"work_id":"dab1fa0a-f903-4ba0-aeaf-6655e16fe01e","year":2021},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.875636Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:89204fbffbd817b26e0b7ec68b13d61a9dac72d8747d539ed866dc3d5a02413b","observation_id":"870736c5-e464-4506-b48d-6a4a916dba84","resolution":{"observed_at":"2026-08-11T11:40:37.162628Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.145745Z","title":"Text-based person search via multi-granularity embedding learning","venue":null,"work_id":"2fbf748d-7901-44bb-ad9a-33dbd347f312","year":2021},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.879255Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:79e745a8d3572ce7b6248532602b15a481eacc7768d3110337647f7eb6ac7f95","observation_id":"ded86603-44c5-4645-8374-7280a06eee96","resolution":{"observed_at":"2026-08-11T11:40:37.150942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.132015Z","title":"Learning semantic-aligned feature representation for text-based person search,","venue":null,"work_id":"2ab808f2-293d-4346-a898-1a66b1e5d683","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.882916Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:dd89cf2b474fe4cffba9bd1fde6913c0cd7de8aed2671a83639edc575ce6849c","observation_id":"ecdb0fd2-89f5-474f-bca7-9eeebacf9971","resolution":{"observed_at":"2026-08-11T11:40:37.137255Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.116747Z","title":"Tipcb: A simple but effective part-based convolutional baseline for text-based person search,","venue":null,"work_id":"ea2d5fcf-68ce-44fc-9917-44c220798f32","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.886665Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:250e40d5edbf2ac0b6a03513d5afb671782c62a5297214c9822b4388b4dd1753","observation_id":"2c38dacd-30e4-4056-adeb-2a8fe6ddc1f4","resolution":{"observed_at":"2026-08-11T11:40:37.121605Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.103388Z","title":"Caibc: Capturing all-round information beyond color for text-based person retrieval,","venue":null,"work_id":"580452a4-0fdf-4f70-a7ac-eb8dc6dda53a","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.890209Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:088aeb2bbfd62fa5be221f00c904cc18f02cea2b742599c787e53cffa5a4fc2f","observation_id":"a6d38b71-c366-4666-b541-90235f75614d","resolution":{"observed_at":"2026-08-11T11:40:37.107666Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.090142Z","title":"Cross-modal co-occurrence attributes alignments for person search by language,","venue":null,"work_id":"63a430c3-1e2b-4be6-8adf-ef1d5d2b7969","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.893976Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:1b40a7bfc6f926dcf681d1f62bee7ec95944f733c8c4d190a17a30293ee2fb0c","observation_id":"04e2d463-0d73-427c-9892-8a5fd6feb2fd","resolution":{"observed_at":"2026-08-11T11:40:37.094475Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.078244Z","title":"Learn- ing granularity-unified representations for text-to-image person re- identification,","venue":null,"work_id":"9598ce4e-ead9-46fc-8f2e-c207499c0341","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.898384Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:af844854228cf28291013e16f208c6e027b82c84ce2c4c13e9669978164838e7","observation_id":"9351ec98-9540-46c2-8dc7-8eba3a930006","resolution":{"observed_at":"2026-08-11T11:40:37.082241Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.063922Z","title":"See finer, see more: Implicit modality alignment for text- based person retrieval,","venue":null,"work_id":"d8d2b953-142f-449a-9f5e-bfa98524730c","year":2022},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.902236Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:51df56020c5dbbc3dcf355248d7a81bc33d67c8680d01b4865dc334abb624b01","observation_id":"ddd867dc-0739-4064-b71f-aa14a3d3a9eb","resolution":{"observed_at":"2026-08-11T11:40:37.068495Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.050248Z","title":"Fine-grained semantic alignment with transferred person-sam for text-based person retrieval,","venue":null,"work_id":"f365f47f-3956-4c9e-af8d-2f35b5330603","year":2024},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.905978Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:97b5df7600d76a88045e008d4996b852a811a2b06a75efe419290108a9095785","observation_id":"47187c3b-22e3-47e7-94ca-2d15c92af268","resolution":{"observed_at":"2026-08-11T11:40:37.054674Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.037279Z","title":"Causality-inspired invariant representation learning for text-based person retrieval,","venue":null,"work_id":"9ec70c8a-80c4-4b11-857c-6da35eb4ea4c","year":2024},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.909747Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:0af98753ce08fe7763a585b3964621e2d78ed8c8fab58cce12b7abee896741aa","observation_id":"d1a380c1-9df1-4d44-8000-df7afdfc6494","resolution":{"observed_at":"2026-08-11T11:40:37.042003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.024016Z","title":"Diverse person: Customize your own dataset for text-based person search,","venue":null,"work_id":"4f00a262-1a38-4149-825a-ef8facf9e2dc","year":2024},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.913660Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:9c50c6e6c78ecf4c17cf09127bf52ffe10d411af49d0fe659aff41662f044415","observation_id":"5646e22e-797c-4196-a0cb-3322c73a915f","resolution":{"observed_at":"2026-08-11T11:40:37.028523Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:37.011199Z","title":"Unifying multi-modal uncertainty modeling and semantic alignment for text-to-image person re-identification,","venue":null,"work_id":"4532aaf3-f4c0-4568-a29b-52d99af6bdee","year":2024},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.917413Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:45166377a345654022d99c085928e6c1e4e0447f88d4a8bc613cd502bcbecb1a","observation_id":"aa3abf2b-a73a-4659-9d19-9799f7993ee8","resolution":{"observed_at":"2026-08-11T11:40:37.015615Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:36.997090Z","title":"Person transfer gan to bridge domain gap for person re-identification,","venue":null,"work_id":"093520d2-24cc-43cd-8128-66b33cc995d3","year":2018},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.920923Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:6b249b6f5b3b76da1e4ec7d77ec1f141d213c5a178e7eb77497a756d29a43b66","observation_id":"57513795-ad5a-4bc5-9197-e252dbb6615a","resolution":{"observed_at":"2026-08-11T11:40:37.002927Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-23T06:30:58.430688+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T11:40:36.924564Z","title":"Distributed representations of words and phrases and their composi- tionality,","venue":null,"work_id":null,"year":2013},"citing_paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-11T11:40:36.924564Z"},"links":{"citing_paper":"/paper/2412.15106"},"observation_digest":"sha256:a221c7dec8b899e9cb233a6af8f96909bacab05365fd89fc2ce855f0f30cfafd","observation_id":"963c600b-3b79-488f-adc0-706344b144ec","resolution":{"observed_at":"2026-08-11T11:40:36.924564Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2412.15106","last_updated":"2024-12-19T17:51:49Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-16T14:55:29.756133Z","submitted_at":"2024-12-19T17:51:49Z","title":"Knowing Where to Focus: Attention-Guided Alignment for Text-based Person Search"},"reference_resolution":{"displayed":43,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":8,"verified_exact":0,"verified_fuzzy":35},"total_outbound_references":43},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-23T06:30:58.430688+00:00","source":"crossref"},{"observed_at":"2026-08-23T06:30:53.778098+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 43 of 43 outbound references and 1 inbound Pith citation observation for arXiv:2412.15106."}