{"as_of":"2026-08-07T08:16:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8a3fc34c52594a4f69e63fd0f360c6cf6f5ea99bc26bcb70cae4df23729b3e04","coverage":[{"denominator":58,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":58,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T13:47:53.840267Z","state":"measured"},{"denominator":59,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":59,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T15:55:21.629729Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.20163","snapshot_observed_at":"2026-08-03T15:55:21.629729Z","title":"Player-centric multimodal prompt generation for large lan- guage model based identity-aware basketball video captioning,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2512.15153","last_updated":"2026-06-04T05:17:07Z","snapshot_observed_at":"2026-08-03T15:55:15.333984Z","submitted_at":"2025-12-17T07:35:03Z","title":"Explainable Action Form Assessment by Exploiting Multimodal Chain-of-Thoughts Reasoning","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-03T15:55:21.629729Z"},"links":{"cited_paper":"/paper/2507.20163","citing_paper":"/paper/2512.15153"},"observation_digest":"sha256:680147c0a9a8eac1eef27b476b95a267e323dc5016b81e1ee6d98ffe71fe74b8","observation_id":"950a8896-23ab-4e88-9383-3528f1d53a28","resolution":{"observed_at":"2026-08-03T15:55:21.629729Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2507.20163/citation-record","integrity":"/paper/2507.20163/integrity","json":"/paper/2507.20163/citation-record.json","paper":"/paper/2507.20163"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-06T13:47:48.338225Z","title":"Gpt-4 technical report","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.338225Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:506e7ee2a9fdedd3331fdea572e0a78d133fa204388e8f2479275c210f832efa","observation_id":"615c7d59-3460-4dbd-ba72-8425a78f764c","resolution":{"observed_at":"2026-08-06T13:47:48.338225Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:03.156038Z","title":"Meteor: An automatic metric for mt evaluation with improved correlation with hu- man judgments","venue":null,"work_id":"b4078688-42e5-4e13-8606-ed900dcc2329","year":2005},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.387302Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:e9356912a597485aa3e5f5417ef89b4de62b0e2368ba09610ff50e41bfa9d3d3","observation_id":"c2c43baf-7352-4256-a8b6-28ee0082f827","resolution":{"observed_at":"2026-08-06T13:48:03.233657Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:02.991379Z","title":"Is space-time attention all you need for video understanding? In ICML, page 4, 2021","venue":null,"work_id":"930b9c77-bff4-4a5d-9804-0563af1ba83e","year":2021},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.471255Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:0f9964fd8b79b27b90f6ad3ae0d5e91fa512458836d8ab6fd4ab9daca123d204","observation_id":"fff18909-4ae3-4d86-8f4e-47e8ba2e2ad9","resolution":{"observed_at":"2026-08-06T13:48:03.098776Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:02.831905Z","title":"Activitynet: A large-scale video benchmark for human activity understanding","venue":null,"work_id":"1a403936-6d76-492d-a73b-352bd6f74c76","year":2015},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.557375Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:edf04c307d8c6875570412690cd90027765e848f4e212e5bb82417d7ca44b6f8","observation_id":"411cb742-a6d4-41c2-babd-432823cd6856","resolution":{"observed_at":"2026-08-06T13:48:02.909167Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:02.671601Z","title":"Quo vadis, action recognition? a new model and the kinetics dataset","venue":null,"work_id":"d14f08d4-512c-4d9e-b2ce-71c24136dda1","year":2017},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.651033Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:f0e26f516909b7b0adfe85c8a6752d512d05bf1317e6d4f2e312dd07276522df","observation_id":"0ec81a08-6cf2-43e0-8e5d-3dadb61c85c3","resolution":{"observed_at":"2026-08-06T13:48:02.760554Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:02.515326Z","title":"Collecting highly paral- lel data for paraphrase evaluation","venue":null,"work_id":"ad5de99d-3e0b-423e-8bb1-4f739e16db56","year":2011},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.689591Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:41093a1710a7ba30ed78f13a889c41a92eaa0df9ed06a4c6ff901e70fc571a9f","observation_id":"66722b3d-12f4-41f1-9a81-a015d7c49748","resolution":{"observed_at":"2026-08-06T13:48:02.593360Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1406.1078","last_updated":"2014-09-03T00:25:02Z","snapshot_observed_at":"2026-07-06T03:45:28.546418Z","submitted_at":"2014-06-03T17:47:08Z","title":"Learning Phrase Representations using RNN Encoder-Decoder for Statistical Machine Translation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1406.1078","snapshot_observed_at":"2026-08-06T13:47:48.721206Z","title":"Learning phrase representations using rnn encoder- decoder for statistical machine translation","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.721206Z"},"links":{"cited_paper":"/paper/1406.1078","citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:c6655f53be3ea17e293efd290bfb85b3ac55910d02e65153a1f31fd92ee56cda","observation_id":"059320fb-9413-43d9-9774-722ab4638539","resolution":{"observed_at":"2026-08-06T13:47:48.721206Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:02.358076Z","title":"Sportsmot: A large multi-object tracking dataset in multiple sports scenes","venue":null,"work_id":"e1cba28c-8dcf-48c2-a7d9-f400d4b0fc47","year":2023},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.739135Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:09188f464fb67d3ab14c9cc92f3d93861d5d8cc3964c9836ca3b7eb68b4fbdfa","observation_id":"e749d5fc-3a1c-493e-90f6-d4cdb1a3f469","resolution":{"observed_at":"2026-08-06T13:48:02.444435Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:02.206843Z","title":"A thou- sand frames in just a few words: Lingual description of videos through latent topics and sparse object stitching","venue":null,"work_id":"2d2c9da9-828e-4eba-8220-0235583c1594","year":2013},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.785487Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:b1e95acd5b15df7f9992edb5d3dafbafb1ac9d1b2abf740419537c6aa258bab8","observation_id":"3ddf3f16-8a17-4b3c-9d57-bac93c3d0a91","resolution":{"observed_at":"2026-08-06T13:48:02.255384Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-06T13:47:48.831059Z","title":"An image is worth 16x16 words: Transformers for im- age recognition at scale","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.831059Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:1bc3ab54426ace8de26001ffd98cf2e35de0bc69cab25845e02217b891453175","observation_id":"1ba5df05-4642-4508-a44d-c8cd9c866ffd","resolution":{"observed_at":"2026-08-06T13:47:48.831059Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:02.065674Z","title":"Soccer captioning: dataset, transformer-based model, and triple-level evaluation","venue":null,"work_id":"436b9b91-2694-4bc1-9c8a-9274df863273","year":2022},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.890391Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:b70c5d53151790eafa2e4d69bd76dbc11c139639c4756c76d84ac743573ed5ec","observation_id":"b0c30544-5258-410e-8fa8-503210534ff5","resolution":{"observed_at":"2026-08-06T13:48:02.121089Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:01.867424Z","title":null,"venue":null,"work_id":"72edb7a8-b284-47db-be9a-73590b322a6c","year":2018},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.942982Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:57ef7d1f8f8287bda2cd52f5a5f53705f8b29d7ea9d94b5ed067dd356546bfa9","observation_id":"28674be7-f982-4455-aefa-ff76dfeaade9","resolution":{"observed_at":"2026-08-06T13:48:01.960339Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:01.701737Z","title":"Deep residual learning for image recognition","venue":null,"work_id":"89f71c22-23bc-421d-a9cb-12ddc3fad6bc","year":2016},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:48.982219Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:52f5e36888fa57512aeb11b32305d11ddbab4bf63d85209fe1b8a38b0a09e5c2","observation_id":"ca1ec4e3-4f73-457f-bbb0-82bcc2a39c01","resolution":{"observed_at":"2026-08-06T13:48:01.789107Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1606.08415","last_updated":"2023-06-06T01:53:32Z","snapshot_observed_at":"2026-07-06T05:01:27.910364Z","submitted_at":"2016-06-27T19:20:40Z","title":"Gaussian Error Linear Units (GELUs)","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.08415","snapshot_observed_at":"2026-08-06T13:47:49.051493Z","title":"Gaussian error linear units (gelus)","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:49.051493Z"},"links":{"cited_paper":"/paper/1606.08415","citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:18e833e5ec6095005bcbbacf8e7483e438be0c87cd5c94b74c6c4951c790a2fb","observation_id":"f9ab540d-86f4-4e49-82d5-b22887b02b66","resolution":{"observed_at":"2026-08-06T13:47:49.051493Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:01.489702Z","title":"Overview of temporal action detection based on deep learning","venue":null,"work_id":"e208e086-a721-4555-b7e5-939c5d0174d0","year":2024},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:49.111293Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:c0ab4f32045e831fe770e1c24b8baaab8313bb47973a0bc9ca8eef90d37d9538","observation_id":"830bdf3f-bf20-47c4-b106-f1922277feca","resolution":{"observed_at":"2026-08-06T13:48:01.578230Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:01.293458Z","title":"Learn- ing to generate move-by-move commentary for chess games from large-scale social forum data","venue":null,"work_id":"e18cfd2d-754b-46bb-abd6-7a914c41c8c6","year":2018},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:49.229047Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:70db6b7c0416ebd3d703a22b56f309686d96ff239a8918eca7ffa70163b2110b","observation_id":"4d3642e8-9f19-4e66-afe3-f34de4e784d0","resolution":{"observed_at":"2026-08-06T13:48:01.371450Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:01.082452Z","title":"Learning seg- ment similarity and alignment in large-scale content based video retrieval","venue":null,"work_id":"0cd8fc17-b15e-41ef-8f5e-7203e54835f8","year":2021},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:49.383887Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:c30e393935f8cdec8d65bb0bcc184262089e1c73861b09ebb4baa992bfb508e9","observation_id":"9b0d2555-6dd6-4586-bb63-a0dd0deded6e","resolution":{"observed_at":"2026-08-06T13:48:01.166288Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:00.949246Z","title":"Automatic baseball commentary generation using deep learning","venue":null,"work_id":"500b47e1-4cbe-4320-a0a9-f94393f7364e","year":2020},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:49.465481Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:d316c44ce9d8187cdc8b2eb0de8c8b62be3a1ece6d7c3dca3ee33303d91aabfd","observation_id":"9163806d-f645-4003-a9b9-4116062688c1","resolution":{"observed_at":"2026-08-06T13:48:01.005079Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1412.6980","last_updated":"2017-01-30T01:27:54Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2014-12-22T13:54:29Z","title":"Adam: A Method for Stochastic Optimization","version":9},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1412.6980","snapshot_observed_at":"2026-08-06T13:47:49.583990Z","title":"Adam: A method for stochastic optimization","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:49.583990Z"},"links":{"cited_paper":"/paper/1412.6980","citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:f890da4bf7d8fc2cfe337a6a939418ca7c23d6c21635d1ca60d08e2d61cf1071","observation_id":"8d357874-4819-4675-b8b5-8f6966fb7123","resolution":{"observed_at":"2026-08-06T13:47:49.583990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:00.783993Z","title":"Video story- telling: Textual summaries for events","venue":null,"work_id":"45f04e0f-20d9-415c-93cb-9a6d1e37e2ff","year":2020},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:49.829969Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:62ef13e6e2f4ef9820d55b0a4f143eba2b6009eb1c37e08b8d02b857dfd73a19","observation_id":"639bf675-0666-4cec-92d7-0ac52890f53c","resolution":{"observed_at":"2026-08-06T13:48:00.869574Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:00.595158Z","title":"Rouge: A package for automatic evaluation of summaries","venue":null,"work_id":"e46c2626-216f-46cc-aa4e-ec4b33273a68","year":2004},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:49.952635Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:3d5171a8e11a4e73426662b5d5ebed3f7fa00b2d1bca3010fc530247a3c2e4d7","observation_id":"3550829e-2298-4d99-8af8-0e6f3c1021da","resolution":{"observed_at":"2026-08-06T13:48:00.680689Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:00.419932Z","title":"Swinbert: End- to-end transformers with sparse attention for video caption- ing","venue":null,"work_id":"c802fc8a-a249-4a4e-92b2-dc36fe8b2f25","year":null},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:50.083025Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:b03907ae5993730a35cec586f6a52297cdd3410351ea11597e963a456a13963a","observation_id":"717f7a0e-aa5a-4888-8221-67085a8a4267","resolution":{"observed_at":"2026-08-06T13:48:00.501771Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:48:00.216121Z","title":"Video swin transformer","venue":null,"work_id":"299dc840-4598-4d2d-82db-86cb0663e533","year":2022},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:50.245337Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:ca2edca2396951693df545afd279f20365139b7b8fa51d66d166779a6a4e60c8","observation_id":"fd23068d-2a81-4dd3-a882-e03f26b28b21","resolution":{"observed_at":"2026-08-06T13:48:00.299165Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:59.985811Z","title":"Llama 3.2 quantized models, 2024","venue":null,"work_id":"11c2be8b-f67e-4ec7-b00d-30bc99392671","year":2024},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:50.372060Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:40506dac3c967adb4971a0472c9b01d299d8f521d96c520f9962fc2db0dd0a2f","observation_id":"bce3b90b-dbca-4d5a-91c8-665ec89a251a","resolution":{"observed_at":"2026-08-06T13:48:00.097823Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:59.744553Z","title":"Soccernet- caption: Dense video captioning for soccer broadcasts com- mentaries","venue":null,"work_id":"a632d33c-c4e3-4edf-8c34-4b4b85e539cc","year":2023},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:50.522585Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:843de5368bf6006254cb10f918199a4c90735e8d38cd96bafd1fd821501c5bfd","observation_id":"ccd780b4-4273-4abd-a016-7267ffcf0f44","resolution":{"observed_at":"2026-08-06T13:47:59.854049Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:59.507445Z","title":"Search-oriented micro-video captioning","venue":null,"work_id":"16115abe-0bba-443c-ac45-5509e279c4a5","year":2022},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:50.650288Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:eca4a32a1abb5c5b99d94707cc4e6126c99ddfb0c0e2747c51eb0190fd4d82ca","observation_id":"0fe48295-4fb5-4b2d-add4-c27f80b8e0b0","resolution":{"observed_at":"2026-08-06T13:47:59.618980Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:59.266742Z","title":"Enhancing visual question answering through question-driven image captions as prompts","venue":null,"work_id":"069c71be-b6c6-4efc-8dcc-6b76834dae05","year":2024},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:50.783054Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:47235330a3b082a0645b5af646beb22b4267f0a8bdc84d386c3b5a8ba3babef0","observation_id":"b68e4fb3-067f-4dec-a121-f022be7d66ce","resolution":{"observed_at":"2026-08-06T13:47:59.373041Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:59.031175Z","title":"Bleu: a method for automatic evaluation of machine translation","venue":null,"work_id":"7417f5e3-1d98-40bf-a55d-60bf5c509507","year":2002},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:51.004697Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:521de1444c7c221b0e0fb6f17e71c6f886bbf5b46fefec511430d83d5bd00b5d","observation_id":"4a3e2810-37bb-4324-8303-65485346603d","resolution":{"observed_at":"2026-08-06T13:47:59.134945Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:58.782596Z","title":"Identity- aware multi-sentence video description","venue":null,"work_id":"531ab6be-30a9-4998-be20-a04aca826f6a","year":2020},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:51.226089Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:97801f43aa80c9d273718b782f005bde0f271a431bc9e583287370a20ce35079","observation_id":"b53fca66-fa12-430d-8a6e-7451aac9892a","resolution":{"observed_at":"2026-08-06T13:47:58.894291Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:58.565892Z","title":"Goal: A challenging knowledge-grounded video captioning benchmark for real- time soccer commentary generation","venue":null,"work_id":"72781168-4527-4eee-905c-639af84b3696","year":2023},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:51.400669Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:ad3e4678eccd972a822a8bee1e7e0cd87aa77789e3e5fa1ca0026c49a03ceb7d","observation_id":"37534dc4-5bcd-4462-9b01-fe8af62f296d","resolution":{"observed_at":"2026-08-06T13:47:58.654739Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:58.383928Z","title":"Sports video captioning via attentive motion representation and group re- lationship modeling","venue":null,"work_id":"6c794133-33d2-4a85-8c55-f82df7110fb7","year":2019},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:51.533250Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:cd5d6d7d3f3676a7d903006e257ecba5eddea0aaf2e99ef877c2ce9c5630bdc4","observation_id":"77af1dcb-5943-4b4e-96b4-73de32742b74","resolution":{"observed_at":"2026-08-06T13:47:58.457349Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:58.186385Z","title":"Language models are unsupervised multitask learners","venue":null,"work_id":"7b6a9b6e-0555-4ffa-ae0d-73087d635b66","year":2019},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:51.721508Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:13bba08aabcb588f68cb2626a4461fc0a38a20904bc73bf362373ad942bfe9e5","observation_id":"f7630cbe-3fa3-4c7e-b011-dcb167702c73","resolution":{"observed_at":"2026-08-06T13:47:58.256585Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:58.025925Z","title":"Learn- ing transferable visual models from natural language super- vision","venue":null,"work_id":"b513dd39-7cd2-4071-8c4a-f5c170af37c2","year":2021},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:51.901836Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:708c880a6fe8d0f80034f83b0dcbcd5347f1e23138c40aa3db7ed40f55e22c36","observation_id":"5fbd5acf-adb7-4f10-9ac1-637ecc7db084","resolution":{"observed_at":"2026-08-06T13:47:58.105707Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.18530","last_updated":"2024-11-18T04:45:19Z","snapshot_observed_at":"2026-08-02T11:20:16.102378Z","submitted_at":"2024-06-26T17:57:25Z","title":"MatchTime: Towards Automatic Soccer Game Commentary Generation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.18530","snapshot_observed_at":"2026-08-06T13:47:52.050904Z","title":"Matchtime: Towards automatic soccer game commentary generation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.050904Z"},"links":{"cited_paper":"/paper/2406.18530","citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:e589227553b588bd7f9fbd671a2d6ebe938091b0edbc2cb778532ca03c734cc7","observation_id":"17fcd46c-ce79-4d66-9dc7-ee5dda44c2ee","resolution":{"observed_at":"2026-08-06T13:47:52.050904Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:57.759396Z","title":"Towards universal soccer video under- standing","venue":null,"work_id":"543e5d56-d9e4-491f-97bc-806af99b6f5b","year":2025},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.195571Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:4bf87d865ab1a895090ca8920cd11f033c0666941a896352833bcbaab5451f77","observation_id":"9ed940e0-1380-4904-adcc-1aaf416794ad","resolution":{"observed_at":"2026-08-06T13:47:57.935512Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:57.616492Z","title":"Grounding action descriptions in videos","venue":null,"work_id":"4598b321-9dbf-456b-9e18-2491d1b6aa82","year":2013},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.288758Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:e1fed44b562081bb8369305ea100e74388ad16cf16d146bfe8138a24cec57f1c","observation_id":"c026b22e-20c4-4f7e-a4c7-94e13b78a84a","resolution":{"observed_at":"2026-08-06T13:47:57.685758Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:57.355443Z","title":"Timechat: A time-sensitive multimodal large language model for long video understanding","venue":null,"work_id":"5d178f2a-a0e7-4830-b839-a55989992554","year":2024},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.384320Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:55a144b29e1e1d4709b11a4935352b510b36dc6b27fc040e4d08f7aae7fa957e","observation_id":"dd8cd761-f1a1-428d-ac77-a9f617db746e","resolution":{"observed_at":"2026-08-06T13:47:57.497245Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:57.202574Z","title":"A dataset for movie description","venue":null,"work_id":"30268fad-733a-4e51-9798-1a86d6c3fad6","year":2015},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.476419Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:781264580201795d99eb7d7ad713d2d79e85d3a42d07639a2258dac486d3dfd0","observation_id":"cccc58cf-1f89-4681-9e80-55edcc540cee","resolution":{"observed_at":"2026-08-06T13:47:57.261849Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:57.004817Z","title":"Accurate and fast compressed video captioning","venue":null,"work_id":"71a0871e-9910-4b3f-8643-6c5be3d93d1d","year":2023},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.570908Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:858681e3a72dbddf6fe44ac71d241751fa2f0a6db8663ce0fbf3118f3036f590","observation_id":"b1a26b85-86e0-4762-87ba-0b6a957a1e6c","resolution":{"observed_at":"2026-08-06T13:47:57.090318Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:56.859281Z","title":"An overview of the tesseract ocr engine","venue":null,"work_id":"1fc1ade8-1c77-41c3-b005-16657fed9952","year":2007},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.629207Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:3237d472e3927efde85b05e4c03f31ff196fc12649d05d50fba16a149245a4c7","observation_id":"358d8369-dea1-4109-938e-9adea259b196","resolution":{"observed_at":"2026-08-06T13:47:56.929106Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:56.711203Z","title":"Clip4caption: Clip for video caption","venue":null,"work_id":"dc3c21f2-2277-455b-8e2b-1d075546a2b5","year":2021},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.681290Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:47925430a3f1088cc2f2eb7c29ae5e83831d54ebf86dcde2ab4931fca531a3fc","observation_id":"f79916d6-9ca0-4e4f-abeb-d43dac9275c1","resolution":{"observed_at":"2026-08-06T13:47:56.762986Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:56.549127Z","title":"Qwen2.5: A party of foundation models, 2024","venue":null,"work_id":"f22b9c6f-b38a-4761-b795-728094767078","year":2024},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.753771Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:17f0712d4ec1170b59dfbeb5ad51915bb9b3ab5ab82a09b5b43b51800e56623f","observation_id":"37b257fc-62a8-4905-8fe4-d1b9b93c9c43","resolution":{"observed_at":"2026-08-06T13:47:56.634122Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-02T11:57:18.735747Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-08-06T13:47:52.792329Z","title":"Llama 2: Open foundation and fine-tuned chat models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.792329Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:ab479a8f0bd2fed63434efef608f96fcc6dc0a5003cb8cf4b2d92f9fdec90b83","observation_id":"6ce10f36-56a2-4295-816e-cb5b157b3675","resolution":{"observed_at":"2026-08-06T13:47:52.792329Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:56.402309Z","title":"Atten- tion is all you need","venue":null,"work_id":"03713372-ad59-4b40-bb9d-47387d3d328e","year":2017},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.859030Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:ea74f422f98409535850f020ce0e8129289c3227ea9d8ee29db12dcf8b61d48d","observation_id":"f50975b2-6527-41d1-b5bb-58210b190ba9","resolution":{"observed_at":"2026-08-06T13:47:56.470095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:56.204036Z","title":"Player tracking and identification in ice hockey.Expert systems with applications, 213:119250, 2023","venue":null,"work_id":"0c7e045b-998a-4af0-8c88-8cd9ed95a460","year":2023},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.919724Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:fababfd1679eb2c11911107d61cb09aab2708c21e00f081c0d15a444a714f311","observation_id":"1c967b64-34d4-4209-935d-c7c83bc22406","resolution":{"observed_at":"2026-08-06T13:47:56.298252Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:56.011839Z","title":"Cider: Consensus-based image description evalua- tion","venue":null,"work_id":"8e1de270-6961-4936-a2ba-fbcbfdc72f5f","year":null},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:52.975748Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:028b49165dbc3d0c061a69a6d1feb353f97d39366cbc1f778d4a23eb352b026e","observation_id":"98ad022a-2a58-4206-80ae-5b22fdb2ef3a","resolution":{"observed_at":"2026-08-06T13:47:56.092821Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:55.879535Z","title":"Omnivid: A generative framework for universal video understanding","venue":null,"work_id":"74ab2c2f-a07b-430a-ac50-0643dc3dd864","year":2024},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.023439Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:f9de025ecfd8c575653046fb94fdef0fd72012cbf1afc940455361f33533d89b","observation_id":"265d17ab-2956-4a69-85f1-c0a1c0511237","resolution":{"observed_at":"2026-08-06T13:47:55.936975Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:55.730408Z","title":"Sports video anal- ysis on large-scale data","venue":null,"work_id":"265577cd-b31e-4438-b04a-11f627f1d53e","year":2022},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.059771Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:65b2608ccaa25d649e5a6e0f6703f45a5a0a204c63637064dedfe80cde27e9dc","observation_id":"371e9520-823e-41f1-aa81-b12c4ac8fd27","resolution":{"observed_at":"2026-08-06T13:47:55.808364Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:55.585828Z","title":"Learning label semantics for weakly supervised group activity recognition","venue":null,"work_id":"f29f5fe0-c374-48b3-b445-2b510b13443b","year":2024},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.106462Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:fe12d0e0d64641a997eca4f9a0ca4a5bd4b403d14544f081dec0b9e4ad90ce68","observation_id":"4e6ee851-9abe-4eab-8126-257f9d148816","resolution":{"observed_at":"2026-08-06T13:47:55.644024Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:55.384085Z","title":"A simple yet effective knowledge guided method for entity-aware video captioning on a basketball benchmark","venue":null,"work_id":"e9847a82-13ff-4422-91a3-912687034f30","year":null},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.202516Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:de9fd82e9dee07c7e88cb738339f5340a6725f515dbe084d26e273f01854722f","observation_id":"d8c8a3ac-9f04-42ef-b5ff-18212b579d07","resolution":{"observed_at":"2026-08-06T13:47:55.465434Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:55.242915Z","title":"Eika: Explicit & im- plicit knowledge-augmented network for entity-aware sports video captioning","venue":null,"work_id":"e7377ff5-6c60-4f8a-8a5d-34d03894193b","year":2025},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.279524Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:039a9a6c1d9c1d7de5f5cbef3c41b5be253f155cd8290be7f1a25a52b0865ae2","observation_id":"c9368bf4-c517-4bea-9e43-b0b69e46abe0","resolution":{"observed_at":"2026-08-06T13:47:55.304502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:55.088688Z","title":"Msr-vtt: A large video description dataset for bridging video and language","venue":null,"work_id":"93c1ffb7-fca8-49c3-ac69-ff80e3bf2609","year":2016},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.363030Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:f0c6381da7642d0a3bbf7dfdc0ab7a7769b0be31c6d9e4dcb0df7addcb256828","observation_id":"93235d0d-b85a-4a05-b083-928108e89e78","resolution":{"observed_at":"2026-08-06T13:47:55.163118Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:54.898061Z","title":"Hierarchical modular network for video captioning","venue":null,"work_id":"fe45b478-d4d8-467b-9ba6-937d7a702ff7","year":2022},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.424019Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:297f18b6b300a0221aca782a7be8936447b1ba885c1464933de1988db45197bf","observation_id":"66f835c6-e6b0-4a1e-9771-1efe8ff04925","resolution":{"observed_at":"2026-08-06T13:47:54.964288Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:54.768054Z","title":"Fine-grained video captioning for sports narrative","venue":null,"work_id":"a734ea6b-6636-4347-b2ba-44f548bf7a20","year":2018},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.530146Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:601695feb350785b9b23856802f3e5c5aba93906c45c92713cb9688e0a2945ac","observation_id":"4357c369-9a65-48c0-9f0f-736c24d4eb78","resolution":{"observed_at":"2026-08-06T13:47:54.829046Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.12140","last_updated":"2023-06-27T11:42:44Z","snapshot_observed_at":"2026-08-04T08:00:42.297781Z","submitted_at":"2023-05-20T08:43:51Z","title":"Movie101: A New Movie Understanding Benchmark","version":2},"cited_work":{"arxiv_id":"2305.12140","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.12140","snapshot_observed_at":"2026-08-06T13:47:53.956828Z","title":"Movie101: A New Movie Understanding Benchmark","venue":"cs.CV","work_id":"88767c7c-55d4-4ec1-aa28-658708d085d9","year":2023},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.617274Z"},"links":{"cited_paper":"/paper/2305.12140","citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:53ac6137671467342d7c3616fbf0e35d24ec6a842e4d3f4eabcda0fb19d2cad6","observation_id":"0fe3ff17-e83f-4ba0-9c7d-568a37f848fb","resolution":{"observed_at":"2026-08-06T13:47:54.028783Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:54.643039Z","title":"Harnessing large language models for training-free video anomaly detection","venue":null,"work_id":"84bb619b-94a4-486b-b8a6-924f37c169d3","year":2024},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.702448Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:2ac00bf3c18f15f8b88ee7cc8424dfd831b4c11dfdc147acd4b5068e26fafa3a","observation_id":"c7834a93-67fc-4937-83c0-de9ac951366d","resolution":{"observed_at":"2026-08-06T13:47:54.711821Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:54.530922Z","title":"A descriptive basketball highlight dataset for automatic commentary gen- eration","venue":null,"work_id":"ac6c8548-9ef3-42fd-9d0f-079abbf8e3a2","year":2024},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.781975Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:23d626f2cfa67dda89a7d92ce71b255064f220de1cb554ac2c253f431e11a03c","observation_id":"ac9f3192-e1d3-4827-9cf4-099253389e48","resolution":{"observed_at":"2026-08-06T13:47:54.589144Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T13:47:54.353753Z","title":"jump ball","venue":null,"work_id":"21f49ee9-6f77-4c51-b25a-20a29988bee4","year":2018},"citing_paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T13:47:53.840267Z"},"links":{"citing_paper":"/paper/2507.20163"},"observation_digest":"sha256:573b66cb37bc42e54f8081b154f5e2b90ab5ac00fc0d3588f8d21db55f658cfd","observation_id":"3d044e4e-301e-400c-90f2-2c0b96895b92","resolution":{"observed_at":"2026-08-06T13:47:54.449519Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2507.20163","last_updated":"2025-07-27T07:30:56Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-06T13:47:46.251861Z","submitted_at":"2025-07-27T07:30:56Z","title":"Player-Centric Multimodal Prompt Generation for Large Language Model Based Identity-Aware Basketball Video Captioning"},"reference_resolution":{"displayed":58,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":8,"verified_exact":1,"verified_fuzzy":49},"total_outbound_references":58},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 58 of 58 outbound references and 1 inbound Pith citation observation for arXiv:2507.20163."}