{"as_of":"2026-08-08T01:16:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:69ff22584b6c84c9c738f7f11358dcac8163927f109248348b3ef2860f28a44e","coverage":[{"denominator":65,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":65,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:11:41.287077Z","state":"measured"},{"denominator":66,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":66,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-28T18:44:39.497333Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-28T20:22:37.768955Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"cited_work":{"arxiv_id":"2506.05395","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.05395","snapshot_observed_at":"2026-06-28T20:22:37.768955Z","title":null,"venue":null,"work_id":"a2348d7b-2cc8-4f45-98c4-6ea6d1ea5c57","year":2025},"citing_paper":{"arxiv_id":"2606.00664","last_updated":"2026-05-30T10:41:34Z","snapshot_observed_at":"2026-08-01T15:41:58.795867Z","submitted_at":"2026-05-30T10:41:34Z","title":"SKIP: Sparse Keyframe Interpolation Paradigm for Efficient Embodied World Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-06-28T18:44:39.497333Z"},"links":{"cited_paper":"/paper/2506.05395","citing_paper":"/paper/2606.00664"},"observation_digest":"sha256:8fc26003272b90e9c58e185f7bacf0d1a272b1327ec8069107896ce108be8113","observation_id":"621b29a9-218b-469c-badc-e48d9b208b65","resolution":{"observed_at":"2026-06-28T20:22:37.770265Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.05395/citation-record","integrity":"/paper/2506.05395/integrity","json":"/paper/2506.05395/citation-record.json","paper":"/paper/2506.05395"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:49.436097Z","title":null,"venue":null,"work_id":"330a80b8-2a66-4b60-b42e-da748ec7702e","year":2015},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.824919Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:1ed4ece679c5b3f1b93b8c438982a057020372a7612b4bd960a0ccefb01adfca","observation_id":"99cfbce8-9a52-4ee3-b98f-0d6468c0b2d1","resolution":{"observed_at":"2026-08-07T11:11:49.556238Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:49.262150Z","title":null,"venue":null,"work_id":"0044161d-80b4-42c8-b1fd-ee8309a77e1d","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.833719Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:a1b614508007c8a00b03fb7953ac81dc5d4462b0a882dcfec253efd6a7352a04","observation_id":"9ee347f2-02ef-4c0c-bfb9-d23023592a8e","resolution":{"observed_at":"2026-08-07T11:11:49.348075Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:49.009072Z","title":null,"venue":null,"work_id":"1360d8ee-7c47-41e6-91b0-706ac34e66f2","year":2018},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.839853Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:e3d4b5087a016fcc1491b54afdbe8ad548eab56c16ed2dde2158c46b2f39c587","observation_id":"cf4534ea-7e88-4c0a-924f-99c9d1654079","resolution":{"observed_at":"2026-08-07T11:11:49.123084Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.869854Z","title":null,"venue":null,"work_id":"a956c8ad-064e-4a0f-98d8-0cc6b6871e8b","year":2018},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.849209Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:5402096d8ef3455fcc9e9cb57a2e9c726c349d5f9b87d555c908e005a292202c","observation_id":"21f9572d-1fe8-4544-8618-eb60246b9a91","resolution":{"observed_at":"2026-08-07T11:11:48.927052Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.750922Z","title":null,"venue":null,"work_id":"6e91c9c1-2c74-4d36-8c80-57602640976a","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.857381Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:cf5718e9480491ab2488d0400f8c020e955823a5675dd8213e97a5455f2c9c91","observation_id":"23853136-1253-44bf-a870-4a4328f48c3e","resolution":{"observed_at":"2026-08-07T11:11:48.804137Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.11249","last_updated":"2021-01-27T08:13:19Z","snapshot_observed_at":"2026-07-06T10:36:00.675514Z","submitted_at":"2021-01-27T08:13:19Z","title":"Efficient Video Summarization Framework using EEG and Eye-tracking Signals","version":1},"cited_work":{"arxiv_id":"2101.11249","doi":null,"metadata_source":"pith","pith_arxiv_id":"2101.11249","snapshot_observed_at":"2026-08-07T11:11:42.054602Z","title":"Efficient Video Summarization Framework using EEG and Eye-tracking Signals","venue":"cs.CV","work_id":"a480bbc7-53ee-4c19-87f6-499e285855c8","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.863251Z"},"links":{"cited_paper":"/paper/2101.11249","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:bbc5080b1f3adcca9565896f569fb50e039d0f3735d168dcef45a1d2e1241f8d","observation_id":"65daa8fb-67f8-409c-a318-e459437439ea","resolution":{"observed_at":"2026-08-07T11:11:42.135561Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.631977Z","title":null,"venue":null,"work_id":"c6c302fd-49f4-4783-b1cc-7b94df755c75","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.873911Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:8c6b34c5705bfec9f8014017b4331c13c84a2e5a98beb5d9fc719bc84e54fd85","observation_id":"01a20a45-f71b-4236-ae2b-e29efe7e3020","resolution":{"observed_at":"2026-08-07T11:11:48.689154Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.398995Z","title":null,"venue":null,"work_id":"8f4013f5-7771-4198-a527-b0113b94a09a","year":2013},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.891595Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:8f59e29e8b326d315cd12e7d16a2d3e2a1c01079dbb30ac89cff8b2ba613f42b","observation_id":"50138b50-0413-4747-a6a3-dd5c8a7d1bfa","resolution":{"observed_at":"2026-08-07T11:11:48.452841Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.12971","last_updated":"2023-10-19T17:59:01Z","snapshot_observed_at":"2026-07-06T16:35:48.413898Z","submitted_at":"2023-10-19T17:59:01Z","title":"CLAIR: Evaluating Image Captions with Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.12971","snapshot_observed_at":"2026-08-07T11:11:40.898559Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.898559Z"},"links":{"cited_paper":"/paper/2310.12971","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:caecdc7364f77c21470619520079217ee51a4dbdc4c9d2571562e80caa363dba","observation_id":"ce457575-69a5-422c-ac8f-0596989bb77d","resolution":{"observed_at":"2026-08-07T11:11:40.898559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13292","last_updated":"2023-05-23T07:48:15Z","snapshot_observed_at":"2026-07-06T15:30:49.735343Z","submitted_at":"2023-05-22T17:51:22Z","title":"VideoLLM: Modeling Video Sequence with Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13292","snapshot_observed_at":"2026-08-07T11:11:40.910714Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.910714Z"},"links":{"cited_paper":"/paper/2305.13292","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:698bbc2ea243341539375d8862892792bdebcd4ca8096f809da706644663ab77","observation_id":"0a03daa1-2d5c-431d-b20d-263c1ca55de5","resolution":{"observed_at":"2026-08-07T11:11:40.910714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.259155Z","title":null,"venue":null,"work_id":"8c71ec38-6e6c-4657-b928-d226a93ce2cf","year":2011},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.923028Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:83b44b4def4f4875ce4db6ae8e89f4592d76965ed5d0a6d98487e6d682a1d2cc","observation_id":"ed2f5c9c-9a40-40b6-bcef-1782e02f3fcf","resolution":{"observed_at":"2026-08-07T11:11:48.323123Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1109/cvpr.2009.5","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:41.333474Z","title":null,"venue":null,"work_id":"a208e36f-32cc-4de8-a605-570838edf9ed","year":2009},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.932393Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:f605011c0e990ea8061e3903c1f3f0d75c865d60a58c30084aa641f86f32b7b8","observation_id":"b1c107cb-d8fe-488f-8b18-2d7691655d9c","resolution":{"observed_at":"2026-08-07T11:11:41.340081Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.163556Z","title":null,"venue":null,"work_id":"f9e6bf9f-17b5-4f95-a278-3fd679703202","year":2003},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.941825Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:d1140376e77bcf6eb5f79769ce200645fa803b9d391220411e40e4d6e03705ec","observation_id":"6f67d612-d5c6-4650-afa0-46e7ca96db70","resolution":{"observed_at":"2026-08-07T11:11:48.195996Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.002797Z","title":null,"venue":null,"work_id":"62fd661c-54a6-4aba-a2ec-01dca04d3ee1","year":1996},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.947137Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:db164902555e8fa6a08e4bd2e38c1c548e52299473adb76001ee3293a95f07f6","observation_id":"37c0b276-d7cc-4f1d-86d6-0234d0a0287b","resolution":{"observed_at":"2026-08-07T11:11:48.075632Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.877543Z","title":null,"venue":null,"work_id":"f6dee9a9-0b7c-41fe-901b-ae30b56dbeba","year":2025},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.956872Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:22177eea5f4e84a095af1fb40647a85bace52d7d0992383ecb285441da2fb1c4","observation_id":"32512852-83f0-4ab9-803c-f808fb0263ae","resolution":{"observed_at":"2026-08-07T11:11:47.933065Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.10173","last_updated":"2023-07-14T16:49:42Z","snapshot_observed_at":"2026-07-06T15:05:05.028097Z","submitted_at":"2023-02-15T19:09:34Z","title":"VideoSum: A Python Library for Surgical Video Summarization","version":2},"cited_work":{"arxiv_id":"2303.10173","doi":null,"metadata_source":"pith","pith_arxiv_id":"2303.10173","snapshot_observed_at":"2026-08-07T11:11:41.867397Z","title":"VideoSum: A Python Library for Surgical Video Summarization","venue":"eess.IV","work_id":"132f6468-5e11-40fe-99d4-70309686b9c4","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.962560Z"},"links":{"cited_paper":"/paper/2303.10173","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:d73358d0f94a2ebd1a4dfd1c98fc7ff6fd9b8292616e4e69c64e34313625d616","observation_id":"9a044658-5cd4-46dd-b437-70a795eb662b","resolution":{"observed_at":"2026-08-07T11:11:41.938502Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.773668Z","title":null,"venue":null,"work_id":"08edff2a-4063-4608-a513-a28ab5ff8b9b","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.968488Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:21aa613b8b04d90bbe9baa980c5e5da878f674d7913b644866d296b4bd7bae0a","observation_id":"da12a42d-f8de-4ab4-9d7f-f445017b3556","resolution":{"observed_at":"2026-08-07T11:11:47.819914Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:40.983011Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.983011Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:eab24258b84a913bd430502403d02d000e467648646c4c2e01c07619e21aa792","observation_id":"d3433d7d-ff72-4c34-9199-d63cea6d9af1","resolution":{"observed_at":"2026-08-07T11:11:40.983011Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.454261Z","title":null,"venue":null,"work_id":"24930598-93c2-422b-87b5-5e9218ecf4a5","year":2016},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.001921Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:b54851834972fec1ab739aa1eb8ed7387140298e23c873c417c3e1a681a59072","observation_id":"46189abc-c034-4146-9748-9d6ee98d558e","resolution":{"observed_at":"2026-08-07T11:11:47.548231Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.271425Z","title":null,"venue":null,"work_id":"91665530-c526-4dd4-88aa-5ae4a353e79b","year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.009725Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:ebc83319855c24321977b4f0e88c8985e5276cc803c93e9948a229b6c5c267b8","observation_id":"5011cb4f-a5a4-4150-aa5d-609ec90f74ce","resolution":{"observed_at":"2026-08-07T11:11:47.353245Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.106798Z","title":null,"venue":null,"work_id":"0bffc55e-b65c-4072-ba15-d755ba99f141","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.019761Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:a3e8e40381a7ae0800ae2df1fe1ff04a472de890bc7934ce0e6dedfaa9d2d59a","observation_id":"0d72a24c-1815-48ca-b5f9-299d474ad21e","resolution":{"observed_at":"2026-08-07T11:11:47.180318Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.943192Z","title":null,"venue":null,"work_id":"10233f15-f129-4a05-8245-df2d35b31e93","year":2014},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.026929Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:9ec1a7b5bb7aeb7ef54b10a2623af503e7ad3b32e69fa3c47fb54179d9f23ea3","observation_id":"0ce271d2-34e6-4573-9d3a-db287984915a","resolution":{"observed_at":"2026-08-07T11:11:47.011659Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.762558Z","title":null,"venue":null,"work_id":"d0c0a120-74b8-4bc8-8eff-5ece94e856e5","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.034496Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:8d42a6cfe1e5c8f1029009755bcb96ad65fca12271562b8ef06bba2545034c4f","observation_id":"b6b05972-cc67-46e6-a87a-83359ffd8299","resolution":{"observed_at":"2026-08-07T11:11:46.833921Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.604922Z","title":null,"venue":null,"work_id":"38189699-0779-4b41-a838-3b17f9d6fc05","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.041918Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:63d29a78cff273dbe3d94f2ddcaf695a37cb97fd7d1fcc8579f6c97b4f46d3a1","observation_id":"e328c37a-83bb-4e4c-a730-039d1cbc5886","resolution":{"observed_at":"2026-08-07T11:11:46.687749Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.437987Z","title":null,"venue":null,"work_id":"6e19d6b0-561c-4db0-8e97-d65f0034543a","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.048219Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:941b55c4e43d2fae01c24ce5e7c5e100c4a8a5652b99959dd7b190de0d96af95","observation_id":"698d7bd8-be6a-454b-8e88-bb1be0825d67","resolution":{"observed_at":"2026-08-07T11:11:46.500924Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.284134Z","title":null,"venue":null,"work_id":"c4dcf223-86b6-4eae-8d7e-e5200f7f08fd","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.057607Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:ffa03099f7dfcca3b3394e075841ff2277dad8ca8e6d27da3a7b0ec9c2044daa","observation_id":"fbe81663-6ea4-4573-a2d6-aece1c26a7a7","resolution":{"observed_at":"2026-08-07T11:11:46.341187Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:41.073628Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.073628Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:c87281794202dc250a1abbdf88970aa4aaaadc7bf33684b63a3dccc0944dde56","observation_id":"f840b952-4aeb-4d0a-a5c2-ec0b16d5dbc8","resolution":{"observed_at":"2026-08-07T11:11:41.073628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.03104","last_updated":"2024-08-10T14:57:37Z","snapshot_observed_at":"2026-07-06T18:40:57.362549Z","submitted_at":"2024-07-03T13:41:44Z","title":"KeyVideoLLM: Towards Large-scale Video Keyframe Selection","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.03104","snapshot_observed_at":"2026-08-07T11:11:41.088190Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.088190Z"},"links":{"cited_paper":"/paper/2407.03104","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:61749ce2bbb377e6275223c9f78ff6c17f565b59f5b88857b22fba49173e9fc0","observation_id":"34c1edd3-b612-4a39-9548-dd7104c7a1fd","resolution":{"observed_at":"2026-08-07T11:11:41.088190Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.757350Z","title":null,"venue":null,"work_id":"cfaef5f8-851d-4128-9855-a3f72426e7cb","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.096040Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:26ebfa008a188ad28e0221a45ff6fceabb140d72dab83e0199e685f89ffaaad9","observation_id":"30d96625-0842-4ed2-9311-fea5ff880d6a","resolution":{"observed_at":"2026-08-07T11:11:45.865795Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.135048Z","title":"INKOM Journal of Informatics, Control Systems, and Computers , 8, 2, 111–116","venue":null,"work_id":"a659e7fc-f6ab-4494-b05d-35e148f611b1","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.064465Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:bd0c492f817c19da7e30e8640198c29264130d75829a788ae66dd9e9bbea39df","observation_id":"a8af1ddb-f4bc-46b1-a991-bdb720a731fa","resolution":{"observed_at":"2026-08-07T11:11:46.209621Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.444232Z","title":null,"venue":null,"work_id":"753f50f6-cb89-424d-b800-a76d5b2fc3c2","year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.115746Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:48e6d28c14b4adbf61918a14331c08f884b886d7ad377db2a77dd20722a655d8","observation_id":"f5624e83-ee82-4f85-9bfc-0979e1cdc6a3","resolution":{"observed_at":"2026-08-07T11:11:45.529856Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2205.14472","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:41.657070Z","title":null,"venue":null,"work_id":"c78d985f-8d87-444b-b44b-80d14bc4b674","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.121903Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:4aed17ff455cb464e49f06705051c3f1c51d2d429859d4fa9aa2ee3b4aeb1496","observation_id":"fa3207c5-8f15-49eb-b2aa-b5d7ffeb7f00","resolution":{"observed_at":"2026-08-07T11:11:41.723881Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.281321Z","title":null,"venue":null,"work_id":"c69bf4ab-5a18-4ac1-b55a-fe38bb26ff03","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.127921Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:4fcc7aa2fcf07419ce97c8f24bbff865d8b011a3bdaab096dbd8fc800b560ee1","observation_id":"f58b63f7-e6a3-45f8-a13c-0260b06cac3d","resolution":{"observed_at":"2026-08-07T11:11:45.348513Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.147860Z","title":null,"venue":null,"work_id":"2d2d9626-bf1e-456f-87fb-50ba44a726e4","year":2025},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.132944Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:ba22e763b5f7b4c7fc6f397800f1477e55bc385b2c54d2a3663495de2cd387db","observation_id":"93856acf-0eeb-454f-b3d4-485c03ca11de","resolution":{"observed_at":"2026-08-07T11:11:45.206142Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.980091Z","title":null,"venue":null,"work_id":"2106a452-d92e-45ea-9d1f-11ca3b67fca4","year":2025},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.138543Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:9d18dcb9d72c5493716f33207bce4ce97dbe458ae6bcbaddaed6e13d4032c62c","observation_id":"daff0687-5277-404a-aa59-ab0618b10ee8","resolution":{"observed_at":"2026-08-07T11:11:45.074622Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.607673Z","title":null,"venue":null,"work_id":"4979b92d-0702-47f2-98e7-ea5a3f5dce92","year":2014},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.109119Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:395362a8886d8e295174431971855ff93f3f7804df7e17f3eb9c43d460635dcc","observation_id":"f4314728-db66-47d2-aba2-120b77bc42dd","resolution":{"observed_at":"2026-08-07T11:11:45.679596Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.696208Z","title":null,"venue":null,"work_id":"4e46550a-5e74-4300-8ec2-bdaec50cd2c0","year":2020},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.150109Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:d92d5b443745239d13aba9e57cfeb8ac44e432b3fb1a73127891385151b3bfef","observation_id":"85a5ef52-6361-41f9-a7f8-d8c5fa36ae85","resolution":{"observed_at":"2026-08-07T11:11:44.747159Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.584697Z","title":null,"venue":null,"work_id":"2dbae7b1-ba0a-4d8e-b841-87d93614bac5","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.156177Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:4604b40201bfb965141a7f2fe2ba765b5cb573f4fc3b1eff5b0df0b495ebf950","observation_id":"4d45c4fb-dba2-47ee-a3f2-c4a749dc1097","resolution":{"observed_at":"2026-08-07T11:11:44.609253Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:41.162226Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.162226Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:49e8d256d352b2f4e7970727ec1303bf5d009792ee6e961670aa287496c3ba71","observation_id":"58f9d618-02b6-4bde-baf3-a4ac1610d85e","resolution":{"observed_at":"2026-08-07T11:11:41.162226Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.447671Z","title":null,"venue":null,"work_id":"cf449f5a-ce85-4fad-b0b6-a59d50adca51","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.168338Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:3fd7f92255bdb6aad1bb6be1876e7906ae39eac4ebf412e3c2362e8df363c0b1","observation_id":"c967baa5-a25d-44a9-ae40-5bb5cb33688d","resolution":{"observed_at":"2026-08-07T11:11:44.501030Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.265207Z","title":null,"venue":null,"work_id":"1f74790f-aa3b-4141-a1c8-e021b2521e09","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.173920Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:11288fa21af2bc64d852e96ce1102d8cdc52dd294248f9598423327d25faac2a","observation_id":"413d38fa-5663-4ba0-a120-75793fa9d2f7","resolution":{"observed_at":"2026-08-07T11:11:44.351325Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.817016Z","title":null,"venue":null,"work_id":"5a6b04d4-0d79-4119-9b88-64b522e9ba7e","year":2014},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.145025Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:f1151bf7701d0b5e472818dbec24f0b2e528ebb90817989ab1566551557f30b6","observation_id":"d69f877d-eb78-440a-8c78-b3f2a65b8fed","resolution":{"observed_at":"2026-08-07T11:11:44.892312Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.114103Z","title":null,"venue":null,"work_id":"5241c636-50e9-4298-b140-e9c251cc1a5a","year":2018},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.185501Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:99ea3c0a6c7767987b2b95ddc481cb7b3cb18154179eafa70afeac4d78349464","observation_id":"ce172a75-48d2-47ae-a8c9-67d83a14a8f3","resolution":{"observed_at":"2026-08-07T11:11:44.181917Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.959494Z","title":null,"venue":null,"work_id":"81136201-cc72-45aa-918b-6dfa937910b2","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.191690Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:d8b89e3e6b68ee02df7fd72ad557288bbcfa3ce8aa4329fcd832c9b34752708d","observation_id":"fa1d919f-94af-4824-b0f3-68e64f999a42","resolution":{"observed_at":"2026-08-07T11:11:43.993073Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.829544Z","title":null,"venue":null,"work_id":"30ae84db-3d7a-4fc7-8c30-59811df3ea9c","year":2020},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.197615Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:bc354c5b95250d13b967c20cc34207881eab05f1a110ada4ee7b16fd5bbf1ac5","observation_id":"b142da1f-3d08-4ded-9ec8-8ffa160ab3ff","resolution":{"observed_at":"2026-08-07T11:11:43.886697Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.667526Z","title":null,"venue":null,"work_id":"34cc62ca-8f12-409d-a9f3-6cb53a3019f1","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.202782Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:97907fc20574f87149f60ef5920035bceac27e26ca2c975e3f430d450b71f3a2","observation_id":"c3ed313b-007c-4bc3-ae8d-945d764c235d","resolution":{"observed_at":"2026-08-07T11:11:43.761042Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.537584Z","title":null,"venue":null,"work_id":"34626edd-d725-4df3-951f-d58e88f20d86","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.213873Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:aec0c0854a26aa15d19bbebd81cad81ec912db6507722e08c36dc9512d7aeeda","observation_id":"2e7e7a2a-8074-4c3a-8e82-2b310b770410","resolution":{"observed_at":"2026-08-07T11:11:43.612449Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1606.05250","last_updated":"2016-10-11T02:42:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2016-06-16T16:36:00Z","title":"SQuAD: 100,000+ Questions for Machine Comprehension of Text","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.05250","snapshot_observed_at":"2026-08-07T11:11:41.179600Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.179600Z"},"links":{"cited_paper":"/paper/1606.05250","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:e49539099f0edd7d7fea4f096da5c022d4737122062e958f40ac07c075493042","observation_id":"a9210d8e-15e9-4536-9840-87b7fa650996","resolution":{"observed_at":"2026-08-07T11:11:41.179600Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.271422Z","title":null,"venue":null,"work_id":"acd85ba1-53d8-4b89-9fd8-87924bf8b876","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.225126Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:a592f7c160debb69effb642b0c1e871b23bb31b81843eaeb87587a03eae7f4e4","observation_id":"575b73a1-6eed-47f4-b84e-3ccbee0371d0","resolution":{"observed_at":"2026-08-07T11:11:43.311304Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.104112Z","title":null,"venue":null,"work_id":"70664d3b-7e01-47e1-b4c7-742dd09f8578","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.231261Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:f8b9df1af3347445b48e15198dfa5efee959d6b0dd6e6444720b73fd636da02a","observation_id":"18ebd25e-27f3-494f-9024-17d9c7206ef8","resolution":{"observed_at":"2026-08-07T11:11:43.186809Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1804.07461","last_updated":"2019-02-22T23:53:34Z","snapshot_observed_at":"2026-07-06T06:34:26.609892Z","submitted_at":"2018-04-20T06:35:04Z","title":"GLUE: A Multi-Task Benchmark and Analysis Platform for Natural Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1804.07461","snapshot_observed_at":"2026-08-07T11:11:41.235939Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.235939Z"},"links":{"cited_paper":"/paper/1804.07461","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:1628440f1bdf5fbcb6ecf4eaeca3843b2324ed951527f712beabcaae86db6905","observation_id":"ca2b91ff-be8f-40f3-a919-79d8d5fdb5c8","resolution":{"observed_at":"2026-08-07T11:11:41.235939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.929246Z","title":null,"venue":null,"work_id":"d7332849-1d18-4728-a8be-b013289249df","year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.243147Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:21a0f50cb8a0662c02f1e2d1fc43ceea9ef3aef53b27869a16cf292619538bee","observation_id":"d9023dd6-662d-408e-9a6c-cb026fd83268","resolution":{"observed_at":"2026-08-07T11:11:43.009275Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.19209","last_updated":"2025-03-14T13:57:16Z","snapshot_observed_at":"2026-08-05T17:33:53.862042Z","submitted_at":"2024-05-29T15:49:09Z","title":"VideoTree: Adaptive Tree-based Video Representation for LLM Reasoning on Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.19209","snapshot_observed_at":"2026-08-07T11:11:41.249261Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.249261Z"},"links":{"cited_paper":"/paper/2405.19209","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:0ce0e84fbcd3932a033799bf53c5aaf3d2a72209188af24ede6d5f8ce9134728","observation_id":"5c494f6b-4435-49f1-9976-d8d62903bfc5","resolution":{"observed_at":"2026-08-07T11:11:41.249261Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.377981Z","title":null,"venue":null,"work_id":"f01b49d1-c669-4c30-b354-d1b3cde23971","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.219099Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:4dd3ed018e5228d3de6823ed2fcc09203acc21d237b5b2be47067181e523de7a","observation_id":"dfb015ca-0bdc-48ae-a64f-2f514181698e","resolution":{"observed_at":"2026-08-07T11:11:43.482117Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.773084Z","title":null,"venue":null,"work_id":"f0ddb2ed-1016-43d4-9a85-8e61710a86e4","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.262182Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:bc4f62670f920da19bdf1bd1589a7b38e39b2b26a28b6560ab0e1d0e9cbec3fa","observation_id":"2af3d5b6-867b-4e6e-be67-f162c3dc4955","resolution":{"observed_at":"2026-08-07T11:11:42.818519Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.623297Z","title":null,"venue":null,"work_id":"afb1dde2-1b8d-451f-b403-a5ed62c302d8","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.267757Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:724414ed8449a157f051f48fa5775b67cf78329064990e910b355b0f12b068e6","observation_id":"f5af5e5b-85c5-43a3-ac02-e79c8bbfc479","resolution":{"observed_at":"2026-08-07T11:11:42.722906Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.533771Z","title":null,"venue":null,"work_id":"d4cd150a-acf5-49cf-a1f2-702786462835","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.274376Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:969822317671e7f31397e17e5fd45b96e68b71b1c11d9195db07fb00be02c73f","observation_id":"ac65a02e-b084-4eea-a5f9-17b415c230e4","resolution":{"observed_at":"2026-08-07T11:11:42.569810Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.240911Z","title":null,"venue":null,"work_id":"fc5a96a8-576d-4eac-a44f-ca6adfdac9f7","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.287077Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:e43eccbd103ed6d113570ac121d7bd2367cc0943c162c83959d14d8114ac1ca5","observation_id":"2596dce1-b6d4-4502-9172-74a9cd7598fc","resolution":{"observed_at":"2026-08-07T11:11:42.314723Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.00476","last_updated":"2021-10-01T15:09:22Z","snapshot_observed_at":"2026-07-06T11:53:31.073250Z","submitted_at":"2021-10-01T15:09:22Z","title":"ResNet strikes back: An improved training procedure in timm","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.00476","snapshot_observed_at":"2026-08-07T11:11:41.255287Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.255287Z"},"links":{"cited_paper":"/paper/2110.00476","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:f637ee42e57c074bf2ebe93364029ef51664dc18c832f0711e103ac9e52020ab","observation_id":"39a45272-c501-499b-81bf-6e6599066c32","resolution":{"observed_at":"2026-08-07T11:11:41.255287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:40.992307Z","title":"In Computer Vision–ECCV 2014: 13th European Conference, Zurich, Switzerland, September 6-12, 2014, Proceedings, Part VII 13","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2014,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.992307Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:bd2a00cca68e06ec0ba2ecf6a9657eff66eb755c8c276b975155672a219f05c7","observation_id":"6b4c05a5-182b-45a6-a4f8-7d98d8bd6c17","resolution":{"observed_at":"2026-08-07T11:11:40.992307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.659369Z","title":"In 2017 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","venue":null,"work_id":"0581a90d-6f78-4871-8696-2922b34cf85e","year":2017},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.976125Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:2133c823a7645b1f465fdfd5923621edbf67ff644ea66e9ff0813270ab01ec03","observation_id":"14a40be6-d566-4d7d-9fb9-5ef82fbd00ec","resolution":{"observed_at":"2026-08-07T11:11:47.703083Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.380191Z","title":"Mathematical Problems in Engineering , 2019, 1, 5217961","venue":null,"work_id":"3f83dd54-1571-40b2-af1a-384908fde1dc","year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.280976Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:6e8a1ee8fc735b1036bd9356d2f17201e133e56abd54c3f0365521e7ac0faead","observation_id":"4b39f2a6-16c5-423f-b8d0-f6eb1cd993ba","resolution":{"observed_at":"2026-08-07T11:11:42.447867Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.933272Z","title":"Pattern Recognition, 111, 107677","venue":null,"work_id":"b4cca118-5f37-4f47-8cf0-5fb0630a459c","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.081280Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:d1ab0e7af6e42ace7ff30e50fbc5cbfc7325e75d01f6ce7b5925e652f4b25d94","observation_id":"b8b5bd29-fbb6-4690-9ebf-076fb949b424","resolution":{"observed_at":"2026-08-07T11:11:46.021545Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.10122","last_updated":"2024-10-01T12:07:31Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-11-16T10:59:44Z","title":"Video-LLaVA: Learning United Visual Representation by Alignment Before Projection","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.10122","snapshot_observed_at":"2026-08-07T11:11:41.103071Z","title":"arXiv preprint arXiv:2311.10122","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.103071Z"},"links":{"cited_paper":"/paper/2311.10122","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:81036dcaf26f5025bfebcf4a2797a50fa32da0b4468e1f2ba52ecfbb2100bf32","observation_id":"00668896-2ef0-4f8f-8d1c-908d4e10382f","resolution":{"observed_at":"2026-08-07T11:11:41.103071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.510518Z","title":"Scientific Reports, 15, 1, 2126","venue":null,"work_id":"eb298eb1-3b02-4fd3-954e-3bcfcfe111f5","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.881716Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:4ec177971866106e6ac9997cc1820bd8981bcb4a9705e89ff009cd72f4564c43","observation_id":"d6d1f0d3-5f41-47da-b6b7-e20e2ad4ccde","resolution":{"observed_at":"2026-08-07T11:11:48.562185Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-07T11:03:01.827723Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations"},"reference_resolution":{"displayed":65,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":56,"verified_exact":4,"verified_fuzzy":5},"total_outbound_references":65},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 65 of 65 outbound references and 1 inbound Pith citation observation for arXiv:2506.05395."}