{"as_of":"2026-08-18T02:28:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:263ce8806808c5e8e66add255375f5f530e2dd45115a98cdcbc761ab857aa8db","coverage":[{"denominator":65,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":65,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:11:41.287077Z","state":"measured"},{"denominator":66,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":66,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-17T06:30:58.91139+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-28T18:44:39.497333Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-28T20:22:37.768955Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"cited_work":{"arxiv_id":"2506.05395","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.05395","snapshot_observed_at":"2026-06-28T20:22:37.768955Z","title":null,"venue":null,"work_id":"a2348d7b-2cc8-4f45-98c4-6ea6d1ea5c57","year":2025},"citing_paper":{"arxiv_id":"2606.00664","last_updated":"2026-05-30T10:41:34Z","snapshot_observed_at":"2026-08-16T01:40:38.393182Z","submitted_at":"2026-05-30T10:41:34Z","title":"SKIP: Sparse Keyframe Interpolation Paradigm for Efficient Embodied World Models","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-06-28T18:44:39.497333Z"},"links":{"cited_paper":"/paper/2506.05395","citing_paper":"/paper/2606.00664"},"observation_digest":"sha256:359607d593dcad0a7a5e1b89e45714bba988b3dfc39e4d635b1ce64c0866a4be","observation_id":"621b29a9-218b-469c-badc-e48d9b208b65","resolution":{"observed_at":"2026-06-28T20:22:37.770265Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.05395/citation-record","integrity":"/paper/2506.05395/integrity","json":"/paper/2506.05395/citation-record.json","paper":"/paper/2506.05395"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:49.436097Z","title":null,"venue":null,"work_id":"330a80b8-2a66-4b60-b42e-da748ec7702e","year":2015},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.824919Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:38843cb5a1d78507cc08ea77b3c03299a78283518f417a3a70b9260d58ee206b","observation_id":"99cfbce8-9a52-4ee3-b98f-0d6468c0b2d1","resolution":{"observed_at":"2026-08-07T11:11:49.556238Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:49.262150Z","title":null,"venue":null,"work_id":"0044161d-80b4-42c8-b1fd-ee8309a77e1d","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.833719Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:a185a05bbaab91f90e351139387aa266e686068a279f8b7a0edc0774085b8018","observation_id":"9ee347f2-02ef-4c0c-bfb9-d23023592a8e","resolution":{"observed_at":"2026-08-07T11:11:49.348075Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:49.009072Z","title":null,"venue":null,"work_id":"1360d8ee-7c47-41e6-91b0-706ac34e66f2","year":2018},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.839853Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:37705a62f781d16014c521c2c9c6bb16358d4ef3a01a3a6f8ce1934ff3ca9da8","observation_id":"cf4534ea-7e88-4c0a-924f-99c9d1654079","resolution":{"observed_at":"2026-08-07T11:11:49.123084Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.869854Z","title":null,"venue":null,"work_id":"a956c8ad-064e-4a0f-98d8-0cc6b6871e8b","year":2018},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.849209Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:d670ae27fe0a5608e037ffe1db6ca2c76547cc12ad54ce6e454db817b839df5f","observation_id":"21f9572d-1fe8-4544-8618-eb60246b9a91","resolution":{"observed_at":"2026-08-07T11:11:48.927052Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.750922Z","title":null,"venue":null,"work_id":"6e91c9c1-2c74-4d36-8c80-57602640976a","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.857381Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:30d6be61b5672a9b1aa6b9df2ea9cb2390074322fa4007db92aaf9b53ddeb45c","observation_id":"23853136-1253-44bf-a870-4a4328f48c3e","resolution":{"observed_at":"2026-08-07T11:11:48.804137Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.11249","last_updated":"2021-01-27T08:13:19Z","snapshot_observed_at":"2026-08-16T18:49:42.961172Z","submitted_at":"2021-01-27T08:13:19Z","title":"Efficient Video Summarization Framework using EEG and Eye-tracking Signals","version":1},"cited_work":{"arxiv_id":"2101.11249","doi":null,"metadata_source":"pith","pith_arxiv_id":"2101.11249","snapshot_observed_at":"2026-08-07T11:11:42.054602Z","title":"Efficient Video Summarization Framework using EEG and Eye-tracking Signals","venue":"cs.CV","work_id":"a480bbc7-53ee-4c19-87f6-499e285855c8","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.863251Z"},"links":{"cited_paper":"/paper/2101.11249","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:058f2b69bbf0d646e3174f740b20c31e6cf3b7fbca55763f7c7102a18986bec9","observation_id":"65daa8fb-67f8-409c-a318-e459437439ea","resolution":{"observed_at":"2026-08-07T11:11:42.135561Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.631977Z","title":null,"venue":null,"work_id":"c6c302fd-49f4-4783-b1cc-7b94df755c75","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.873911Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:5c03807885565c9a18c468f2bee31f19d8e298738e3bcd0379bb97223fcd19b3","observation_id":"01a20a45-f71b-4236-ae2b-e29efe7e3020","resolution":{"observed_at":"2026-08-07T11:11:48.689154Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.398995Z","title":null,"venue":null,"work_id":"8f4013f5-7771-4198-a527-b0113b94a09a","year":2013},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.891595Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:665172c8a4b44f036e4d9b4f8b65a7f63995f4fbdf06067eb268adff2665cd66","observation_id":"50138b50-0413-4747-a6a3-dd5c8a7d1bfa","resolution":{"observed_at":"2026-08-07T11:11:48.452841Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.12971","last_updated":"2023-10-19T17:59:01Z","snapshot_observed_at":"2026-08-16T14:50:37.149811Z","submitted_at":"2023-10-19T17:59:01Z","title":"CLAIR: Evaluating Image Captions with Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.12971","snapshot_observed_at":"2026-08-07T11:11:40.898559Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.898559Z"},"links":{"cited_paper":"/paper/2310.12971","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:852d6182b619ffed0c7f55c0c21393be09863d9d2cf7bfc8b7b666c082174a57","observation_id":"ce457575-69a5-422c-ac8f-0596989bb77d","resolution":{"observed_at":"2026-08-07T11:11:40.898559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.13292","last_updated":"2023-05-23T07:48:15Z","snapshot_observed_at":"2026-08-16T15:30:58.546906Z","submitted_at":"2023-05-22T17:51:22Z","title":"VideoLLM: Modeling Video Sequence with Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.13292","snapshot_observed_at":"2026-08-07T11:11:40.910714Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.910714Z"},"links":{"cited_paper":"/paper/2305.13292","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:a228195b532e102fa52e417bcd6b2d0165960fad7a6276e1ac59ad8c1da4f4fa","observation_id":"0a03daa1-2d5c-431d-b20d-263c1ca55de5","resolution":{"observed_at":"2026-08-07T11:11:40.910714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.259155Z","title":null,"venue":null,"work_id":"8c71ec38-6e6c-4657-b928-d226a93ce2cf","year":2011},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.923028Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:aea9cd751aa79007fdb6f29b41f291a73b3cfd784202a3e5b9e28774e2fe1281","observation_id":"ed2f5c9c-9a40-40b6-bcef-1782e02f3fcf","resolution":{"observed_at":"2026-08-07T11:11:48.323123Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1109/cvpr.2009.5","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:41.333474Z","title":null,"venue":null,"work_id":"a208e36f-32cc-4de8-a605-570838edf9ed","year":2009},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.932393Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:3ae88efcb6b7e28619d8036dd29453f46e20b5fa790712e1071a4a14e20bd849","observation_id":"b1c107cb-d8fe-488f-8b18-2d7691655d9c","resolution":{"observed_at":"2026-08-07T11:11:41.340081Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.163556Z","title":null,"venue":null,"work_id":"f9e6bf9f-17b5-4f95-a278-3fd679703202","year":2003},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.941825Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:2024d00416f5b6c96996f509bcc97c67f22215df07af522d40b9c57e95aceeb4","observation_id":"6f67d612-d5c6-4650-afa0-46e7ca96db70","resolution":{"observed_at":"2026-08-07T11:11:48.195996Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.002797Z","title":null,"venue":null,"work_id":"62fd661c-54a6-4aba-a2ec-01dca04d3ee1","year":1996},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.947137Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:eef04e07d67dd90203e1dc72da3c95a4481792a9747cc20706eb321752135c94","observation_id":"37c0b276-d7cc-4f1d-86d6-0234d0a0287b","resolution":{"observed_at":"2026-08-07T11:11:48.075632Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.877543Z","title":null,"venue":null,"work_id":"f6dee9a9-0b7c-41fe-901b-ae30b56dbeba","year":2025},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.956872Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:13a10accb69c0347b33931a485f54f6a22f0d590051bd9698bbc7c934101194d","observation_id":"32512852-83f0-4ab9-803c-f808fb0263ae","resolution":{"observed_at":"2026-08-07T11:11:47.933065Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.10173","last_updated":"2023-07-14T16:49:42Z","snapshot_observed_at":"2026-08-16T15:55:04.357083Z","submitted_at":"2023-02-15T19:09:34Z","title":"VideoSum: A Python Library for Surgical Video Summarization","version":2},"cited_work":{"arxiv_id":"2303.10173","doi":null,"metadata_source":"pith","pith_arxiv_id":"2303.10173","snapshot_observed_at":"2026-08-07T11:11:41.867397Z","title":"VideoSum: A Python Library for Surgical Video Summarization","venue":"eess.IV","work_id":"132f6468-5e11-40fe-99d4-70309686b9c4","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.962560Z"},"links":{"cited_paper":"/paper/2303.10173","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:668532134cc1ff29845c6a1d2914bf58210954ee957d45086b29d6e74b36ea72","observation_id":"9a044658-5cd4-46dd-b437-70a795eb662b","resolution":{"observed_at":"2026-08-07T11:11:41.938502Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.773668Z","title":null,"venue":null,"work_id":"08edff2a-4063-4608-a513-a28ab5ff8b9b","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.968488Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:40725aadf53fd00f81b0a665982bd333731733c7f11c0aa9b71d486b5d84fcf3","observation_id":"da12a42d-f8de-4ab4-9d7f-f445017b3556","resolution":{"observed_at":"2026-08-07T11:11:47.819914Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:40.983011Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.983011Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:9caf3dc3ea634d6c840a8112f0e5943426b3a27da25c3962c2bb61f25badfac5","observation_id":"d3433d7d-ff72-4c34-9199-d63cea6d9af1","resolution":{"observed_at":"2026-08-07T11:11:40.983011Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.454261Z","title":null,"venue":null,"work_id":"24930598-93c2-422b-87b5-5e9218ecf4a5","year":2016},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.001921Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:0d113f82eb3fb17fb359d36dd014a8a44da8b8d74973dc18d8c52f7ba1bd41b1","observation_id":"46189abc-c034-4146-9748-9d6ee98d558e","resolution":{"observed_at":"2026-08-07T11:11:47.548231Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.271425Z","title":null,"venue":null,"work_id":"91665530-c526-4dd4-88aa-5ae4a353e79b","year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.009725Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:2095eea581063458a7254698412c76e67e92486ff012d65cc16ca6d28921a2a6","observation_id":"5011cb4f-a5a4-4150-aa5d-609ec90f74ce","resolution":{"observed_at":"2026-08-07T11:11:47.353245Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.106798Z","title":null,"venue":null,"work_id":"0bffc55e-b65c-4072-ba15-d755ba99f141","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.019761Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:790166d4de04dc5a9be47b4983f31914f62195bbd2a03eee1c6766daf4881d73","observation_id":"0d72a24c-1815-48ca-b5f9-299d474ad21e","resolution":{"observed_at":"2026-08-07T11:11:47.180318Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.943192Z","title":null,"venue":null,"work_id":"10233f15-f129-4a05-8245-df2d35b31e93","year":2014},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.026929Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:bdef7081176ac2f27661e239b1739c727c977b78c1b150d5358a4c047bae3463","observation_id":"0ce271d2-34e6-4573-9d3a-db287984915a","resolution":{"observed_at":"2026-08-07T11:11:47.011659Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.762558Z","title":null,"venue":null,"work_id":"d0c0a120-74b8-4bc8-8eff-5ece94e856e5","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.034496Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:517a122a62b1028f849b16c34289f97d7cebbb554cfcb24a8d0b9aa6a4b54934","observation_id":"b6b05972-cc67-46e6-a87a-83359ffd8299","resolution":{"observed_at":"2026-08-07T11:11:46.833921Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.604922Z","title":null,"venue":null,"work_id":"38189699-0779-4b41-a838-3b17f9d6fc05","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.041918Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:159a88bc7e13345a0616c8ea092a5a73686df4cf6ab70f29273adb40b885d55d","observation_id":"e328c37a-83bb-4e4c-a730-039d1cbc5886","resolution":{"observed_at":"2026-08-07T11:11:46.687749Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.437987Z","title":null,"venue":null,"work_id":"6e19d6b0-561c-4db0-8e97-d65f0034543a","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.048219Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:92e4395a4ae287a6b2d07a10d4a9725334522d0f022d8929e711b63ce17bb4c4","observation_id":"698d7bd8-be6a-454b-8e88-bb1be0825d67","resolution":{"observed_at":"2026-08-07T11:11:46.500924Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.284134Z","title":null,"venue":null,"work_id":"c4dcf223-86b6-4eae-8d7e-e5200f7f08fd","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.057607Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:f37050aece820e5789b999fc208ed814afd10c901ffa2011ee6c80683f5d8bae","observation_id":"fbe81663-6ea4-4573-a2d6-aece1c26a7a7","resolution":{"observed_at":"2026-08-07T11:11:46.341187Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:41.073628Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.073628Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:6c6f3985a6d0f6fb93b35501a67cfab7f976c5125fe6a51ddae0d228d8ed3cbc","observation_id":"f840b952-4aeb-4d0a-a5c2-ec0b16d5dbc8","resolution":{"observed_at":"2026-08-07T11:11:41.073628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.03104","last_updated":"2024-08-10T14:57:37Z","snapshot_observed_at":"2026-08-16T13:37:23.021128Z","submitted_at":"2024-07-03T13:41:44Z","title":"KeyVideoLLM: Towards Large-scale Video Keyframe Selection","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.03104","snapshot_observed_at":"2026-08-07T11:11:41.088190Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.088190Z"},"links":{"cited_paper":"/paper/2407.03104","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:3ac09dd7a2a4eb2e877892248a056008cd3a5357319616af7840fc0c6c77ca2e","observation_id":"34c1edd3-b612-4a39-9548-dd7104c7a1fd","resolution":{"observed_at":"2026-08-07T11:11:41.088190Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.757350Z","title":null,"venue":null,"work_id":"cfaef5f8-851d-4128-9855-a3f72426e7cb","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.096040Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:c9b026524f87d2aa5e342d8fa6f1d6a29b46c972e51cdeaf2c2ad2ce634a897b","observation_id":"30d96625-0842-4ed2-9311-fea5ff880d6a","resolution":{"observed_at":"2026-08-07T11:11:45.865795Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:46.135048Z","title":"INKOM Journal of Informatics, Control Systems, and Computers , 8, 2, 111–116","venue":null,"work_id":"a659e7fc-f6ab-4494-b05d-35e148f611b1","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.064465Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:e22af766d668dba5af04b75141a7697b11003a480d4f5fb690ad8fea12c36e49","observation_id":"a8af1ddb-f4bc-46b1-a991-bdb720a731fa","resolution":{"observed_at":"2026-08-07T11:11:46.209621Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.444232Z","title":null,"venue":null,"work_id":"753f50f6-cb89-424d-b800-a76d5b2fc3c2","year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.115746Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:f4664692a85daeabd0e016d63ac8ab97519fb71de7c6496deb38d4c6e20140c2","observation_id":"f5624e83-ee82-4f85-9bfc-0979e1cdc6a3","resolution":{"observed_at":"2026-08-07T11:11:45.529856Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2205.14472","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:41.657070Z","title":null,"venue":null,"work_id":"c78d985f-8d87-444b-b44b-80d14bc4b674","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.121903Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:7600b1324d730c97a15b7527a2d446971f99bf453c3b9a781dda4dcb06aee63e","observation_id":"fa3207c5-8f15-49eb-b2aa-b5d7ffeb7f00","resolution":{"observed_at":"2026-08-07T11:11:41.723881Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.281321Z","title":null,"venue":null,"work_id":"c69bf4ab-5a18-4ac1-b55a-fe38bb26ff03","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.127921Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:452c512c15dc6e9a058a06f67b148273fc8990c27f8e1b47705e00aa8ed0beb0","observation_id":"f58b63f7-e6a3-45f8-a13c-0260b06cac3d","resolution":{"observed_at":"2026-08-07T11:11:45.348513Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.147860Z","title":null,"venue":null,"work_id":"2d2d9626-bf1e-456f-87fb-50ba44a726e4","year":2025},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.132944Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:5b36937f207345ffbe372243881d1d9853d483eb80b4e51027f5deee7802abd9","observation_id":"93856acf-0eeb-454f-b3d4-485c03ca11de","resolution":{"observed_at":"2026-08-07T11:11:45.206142Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.980091Z","title":null,"venue":null,"work_id":"2106a452-d92e-45ea-9d1f-11ca3b67fca4","year":2025},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.138543Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:e6b1ce57394cf6eb44d001a2b52931f6be95d581f2f2f4c817f77149cdc05f59","observation_id":"daff0687-5277-404a-aa59-ab0618b10ee8","resolution":{"observed_at":"2026-08-07T11:11:45.074622Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.607673Z","title":null,"venue":null,"work_id":"4979b92d-0702-47f2-98e7-ea5a3f5dce92","year":2014},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.109119Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:7fb06ccddf6f66fd99d787aa81e088fe8d44e3452b12657d6b0e75fc314c102d","observation_id":"f4314728-db66-47d2-aba2-120b77bc42dd","resolution":{"observed_at":"2026-08-07T11:11:45.679596Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.696208Z","title":null,"venue":null,"work_id":"4e46550a-5e74-4300-8ec2-bdaec50cd2c0","year":2020},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.150109Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:fb77025c7a5908fbe34c88b14310731fb24d1e42669f0c002ead3145620bd3eb","observation_id":"85a5ef52-6361-41f9-a7f8-d8c5fa36ae85","resolution":{"observed_at":"2026-08-07T11:11:44.747159Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.584697Z","title":null,"venue":null,"work_id":"2dbae7b1-ba0a-4d8e-b841-87d93614bac5","year":2022},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.156177Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:b5d3efd73de5264d60499a9374d561564a92258a86b59041eb771eb9b84beae8","observation_id":"4d45c4fb-dba2-47ee-a3f2-c4a749dc1097","resolution":{"observed_at":"2026-08-07T11:11:44.609253Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:41.162226Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.162226Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:5d1dee73c9e7d488352f0a25b7b891a1492a55f058c5308ea5581f0336932b44","observation_id":"58f9d618-02b6-4bde-baf3-a4ac1610d85e","resolution":{"observed_at":"2026-08-07T11:11:41.162226Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.447671Z","title":null,"venue":null,"work_id":"cf449f5a-ce85-4fad-b0b6-a59d50adca51","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.168338Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:2dd953124aa02c138bda37da34b8e648e862c44f9aaaa8afcfe324d10b229ab9","observation_id":"c967baa5-a25d-44a9-ae40-5bb5cb33688d","resolution":{"observed_at":"2026-08-07T11:11:44.501030Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.265207Z","title":null,"venue":null,"work_id":"1f74790f-aa3b-4141-a1c8-e021b2521e09","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.173920Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:fe2c3128a8529d6a398e6a2f1385bf95e198c654f81a6902a8f7e5dceb9b238d","observation_id":"413d38fa-5663-4ba0-a120-75793fa9d2f7","resolution":{"observed_at":"2026-08-07T11:11:44.351325Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.817016Z","title":null,"venue":null,"work_id":"5a6b04d4-0d79-4119-9b88-64b522e9ba7e","year":2014},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.145025Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:2b4d677428e2be3d90f1ec563d1c0b8c13432f269f2f1ace0a32d9cfddfb48a2","observation_id":"d69f877d-eb78-440a-8c78-b3f2a65b8fed","resolution":{"observed_at":"2026-08-07T11:11:44.892312Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:44.114103Z","title":null,"venue":null,"work_id":"5241c636-50e9-4298-b140-e9c251cc1a5a","year":2018},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.185501Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:72d0d3c2fde02f18b9d3081e1bb198f618ba7b5a4374072cce969d04aea41dc3","observation_id":"ce172a75-48d2-47ae-a8c9-67d83a14a8f3","resolution":{"observed_at":"2026-08-07T11:11:44.181917Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.959494Z","title":null,"venue":null,"work_id":"81136201-cc72-45aa-918b-6dfa937910b2","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.191690Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:3bc3f5a6bfca25fb9c58d19b8bc4ede219bf05c78ff6efe1bca3b3b6ab15500c","observation_id":"fa1d919f-94af-4824-b0f3-68e64f999a42","resolution":{"observed_at":"2026-08-07T11:11:43.993073Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.829544Z","title":null,"venue":null,"work_id":"30ae84db-3d7a-4fc7-8c30-59811df3ea9c","year":2020},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.197615Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:7e66857863138ed36ac6eaf3e7bed2d714f153dc34a0caf679395bb91d665445","observation_id":"b142da1f-3d08-4ded-9ec8-8ffa160ab3ff","resolution":{"observed_at":"2026-08-07T11:11:43.886697Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.667526Z","title":null,"venue":null,"work_id":"34cc62ca-8f12-409d-a9f3-6cb53a3019f1","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.202782Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:cac1e4aec6ecd85ced99f18495445e8c86a8d2acd7b5d270c48e5e9c04aebe32","observation_id":"c3ed313b-007c-4bc3-ae8d-945d764c235d","resolution":{"observed_at":"2026-08-07T11:11:43.761042Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.537584Z","title":null,"venue":null,"work_id":"34626edd-d725-4df3-951f-d58e88f20d86","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.213873Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:db2302a028fa899ba4b3ba74d7e019a0bcbb08f5506d4258355595ec85e8afaa","observation_id":"2e7e7a2a-8074-4c3a-8e82-2b310b770410","resolution":{"observed_at":"2026-08-07T11:11:43.612449Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1606.05250","last_updated":"2016-10-11T02:42:36Z","snapshot_observed_at":"2026-08-17T14:32:22.468812Z","submitted_at":"2016-06-16T16:36:00Z","title":"SQuAD: 100,000+ Questions for Machine Comprehension of Text","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1606.05250","snapshot_observed_at":"2026-08-07T11:11:41.179600Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.179600Z"},"links":{"cited_paper":"/paper/1606.05250","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:e5e88edf0d2b8ebc6c46e0f429e660eddcd2d88b0b249fce08dc3a823c639b5f","observation_id":"a9210d8e-15e9-4536-9840-87b7fa650996","resolution":{"observed_at":"2026-08-07T11:11:41.179600Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.271422Z","title":null,"venue":null,"work_id":"acd85ba1-53d8-4b89-9fd8-87924bf8b876","year":2023},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.225126Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:8671f0566187ce070f01b4cce9e9273012b1c0b6f14fa3191b58f177310491e1","observation_id":"575b73a1-6eed-47f4-b84e-3ccbee0371d0","resolution":{"observed_at":"2026-08-07T11:11:43.311304Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.104112Z","title":null,"venue":null,"work_id":"70664d3b-7e01-47e1-b4c7-742dd09f8578","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.231261Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:3011382682c32d9e0e6afe8cd8f44eb1197e27aeddc865cfd018d5a33b19effd","observation_id":"18ebd25e-27f3-494f-9024-17d9c7206ef8","resolution":{"observed_at":"2026-08-07T11:11:43.186809Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1804.07461","last_updated":"2019-02-22T23:53:34Z","snapshot_observed_at":"2026-08-16T09:50:11.319379Z","submitted_at":"2018-04-20T06:35:04Z","title":"GLUE: A Multi-Task Benchmark and Analysis Platform for Natural Language Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1804.07461","snapshot_observed_at":"2026-08-07T11:11:41.235939Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.235939Z"},"links":{"cited_paper":"/paper/1804.07461","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:6adaa46d57ad0042fa324ca1ef4f4ed65cc77edf46bc370aaab124f38436957f","observation_id":"ca2b91ff-be8f-40f3-a919-79d8d5fdb5c8","resolution":{"observed_at":"2026-08-07T11:11:41.235939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.929246Z","title":null,"venue":null,"work_id":"d7332849-1d18-4728-a8be-b013289249df","year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.243147Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:1968653e79e33c2a5ee218f9bdf40ef85d3f920b77567db0f63e38750b821b90","observation_id":"d9023dd6-662d-408e-9a6c-cb026fd83268","resolution":{"observed_at":"2026-08-07T11:11:43.009275Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.19209","last_updated":"2025-03-14T13:57:16Z","snapshot_observed_at":"2026-08-17T21:35:27.995630Z","submitted_at":"2024-05-29T15:49:09Z","title":"VideoTree: Adaptive Tree-based Video Representation for LLM Reasoning on Long Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.19209","snapshot_observed_at":"2026-08-07T11:11:41.249261Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.249261Z"},"links":{"cited_paper":"/paper/2405.19209","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:526c686b8c6448a3b6dec625e8f5427db90f1f7f96a6532fc9c5ffdc3d64cc8f","observation_id":"5c494f6b-4435-49f1-9976-d8d62903bfc5","resolution":{"observed_at":"2026-08-07T11:11:41.249261Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:43.377981Z","title":null,"venue":null,"work_id":"f01b49d1-c669-4c30-b354-d1b3cde23971","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.219099Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:224ed60aeb79b1d447c259cebc981856e279c90886fd594fcf9eb5eee402c7c5","observation_id":"dfb015ca-0bdc-48ae-a64f-2f514181698e","resolution":{"observed_at":"2026-08-07T11:11:43.482117Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.773084Z","title":null,"venue":null,"work_id":"f0ddb2ed-1016-43d4-9a85-8e61710a86e4","year":2024},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.262182Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:a884f7150e3c1103bd5a467fd33f880b31d8e15bca7993d6c736ce936f6cfc76","observation_id":"2af3d5b6-867b-4e6e-be67-f162c3dc4955","resolution":{"observed_at":"2026-08-07T11:11:42.818519Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.623297Z","title":null,"venue":null,"work_id":"afb1dde2-1b8d-451f-b403-a5ed62c302d8","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.267757Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:caba00f052590b70225a4065d86140ae87d6d6acf5b654762f3f03a8e09e34ca","observation_id":"f5af5e5b-85c5-43a3-ac02-e79c8bbfc479","resolution":{"observed_at":"2026-08-07T11:11:42.722906Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.533771Z","title":null,"venue":null,"work_id":"d4cd150a-acf5-49cf-a1f2-702786462835","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.274376Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:1ccab0bdd9b5a7abe0d30ac6bef485d2585abdafbb877365c7dc7cc346f7e6df","observation_id":"ac65a02e-b084-4eea-a5f9-17b415c230e4","resolution":{"observed_at":"2026-08-07T11:11:42.569810Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.240911Z","title":null,"venue":null,"work_id":"fc5a96a8-576d-4eac-a44f-ca6adfdac9f7","year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.287077Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:d2da86ea4c2f4b787834086978967f8afa176973a2574ecca7c8a97532679fbb","observation_id":"2596dce1-b6d4-4502-9172-74a9cd7598fc","resolution":{"observed_at":"2026-08-07T11:11:42.314723Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2110.00476","last_updated":"2021-10-01T15:09:22Z","snapshot_observed_at":"2026-08-16T17:52:28.101362Z","submitted_at":"2021-10-01T15:09:22Z","title":"ResNet strikes back: An improved training procedure in timm","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.00476","snapshot_observed_at":"2026-08-07T11:11:41.255287Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.255287Z"},"links":{"cited_paper":"/paper/2110.00476","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:7c33866fc0b67990c4864a43da2bf6f365db349c6c06c6ffdc8efcab0a9e44be","observation_id":"39a45272-c501-499b-81bf-6e6599066c32","resolution":{"observed_at":"2026-08-07T11:11:41.255287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:40.992307Z","title":"In Computer Vision–ECCV 2014: 13th European Conference, Zurich, Switzerland, September 6-12, 2014, Proceedings, Part VII 13","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2014,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.992307Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:24d6d786e3cb784d50ee68a7240987adf5ef2b31616fe696d001c88c1bdb06f6","observation_id":"6b4c05a5-182b-45a6-a4f8-7d98d8bd6c17","resolution":{"observed_at":"2026-08-07T11:11:40.992307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:47.659369Z","title":"In 2017 IEEE International Conference on Acoustics, Speech and Signal Processing (ICASSP)","venue":null,"work_id":"0581a90d-6f78-4871-8696-2922b34cf85e","year":2017},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2017,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.976125Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:617ff6684c16674899302d59c6761aa48805dbe2c0c0d1ab8eaac2f58f1ddbd1","observation_id":"14a40be6-d566-4d7d-9fb9-5ef82fbd00ec","resolution":{"observed_at":"2026-08-07T11:11:47.703083Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:42.380191Z","title":"Mathematical Problems in Engineering , 2019, 1, 5217961","venue":null,"work_id":"3f83dd54-1571-40b2-af1a-384908fde1dc","year":2019},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2019,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.280976Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:da2d780e018806e71a20d17f838497b85124abdfc91628f14deda4c0bf8d1c01","observation_id":"4b39f2a6-16c5-423f-b8d0-f6eb1cd993ba","resolution":{"observed_at":"2026-08-07T11:11:42.447867Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:45.933272Z","title":"Pattern Recognition, 111, 107677","venue":null,"work_id":"b4cca118-5f37-4f47-8cf0-5fb0630a459c","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.081280Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:fa85453d7267a9a93f392189993ecdda063449d1949b0755c288c95870fb2d4e","observation_id":"b8b5bd29-fbb6-4690-9ebf-076fb949b424","resolution":{"observed_at":"2026-08-07T11:11:46.021545Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.10122","last_updated":"2024-10-01T12:07:31Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-11-16T10:59:44Z","title":"Video-LLaVA: Learning United Visual Representation by Alignment Before Projection","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.10122","snapshot_observed_at":"2026-08-07T11:11:41.103071Z","title":"arXiv preprint arXiv:2311.10122","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:41.103071Z"},"links":{"cited_paper":"/paper/2311.10122","citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:1a1b2a888f030c3163da2f7408ff316b5dc07478b324b7ba4d900532d693adbd","observation_id":"00668896-2ef0-4f8f-8d1c-908d4e10382f","resolution":{"observed_at":"2026-08-07T11:11:41.103071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T11:11:48.510518Z","title":"Scientific Reports, 15, 1, 2126","venue":null,"work_id":"eb298eb1-3b02-4fd3-954e-3bcfcfe111f5","year":null},"citing_paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-07T11:11:40.881716Z"},"links":{"citing_paper":"/paper/2506.05395"},"observation_digest":"sha256:df7c8c6a8d8da29558cd1de2133534b2a57668026e70d58de8f5d6c3d54a85c6","observation_id":"d6d1f0d3-5f41-47da-b6b7-e20e2ad4ccde","resolution":{"observed_at":"2026-08-07T11:11:48.562185Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-17T06:30:58.91139+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.05395","last_updated":"2025-09-02T17:50:58Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-17T04:14:54.934984Z","submitted_at":"2025-06-03T19:44:49Z","title":"TriPSS: A Tri-Modal Keyframe Extraction Framework Using Perceptual, Structural, and Semantic Representations"},"reference_resolution":{"displayed":65,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":56,"verified_exact":4,"verified_fuzzy":5},"total_outbound_references":65},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-17T06:30:58.91139+00:00","source":"crossref"},{"observed_at":"2026-08-17T06:30:54.323127+00:00","source":"retraction_watch"}],"thesis":"As of 18 August 2026, this Paper Citation Record lists 65 of 65 outbound references and 1 inbound Pith citation observation for arXiv:2506.05395."}