{"as_of":"2026-08-20T07:38:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:76fb6afb41f6b7d349921660150ee3aba516d462d3105a423003b42b44815bc0","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":6,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":6,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":6,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":6,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T20:47:06.627050Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.11695","last_updated":"2025-08-20T09:43:41Z","snapshot_observed_at":"2026-08-16T12:40:39.959641Z","submitted_at":"2025-04-16T01:40:06Z","title":"Interpreting the linear structure of vision-language model embedding spaces","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.11695","snapshot_observed_at":"2026-08-06T20:47:06.627050Z","title":"Interpreting the linear structure of vision-language model embedding spaces","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.01790","last_updated":"2025-07-02T15:15:14Z","snapshot_observed_at":"2026-08-18T14:40:27.199280Z","submitted_at":"2025-07-02T15:15:14Z","title":"How Do Vision-Language Models Process Conflicting Information Across Modalities?","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-06T20:47:06.627050Z"},"links":{"cited_paper":"/paper/2504.11695","citing_paper":"/paper/2507.01790"},"observation_digest":"sha256:9b17f357297caa9d598c50bd588b42d6bb92b0309a08237fd78facbba96e30c3","observation_id":"f206f484-5183-40d2-a623-9fcab1b725fd","resolution":{"observed_at":"2026-08-06T20:47:06.627050Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.11695","last_updated":"2025-08-20T09:43:41Z","snapshot_observed_at":"2026-08-16T12:40:39.959641Z","submitted_at":"2025-04-16T01:40:06Z","title":"Interpreting the linear structure of vision-language model embedding spaces","version":4},"cited_work":{"arxiv_id":"2504.11695","doi":"10.48550/arxiv.2504.11695","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.11695","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"In- terpreting the linear structure of vision-language model embedding spaces","venue":"ArXiv.org","work_id":"8f96bad0-8e97-40bf-9976-b3a6b2fd1d77","year":2025},"citing_paper":{"arxiv_id":"2604.14363","last_updated":"2026-08-08T10:21:17Z","snapshot_observed_at":"2026-08-15T06:52:09.603067Z","submitted_at":"2026-04-15T19:26:30Z","title":"The Cost of Language: Centroid Erasure Exposes and Exploits Modal Competition in Multimodal Language Models","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-10T13:25:27.762910Z"},"links":{"cited_paper":"/paper/2504.11695","citing_paper":"/paper/2604.14363"},"observation_digest":"sha256:0a5f9344716074d227e4f7ea056685be91283d3a9a7a01a7241e8a306bbc7ac7","observation_id":"0f0191de-dad2-4a2c-86bc-443d10e75f95","resolution":{"observed_at":"2026-05-10T13:35:26.787193Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.11695","last_updated":"2025-08-20T09:43:41Z","snapshot_observed_at":"2026-08-16T12:40:39.959641Z","submitted_at":"2025-04-16T01:40:06Z","title":"Interpreting the linear structure of vision-language model embedding spaces","version":4},"cited_work":{"arxiv_id":"2504.11695","doi":"10.48550/arxiv.2504.11695","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.11695","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"In- terpreting the linear structure of vision-language model embedding spaces","venue":"ArXiv.org","work_id":"8f96bad0-8e97-40bf-9976-b3a6b2fd1d77","year":2025},"citing_paper":{"arxiv_id":"2605.06640","last_updated":"2026-05-07T17:51:13Z","snapshot_observed_at":"2026-08-15T01:54:05.211553Z","submitted_at":"2026-05-07T17:51:13Z","title":"Concept-Based Abductive and Contrastive Explanations for Behaviors of Vision Models","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-08T12:05:51.787664Z"},"links":{"cited_paper":"/paper/2504.11695","citing_paper":"/paper/2605.06640"},"observation_digest":"sha256:b21a0199b919506025b345a1c97c50a6f132f586dae5c1d54633214f8c2849aa","observation_id":"037f391e-0b7b-4463-8738-050cab934dbe","resolution":{"observed_at":"2026-05-11T19:21:09.414903Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.11695","last_updated":"2025-08-20T09:43:41Z","snapshot_observed_at":"2026-08-16T12:40:39.959641Z","submitted_at":"2025-04-16T01:40:06Z","title":"Interpreting the linear structure of vision-language model embedding spaces","version":4},"cited_work":{"arxiv_id":"2504.11695","doi":"10.48550/arxiv.2504.11695","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.11695","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"In- terpreting the linear structure of vision-language model embedding spaces","venue":"ArXiv.org","work_id":"8f96bad0-8e97-40bf-9976-b3a6b2fd1d77","year":2025},"citing_paper":{"arxiv_id":"2605.11107","last_updated":"2026-05-11T18:13:05Z","snapshot_observed_at":"2026-08-11T03:45:15.335189Z","submitted_at":"2026-05-11T18:13:05Z","title":"Birds of a Feather Flock Together: Background-Invariant Representations via Linear Structure in VLMs","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-13T07:26:02.947081Z"},"links":{"cited_paper":"/paper/2504.11695","citing_paper":"/paper/2605.11107"},"observation_digest":"sha256:92d79199ade0d3489c53c415e3bec9c991474f1407a0a30ae7ad9540f78b4013","observation_id":"9044e343-c697-4cfc-95b9-730378362885","resolution":{"observed_at":"2026-05-13T07:27:29.719184Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.11695","last_updated":"2025-08-20T09:43:41Z","snapshot_observed_at":"2026-08-16T12:40:39.959641Z","submitted_at":"2025-04-16T01:40:06Z","title":"Interpreting the linear structure of vision-language model embedding spaces","version":4},"cited_work":{"arxiv_id":"2504.11695","doi":"10.48550/arxiv.2504.11695","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.11695","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"In- terpreting the linear structure of vision-language model embedding spaces","venue":"ArXiv.org","work_id":"8f96bad0-8e97-40bf-9976-b3a6b2fd1d77","year":2025},"citing_paper":{"arxiv_id":"2606.30815","last_updated":"2026-06-29T18:42:03Z","snapshot_observed_at":"2026-08-12T19:50:48.405292Z","submitted_at":"2026-06-29T18:42:03Z","title":"When transformers learn \"impossible\" languages, what do they learn?","version":1},"reference_index":227,"source":"arxiv_source","source_observed_at":"2026-07-01T02:13:58.839175Z"},"links":{"cited_paper":"/paper/2504.11695","citing_paper":"/paper/2606.30815"},"observation_digest":"sha256:345efe7731cbe26286eb472a209683c9b67003933f62ba81ae2aa44e7fc1c040","observation_id":"2dccbba9-2db7-4122-9b51-1dfa0d37078a","resolution":{"observed_at":"2026-07-01T02:15:14.343884Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-20T06:33:59.587034+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.11695","last_updated":"2025-08-20T09:43:41Z","snapshot_observed_at":"2026-08-16T12:40:39.959641Z","submitted_at":"2025-04-16T01:40:06Z","title":"Interpreting the linear structure of vision-language model embedding spaces","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.11695","snapshot_observed_at":"2026-08-01T03:02:06.469837Z","title":"Interpreting the linear structure of vision-language model embedding spaces , url =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.25271","last_updated":"2026-07-28T04:18:49Z","snapshot_observed_at":"2026-08-13T02:12:02.677620Z","submitted_at":"2026-07-28T04:18:49Z","title":"Bridging Compute- and Data-Optimal Pretraining","version":1},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-08-01T03:02:06.469837Z"},"links":{"cited_paper":"/paper/2504.11695","citing_paper":"/paper/2607.25271"},"observation_digest":"sha256:19cfb7de927c2f8eeecefc38db82960d03ad152b5c16866a40ce9b78f6887935","observation_id":"1a17e2ba-4f42-4226-a144-4904e2615cc4","resolution":{"observed_at":"2026-08-01T03:02:06.469837Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2504.11695/citation-record","integrity":"/paper/2504.11695/integrity","json":"/paper/2504.11695/citation-record.json","paper":"/paper/2504.11695"},"outbound":[],"paper":{"arxiv_id":"2504.11695","last_updated":"2025-08-20T09:43:41Z","latest_version":4,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-16T12:40:39.959641Z","submitted_at":"2025-04-16T01:40:06Z","title":"Interpreting the linear structure of vision-language model embedding spaces"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 6 inbound Pith citation observations for arXiv:2504.11695."}