{"as_of":"2026-08-09T14:29:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:fc022e2805deaaa739f14ba63d9e8221e64968241908b8b6888689a347a31b52","coverage":[{"denominator":32,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":32,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-09T10:11:51.810083Z","state":"measured"},{"denominator":35,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":35,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T23:35:06.626592Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.03032","snapshot_observed_at":"2026-08-06T23:35:06.626592Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17673","last_updated":"2025-06-21T10:18:25Z","snapshot_observed_at":"2026-08-09T09:01:02.896089Z","submitted_at":"2025-06-21T10:18:25Z","title":"FaithfulSAE: Towards Capturing Faithful Features with Sparse Autoencoders without External Dataset Dependencies","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-06T23:35:06.626592Z"},"links":{"cited_paper":"/paper/2502.03032","citing_paper":"/paper/2506.17673"},"observation_digest":"sha256:9397cb9655daa49e99eb4a0b93f39fc0c2c3d77b0fd4b4c45b79f6d639e25a1d","observation_id":"ddfb11d5-5e35-48ce-a91a-85dc7bfd7db5","resolution":{"observed_at":"2026-08-06T23:35:06.626592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.03032","snapshot_observed_at":"2026-08-06T23:03:04.984298Z","title":"Jack Lindsey, Adly Templeton, Jonathan Marcus, Tom Conerly, Joshua Batson, and Chris Olah","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.20040","last_updated":"2026-06-09T21:19:14Z","snapshot_observed_at":"2026-08-09T09:00:18.331911Z","submitted_at":"2025-06-24T22:43:36Z","title":"Cross-Layer Discrete Concept Discovery for Interpreting Language Models","version":3},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T23:03:04.984298Z"},"links":{"cited_paper":"/paper/2502.03032","citing_paper":"/paper/2506.20040"},"observation_digest":"sha256:ff0c3e25cda111d7d2643b797dcea92b5a2c2df212e777dcc6beca10d67c4ee1","observation_id":"99c2055d-f6ed-4653-a0f4-30634d1cf8bd","resolution":{"observed_at":"2026-08-06T23:03:04.984298Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"cited_work":{"arxiv_id":"2502.03032","doi":"10.48550/arxiv.2502.03032","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.03032","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"2025 , eprint =","venue":"ArXiv.org","work_id":"2c10f33c-ee58-401e-883a-45c6a0d3377a","year":2025},"citing_paper":{"arxiv_id":"2606.12138","last_updated":"2026-06-10T14:32:57Z","snapshot_observed_at":"2026-08-07T02:28:21.915029Z","submitted_at":"2026-06-10T14:32:57Z","title":"Unstable Features, Reproducible Subspaces: Understanding Seed Dependence in Sparse Autoencoders","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-06-27T10:39:51.615710Z"},"links":{"cited_paper":"/paper/2502.03032","citing_paper":"/paper/2606.12138"},"observation_digest":"sha256:6e80b5d6f59c01c2ca9c64fe35eb5e6a5e3a41271b9d539c602bd09f36589398","observation_id":"9504ce52-0529-4479-a0d4-d0e2c48eabff","resolution":{"observed_at":"2026-06-27T10:40:49.821238Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2502.03032/citation-record","integrity":"/paper/2502.03032/integrity","json":"/paper/2502.03032/citation-record.json","paper":"/paper/2502.03032"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.468511Z","title":"Mechanistic permutability: Match features across layers","venue":null,"work_id":"1b07ce65-4afd-4c6e-bd7b-d38c773c16f4","year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.438449Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:d6ce3faeeb09ed491b7af22e57a0f27a99e8a8841fe04aca6441c53f9988ae55","observation_id":"1b5f570f-fced-4b51-a3a2-55745b003cab","resolution":{"observed_at":"2026-08-09T10:11:52.471703Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.08869","last_updated":"2024-11-17T22:45:45Z","snapshot_observed_at":"2026-08-06T08:52:24.587400Z","submitted_at":"2024-10-11T14:46:49Z","title":"Evolution of SAE Features Across Layers in LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.08869","snapshot_observed_at":"2026-08-09T10:11:51.443311Z","title":"Evolution of sae features across layers in llms, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.443311Z"},"links":{"cited_paper":"/paper/2410.08869","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:ba76ff504ba99603c944ffa5dd329de43682f18421b3eae0e4b6f2099d12531a","observation_id":"8e2a38cf-1e5b-420a-9920-32591196071b","resolution":{"observed_at":"2026-08-09T10:11:51.443311Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:51.447543Z","title":"E., Hume, T., Carter, S., Henighan, T., and Olah, C","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.447543Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:4f70128866e752b759587df02a05f2053f9f908b6c41034b839d786eb9bba26c","observation_id":"02e1601b-20f5-4212-9477-e64947195964","resolution":{"observed_at":"2026-08-09T10:11:51.447543Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.06410","last_updated":"2024-12-09T11:39:00Z","snapshot_observed_at":"2026-08-01T21:49:25.316353Z","submitted_at":"2024-12-09T11:39:00Z","title":"BatchTopK Sparse Autoencoders","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.06410","snapshot_observed_at":"2026-08-09T10:11:51.452310Z","title":"Batchtopk sparse autoencoders","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.452310Z"},"links":{"cited_paper":"/paper/2412.06410","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:c5547ab2f47ea5f54ad7849554d740614370f996b053de99c5e132da88c62c32","observation_id":"4d55c472-3b55-4205-a041-12f6c98efdfd","resolution":{"observed_at":"2026-08-09T10:11:51.452310Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.02193","last_updated":"2024-11-21T12:10:54Z","snapshot_observed_at":"2026-08-09T11:08:29.370811Z","submitted_at":"2024-11-04T15:46:20Z","title":"Improving Steering Vectors by Targeting Sparse Autoencoder Features","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.02193","snapshot_observed_at":"2026-08-09T10:11:51.456407Z","title":"Improving steering vectors by targeting sparse autoencoder features, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.456407Z"},"links":{"cited_paper":"/paper/2411.02193","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:dde783209a8bf05f201ec2e499a57f2a26fc0268de6dede672dfeed8c8a3fd9c","observation_id":"8e0726db-180b-438b-a167-df0d0d9fc048","resolution":{"observed_at":"2026-08-09T10:11:51.456407Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.451785Z","title":"N., Lynch, A., Heimersheim, S., and Garriga-Alonso, A","venue":null,"work_id":"8ece01ec-65c9-413b-a42e-27c2b395d9e7","year":2023},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.461143Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:6d7ff1f3894b4d7c46a63e44ad6cee9bfb8d562830e18396d3deffd831099258","observation_id":"04efe5ec-a6ff-4f72-bb0d-ed8023baf3e8","resolution":{"observed_at":"2026-08-09T10:11:52.455008Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.08600","last_updated":"2023-10-04T13:17:38Z","snapshot_observed_at":"2026-07-06T16:19:05.495349Z","submitted_at":"2023-09-15T17:56:55Z","title":"Sparse Autoencoders Find Highly Interpretable Features in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.08600","snapshot_observed_at":"2026-08-09T10:11:51.465409Z","title":"Sparse autoencoders find highly interpretable features in language models, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.465409Z"},"links":{"cited_paper":"/paper/2309.08600","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:3923e158d99ccd7404d520a69afae110a34672e0fcd50e69420cfb4d32f5a9c4","observation_id":"f1fd3a9a-740a-4107-af88-2d7bb61ce693","resolution":{"observed_at":"2026-08-09T10:11:51.465409Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11944","last_updated":"2024-11-06T22:37:30Z","snapshot_observed_at":"2026-08-04T08:40:10.913790Z","submitted_at":"2024-06-17T17:49:00Z","title":"Transcoders Find Interpretable LLM Feature Circuits","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.11944","snapshot_observed_at":"2026-08-09T10:11:51.469627Z","title":"Transcoders find interpretable llm feature circuits","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.469627Z"},"links":{"cited_paper":"/paper/2406.11944","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:48b46ed7d6e418a67a6194986dcf030bd399cae6f3cf33ab282e63d2b83b0ed9","observation_id":"b7aadacf-6725-4519-b232-02aba34b30f3","resolution":{"observed_at":"2026-08-09T10:11:51.469627Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.07759","last_updated":"2023-05-24T23:30:43Z","snapshot_observed_at":"2026-08-04T10:35:00.917001Z","submitted_at":"2023-05-12T20:56:48Z","title":"TinyStories: How Small Can Language Models Be and Still Speak Coherent English?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.07759","snapshot_observed_at":"2026-08-09T10:11:51.474085Z","title":"and Li, Y","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.474085Z"},"links":{"cited_paper":"/paper/2305.07759","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:d8afc706f3dbb8fb5e96d4baf9c343cec9c17273ea55fea258f1da1ec6122674","observation_id":"eaf86cd0-e337-456e-8a56-29e124281839","resolution":{"observed_at":"2026-08-09T10:11:51.474085Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.440334Z","title":"A mathematical framework for transformer circuits, 2021","venue":null,"work_id":"8bc73d3f-1bb4-476f-bcc6-ca4de0a74115","year":2021},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.491994Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:62857824ef499e7401660a318a058ed4c8b761db63410de3d3227879f225bfa0","observation_id":"7178c337-bf4e-46a1-a006-913578813a84","resolution":{"observed_at":"2026-08-09T10:11:52.444187Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:51.552647Z","title":"J., Liao, I., Gurnee, W., and Tegmark, M","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.552647Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:4bfdfcc17cdf802b5897926a6a4f656854c9688f72bc30d7811076975fa94c4a","observation_id":"369128b1-19b1-450b-a883-e88add6235ad","resolution":{"observed_at":"2026-08-09T10:11:51.552647Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.419952Z","title":"D., Tillman, H., Goh, G., Troll, R., Radford, A., Sutskever, I., Leike, J., and Wu, J","venue":null,"work_id":"d7d11e7b-edbc-484f-a8c5-58781d2355c2","year":2025},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.612910Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:97284cc4a259f3068c40c53c31b0da3f1508b7cf18a1acf332f28a98545d2602","observation_id":"0fbea687-ec13-4507-a204-5298dea91d5b","resolution":{"observed_at":"2026-08-09T10:11:52.424747Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.13868","last_updated":"2024-07-21T11:42:32Z","snapshot_observed_at":"2026-08-02T09:03:28.447459Z","submitted_at":"2024-05-22T17:50:04Z","title":"Automatically Identifying Local and Global Circuits with Linear Computation Graphs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.13868","snapshot_observed_at":"2026-08-09T10:11:51.652517Z","title":"Automatically identifying local and global circuits with linear computation graphs","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.652517Z"},"links":{"cited_paper":"/paper/2405.13868","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:2c868a7fb0f28110a5537016b56ca6bda7b1ee26d04d1f7692cce4e32f0ba617","observation_id":"522e24f5-217d-44a6-b183-1b8aff6f0eab","resolution":{"observed_at":"2026-08-09T10:11:51.652517Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00118","last_updated":"2024-10-02T15:22:49Z","snapshot_observed_at":"2026-08-02T16:20:09.773989Z","submitted_at":"2024-07-31T19:13:07Z","title":"Gemma 2: Improving Open Language Models at a Practical Size","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.00118","snapshot_observed_at":"2026-08-09T10:11:51.712221Z","title":"Gemma 2: Improving open language models at a practical size, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.712221Z"},"links":{"cited_paper":"/paper/2408.00118","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:ea03fb1799806a4ce82822ba766adbd941d23f9df122c5436f3be08e77d93b72","observation_id":"2e7c6fd6-46b2-4f10-8fc0-fe92a225119f","resolution":{"observed_at":"2026-08-09T10:11:51.712221Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2024.blackboxnlp-1.32","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:51.883387Z","title":"Accelerating sparse autoencoder training via layer-wise transfer learning in large language models","venue":null,"work_id":"f262a5e0-ca52-4a3d-8e87-30c097c5a05b","year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.749269Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:e7fdbb7bea6eca048c1ce027dbee5080db9339ce2eaa1faf2360c62adaed7bcb","observation_id":"98b75e5c-0694-4eb6-9e01-3d12c397fe86","resolution":{"observed_at":"2026-08-09T10:11:51.966007Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.02207","last_updated":"2024-03-04T18:25:29Z","snapshot_observed_at":"2026-08-06T19:42:39.893943Z","submitted_at":"2023-10-03T17:06:52Z","title":"Language Models Represent Space and Time","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.02207","snapshot_observed_at":"2026-08-09T10:11:51.752448Z","title":"and Tegmark, M","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.752448Z"},"links":{"cited_paper":"/paper/2310.02207","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:a08f5c762eca2958fa14c5cb8eeeedfdb2e5dfd8ff3b2f3b3ce4dc441d3b50f6","observation_id":"91727d6c-73c4-4a18-b1fc-7659dbd1fad0","resolution":{"observed_at":"2026-08-09T10:11:51.752448Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.407376Z","title":"Finding neurons in a haystack: Case studies with sparse probing","venue":null,"work_id":"1bd72647-e628-4926-991e-fdb400dcb115","year":2023},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.756129Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:682d84f317d989c57f50140ec44015d7c54a4456dcfcd35298a73275f2afae62","observation_id":"7226ec3c-01f7-4cce-abce-cef58a291205","resolution":{"observed_at":"2026-08-09T10:11:52.411395Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.20526","last_updated":"2024-10-27T17:33:49Z","snapshot_observed_at":"2026-08-05T02:09:50.000836Z","submitted_at":"2024-10-27T17:33:49Z","title":"Llama Scope: Extracting Millions of Features from Llama-3.1-8B with Sparse Autoencoders","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.20526","snapshot_observed_at":"2026-08-09T10:11:51.760053Z","title":"Llama scope: Extracting millions of features from llama-3.1-8b with sparse autoencoders","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.760053Z"},"links":{"cited_paper":"/paper/2410.20526","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:dea7df05c1d66ea66e5bfe40e0ff16e7d5181f237946e912535b79a055e8db45","observation_id":"c5dd6416-b854-4c5a-80fe-6d72137b05a6","resolution":{"observed_at":"2026-08-09T10:11:51.760053Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.395542Z","title":"Random open problems","venue":null,"work_id":"5c2c7857-138b-4705-b9ab-a607a789433a","year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.763568Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:3ee36ec778f5b1c60b54deb085c81403ad4b165aa33795aec9f3933f18b0d4e5","observation_id":"d64d0a96-a696-4745-80fe-2af859f96d1f","resolution":{"observed_at":"2026-08-09T10:11:52.399370Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.05147","last_updated":"2024-08-19T07:51:05Z","snapshot_observed_at":"2026-08-04T11:44:14.524984Z","submitted_at":"2024-08-09T16:06:42Z","title":"Gemma Scope: Open Sparse Autoencoders Everywhere All At Once on Gemma 2","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.05147","snapshot_observed_at":"2026-08-09T10:11:51.766856Z","title":"Gemma scope: Open sparse autoencoders everywhere all at once on gemma 2","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.766856Z"},"links":{"cited_paper":"/paper/2408.05147","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:bfdaaacd83e437860063077ec2791d8e98bc4f643819f474668a0a1125e6e01f","observation_id":"3eaf50f9-ab4e-48a8-a132-b22b001828b1","resolution":{"observed_at":"2026-08-09T10:11:51.766856Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.384081Z","title":"Sparse crosscoders for cross-layer features and model diffing, 2024","venue":null,"work_id":"abfcb203-55e7-441d-b3bd-f11d4534fdd5","year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.770388Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:3d2021a73e45abf362d77831d75bcd0db58e3a3077e19cf68a4c60f2baeb322c","observation_id":"c132b4d8-b1b0-47e4-ad7c-4f229fd93528","resolution":{"observed_at":"2026-08-09T10:11:52.388484Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1312.5663","last_updated":"2014-03-22T17:12:07Z","snapshot_observed_at":"2026-08-01T16:20:01.675800Z","submitted_at":"2013-12-19T17:46:46Z","title":"k-Sparse Autoencoders","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1312.5663","snapshot_observed_at":"2026-08-09T10:11:51.773693Z","title":"and Frey, B","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.773693Z"},"links":{"cited_paper":"/paper/1312.5663","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:05fe5bb62091cc2085bbc203916842b0042dfbe925d1d54c2b5dc5ec41976df2","observation_id":"b6f5904a-073e-4b41-8df6-b5af9e22a9b1","resolution":{"observed_at":"2026-08-09T10:11:51.773693Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06824","last_updated":"2024-08-19T01:18:41Z","snapshot_observed_at":"2026-07-06T16:30:37.867641Z","submitted_at":"2023-10-10T17:54:39Z","title":"The Geometry of Truth: Emergent Linear Structure in Large Language Model Representations of True/False Datasets","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06824","snapshot_observed_at":"2026-08-09T10:11:51.777585Z","title":"and Tegmark, M","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.777585Z"},"links":{"cited_paper":"/paper/2310.06824","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:f42f441179878991fec82445ff7293d6849a68063ab61d5142a4aa370b96b157","observation_id":"02ba2ab7-be1f-4936-ab0d-f1fef178fe0b","resolution":{"observed_at":"2026-08-09T10:11:51.777585Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.371708Z","title":"J., Belinkov, Y., Bau, D., and Mueller, A","venue":null,"work_id":"128d5b9a-5297-429d-95fe-c434e44c2db4","year":2025},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.781326Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:5f5f36eb98fb37b15d68c3dbaf0288acf9623fc800bd59eae69a38da17e42f88","observation_id":"93bbec5a-e5f4-4230-90fe-d52523a6e18c","resolution":{"observed_at":"2026-08-09T10:11:52.375920Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.15771","last_updated":"2023-07-28T19:13:26Z","snapshot_observed_at":"2026-08-06T01:59:30.696076Z","submitted_at":"2023-07-28T19:13:26Z","title":"The Hydra Effect: Emergent Self-repair in Language Model Computations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.15771","snapshot_observed_at":"2026-08-09T10:11:51.784564Z","title":"The hydra effect: Emergent self-repair in language model computations, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.784564Z"},"links":{"cited_paper":"/paper/2307.15771","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:8c735fcf73db8f0caeb638713ff0df1099f46e694cff5619f6efd5ed9c6b5f06","observation_id":"61068ed2-927c-4436-a78c-a9784e7cfa1b","resolution":{"observed_at":"2026-08-09T10:11:51.784564Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.360664Z","title":"Linguistic regularities in continuous space word representations","venue":null,"work_id":"a5dfeb03-3635-4d30-a39f-4579a99976be","year":2013},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.788166Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:937b5480a96644c0c72e50bd22f7ab33087e33ebd3286934419833dc611da410","observation_id":"56267dcb-dd2e-4000-a611-11d1caf9b6a5","resolution":{"observed_at":"2026-08-09T10:11:52.364819Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:51.791856Z","title":"B., Lozhkov, A., Mitchell, M., Raffel, C., Werra, L","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.791856Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:49a9894c8a75a250e14933ff0d48a0d599e49918302c852b260324a2726d5ca7","observation_id":"858e0fbc-f649-4593-8231-15c248d939dc","resolution":{"observed_at":"2026-08-09T10:11:51.791856Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.14435","last_updated":"2024-08-01T17:42:04Z","snapshot_observed_at":"2026-08-08T18:10:15.134088Z","submitted_at":"2024-07-19T16:07:19Z","title":"Jumping Ahead: Improving Reconstruction Fidelity with JumpReLU Sparse Autoencoders","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.14435","snapshot_observed_at":"2026-08-09T10:11:51.796172Z","title":"Jumping ahead: Improving reconstruction fidelity with jumprelu sparse autoencoders","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.796172Z"},"links":{"cited_paper":"/paper/2407.14435","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:406d468a5701aa26ef6a08a8415fd8d44c09d366f5d8f9b375725a002b429cdf","observation_id":"5e45417b-291a-4d4a-bced-9d75e75a561d","resolution":{"observed_at":"2026-08-09T10:11:51.796172Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:51.799408Z","title":"L., McDougall, C., MacDiarmid, M., Freeman, C","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.799408Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:6737dd4f84f1778a707028f0a49ecd7f6cf4b626ee6ec6175ddbbe8630906326","observation_id":"b6425fa9-13d0-4158-93f9-df0ecb62edbb","resolution":{"observed_at":"2026-08-09T10:11:51.799408Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:52.336989Z","title":"Towards universality: Studying mechanistic similarity across language model architectures","venue":null,"work_id":"6ff3de2c-8832-4787-831d-0cbca6358837","year":2025},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.803092Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:e58eeef86d3b43837bd9b781c7448ecbb590e8cd1248400fa121296bea580706","observation_id":"ad94b412-6405-4d55-bf05-16fa3d71fbee","resolution":{"observed_at":"2026-08-09T10:11:52.340927Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.07625","last_updated":"2025-07-22T09:17:48Z","snapshot_observed_at":"2026-07-06T17:28:52.196375Z","submitted_at":"2024-02-12T13:09:21Z","title":"Autonomous Data Selection with Zero-shot Generative Classifiers for Mathematical Texts","version":7},"cited_work":{"arxiv_id":"2402.07625","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.07625","snapshot_observed_at":"2026-08-09T10:11:52.071266Z","title":"Autonomous Data Selection with Zero-shot Generative Classifiers for Mathematical Texts","venue":"cs.CL","work_id":"b27b9736-9f8a-443a-8f3d-8fa5457b58b5","year":2024},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.806394Z"},"links":{"cited_paper":"/paper/2402.07625","citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:4ba8cfffc7a78d9b837d2739dbebecb02440e73a0c4ee53d10fb4278851e886b","observation_id":"4fe54092-f937-4040-8c7d-294d134c76d3","resolution":{"observed_at":"2026-08-09T10:11:52.174690Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-09T10:11:51.810083Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models","version":3},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-09T10:11:51.810083Z"},"links":{"citing_paper":"/paper/2502.03032"},"observation_digest":"sha256:7edc580f0d2ac4630ecf4524ee5009b011b6199954e6793647b3fcc567e597fc","observation_id":"89cf2cd1-53ed-49ba-ae76-d9963ae7c568","resolution":{"observed_at":"2026-08-09T10:11:51.810083Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2502.03032","last_updated":"2025-07-24T21:16:27Z","latest_version":3,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-09T11:09:56.472014Z","submitted_at":"2025-02-05T09:39:34Z","title":"Analyze Feature Flow to Enhance Interpretation and Steering in Language Models"},"reference_resolution":{"displayed":32,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":20,"verified_exact":2,"verified_fuzzy":10},"total_outbound_references":32},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 32 of 32 outbound references and 3 inbound Pith citation observations for arXiv:2502.03032."}