{"as_of":"2026-08-21T09:01:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:316c904469470a3e9822f80dd0ee8feedf6f35859ef4b8cb82238d3609ab2974","coverage":[{"denominator":45,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":45,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-15T20:57:38.204999Z","state":"measured"},{"denominator":49,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":49,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":4,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":4,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T00:22:51.856819Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-01T14:05:47.019328Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"cited_work":{"arxiv_id":"2505.11615","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.11615","snapshot_observed_at":"2026-07-01T14:05:47.019328Z","title":"Steering risk preferences in large language models by aligning behavioral and neural representations.arXiv preprint arXiv:2505.11615, 2025","venue":null,"work_id":"587b9f18-5be0-48e2-9140-aafb14adbfed","year":2025},"citing_paper":{"arxiv_id":"2605.08556","last_updated":"2026-05-08T23:26:35Z","snapshot_observed_at":"2026-08-12T22:52:04.451328Z","submitted_at":"2026-05-08T23:26:35Z","title":"Can Revealed Preferences Clarify LLM Alignment and Steering?","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-12T02:10:50.298626Z"},"links":{"cited_paper":"/paper/2505.11615","citing_paper":"/paper/2605.08556"},"observation_digest":"sha256:ab7250c7af2db0f12fa0d158abdd2db513de8becd584b04f5aee26480806bddd","observation_id":"32a752d9-d242-4873-85e1-975b11a05e71","resolution":{"observed_at":"2026-05-12T02:11:15.341384Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"cited_work":{"arxiv_id":"2505.11615","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.11615","snapshot_observed_at":"2026-07-01T14:05:47.019328Z","title":"Steering risk preferences in large language models by aligning behavioral and neural representations.arXiv preprint arXiv:2505.11615, 2025","venue":null,"work_id":"587b9f18-5be0-48e2-9140-aafb14adbfed","year":2025},"citing_paper":{"arxiv_id":"2606.05194","last_updated":"2026-07-08T19:54:16Z","snapshot_observed_at":"2026-08-06T06:51:53.907927Z","submitted_at":"2026-05-11T21:09:00Z","title":"Temporal Preference Concepts and their Functions in a Large Language Model","version":1},"reference_index":123,"source":"pdf_text","source_observed_at":"2026-06-30T22:16:47.743387Z"},"links":{"cited_paper":"/paper/2505.11615","citing_paper":"/paper/2606.05194"},"observation_digest":"sha256:55b3456ca224a315c3c42728e6c2593006aed6dac0b985649eee51dc170b80f2","observation_id":"2f072f0e-991e-4ae4-b59e-6c463f16f722","resolution":{"observed_at":"2026-07-01T14:05:47.020963Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.11615","snapshot_observed_at":"2026-07-12T17:03:44.315006Z","title":"Steering risk preferences in large language models by aligning behavioral and neural representations.arXiv preprint arXiv:2505.11615, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2606.05194","last_updated":"2026-07-08T19:54:16Z","snapshot_observed_at":"2026-08-06T06:51:53.907927Z","submitted_at":"2026-05-11T21:09:00Z","title":"Temporal Preference Concepts and their Functions in a Large Language Model","version":2},"reference_index":123,"source":"pdf_text","source_observed_at":"2026-07-12T17:03:44.315006Z"},"links":{"cited_paper":"/paper/2505.11615","citing_paper":"/paper/2606.05194"},"observation_digest":"sha256:cb53dcf10bea2a17651a2c649094050dfb53a4f9a699d52a015b5335ebd1d49d","observation_id":"d7d3974d-23a8-4a40-9828-b8d6b5b99a96","resolution":{"observed_at":"2026-07-12T17:03:44.315006Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.11615","snapshot_observed_at":"2026-08-06T00:22:51.856819Z","title":"Griffiths","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.01361","last_updated":"2026-08-02T16:25:08Z","snapshot_observed_at":"2026-08-20T02:29:17.423934Z","submitted_at":"2026-08-02T16:25:08Z","title":"High-Stakes Decisions with Language Models: Insights from Emergency Triage","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T00:22:51.856819Z"},"links":{"cited_paper":"/paper/2505.11615","citing_paper":"/paper/2608.01361"},"observation_digest":"sha256:8780a6e893a328ba043b7261ce1e260d02749f2090e953ab3c1c13a4cde9d853","observation_id":"1ba6ebd8-f0f1-4214-8f72-164936310690","resolution":{"observed_at":"2026-08-06T00:22:51.856819Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2505.11615/citation-record","integrity":"/paper/2505.11615/integrity","json":"/paper/2505.11615/citation-record.json","paper":"/paper/2505.11615"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.695619Z","title":"Leveraging large language models to estimate clinically relevant psychological constructs in psychotherapy transcripts.medRxiv, pages 2025–03, 2025","venue":null,"work_id":"b82fabdf-ab22-47d3-a401-3c885f4b4d4c","year":2025},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.042071Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:b3415743e034df17ad1c463578a6483949be6cc8dbfa357b1c920fba0a53ba4d","observation_id":"dd85b3ef-2114-4bed-a22c-db25d866730b","resolution":{"observed_at":"2026-08-15T20:57:38.700089Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.684438Z","title":"Monte Carlo calculations of the radial distribution functions for a proton-electron plasma.Australian Journal of Physics, 18(2):119–134, 1965","venue":null,"work_id":"8746f4c4-7e5c-4e76-b47c-2d85997921f0","year":1965},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.046403Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:ddee4e6b465b5562d97995cd7a596c5e934057c4e6d8dbe78beb392aaabe1447","observation_id":"9d336197-5dc4-48a8-8641-d8dc879ead57","resolution":{"observed_at":"2026-08-15T20:57:38.687925Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.673758Z","title":"Exploring variability in risk taking with large language models.Journal of Experimental Psychology: General, 153(7):1838, 2024","venue":null,"work_id":"209887b1-c277-497f-b441-cfd8c9e0b194","year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.050588Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:2868319fc1d47ef89e8619bc02417fbbfdcd220ee507abd750405b5bc76fb88f","observation_id":"491adfb3-b5de-43da-9844-e8724392b84e","resolution":{"observed_at":"2026-08-15T20:57:38.677602Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.054128Z","title":"Language models are few-shot learners.Advances in Neural Information Processing Systems, 33:1877–1901, 2020","venue":null,"work_id":null,"year":1901},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.054128Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:af4d96def54cb10f341277abaf336c50d021c1bd8888482fe19e7fd4b901de91","observation_id":"9970f6f9-2d5f-4363-b3d8-64d433320fd1","resolution":{"observed_at":"2026-08-15T20:57:38.054128Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17284","last_updated":"2025-05-28T17:16:40Z","snapshot_observed_at":"2026-08-20T11:17:51.979203Z","submitted_at":"2024-11-26T10:13:39Z","title":"AutoElicit: Using Large Language Models for Expert Prior Elicitation in Predictive Modelling","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17284","snapshot_observed_at":"2026-08-15T20:57:38.058612Z","title":"Using large language models for expert prior elicitation in predictive modelling.arXiv preprint arXiv:2411.17284, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.058612Z"},"links":{"cited_paper":"/paper/2411.17284","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:f0d2c1d27a7c2b9a773c2944e72974c7dad1f8861e3dc10114cff4f30e298414","observation_id":"9edd798b-3440-44c4-b33b-8cb91d5eb816","resolution":{"observed_at":"2026-08-15T20:57:38.058612Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.656105Z","title":"Bias correction of learned generative models using likelihood-free importance weighting.Advances in Neural Information Processing Systems, 32, 2019","venue":null,"work_id":"c817c454-cbe4-4bd4-9738-281d4fca89dc","year":2019},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.064913Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:1b872c98f8baae2541628f2a255bf0d057c33d8bb6ec2058b7f47b1c41c2df74","observation_id":"628dd1f0-b78a-4343-aa27-1a5505850a73","resolution":{"observed_at":"2026-08-15T20:57:38.660166Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.643348Z","title":"Predictions about indifference curves inside the unit triangle: A test of variants of expected utility theory.Journal of Economic Behavior & Organization, 18(3):391–414, 1992","venue":null,"work_id":"2dc9ccce-29d1-4d0f-a80a-a92b9bcd8e0d","year":1992},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.069914Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:471ef5c56414fe0d0611e2c2d37aa85500858f29502ab60528975cba67802f87","observation_id":"9a5b210b-6aac-4a71-ac72-d4e495a947d2","resolution":{"observed_at":"2026-08-15T20:57:38.648332Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.631967Z","title":"Gibbs sampling with people","venue":null,"work_id":"dac29214-7486-4eea-91c2-139d71aba3fc","year":2020},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.073058Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:fe94c893a8f57253f1f507bca5b5b30c7a070211c46e610541e3491434a3fa9c","observation_id":"513b9545-ad28-40c2-9084-fedc9abbe8e1","resolution":{"observed_at":"2026-08-15T20:57:38.635550Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.076234Z","title":"Taylor & Francis, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.076234Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:9e558345e25cc01caa75be98d5abf606dc492ef4cbb505b977c137d4847a2d30","observation_id":"6e124fb8-f600-48d1-b1fe-92cf4a1f5d21","resolution":{"observed_at":"2026-08-15T20:57:38.076234Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2008.02275","last_updated":"2023-02-17T16:08:22Z","snapshot_observed_at":"2026-08-14T10:58:48.301354Z","submitted_at":"2020-08-05T17:59:16Z","title":"Aligning AI With Shared Human Values","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2008.02275","snapshot_observed_at":"2026-08-15T20:57:38.079833Z","title":"Aligning AI with shared human values.arXiv preprint arXiv:2008.02275, 2020","venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.079833Z"},"links":{"cited_paper":"/paper/2008.02275","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:83626a0a9f4625bf1bdbadf89513cafeeb9137d8ea565763195202a8915222dc","observation_id":"5579810b-25df-409b-b350-7765dbb32bcc","resolution":{"observed_at":"2026-08-15T20:57:38.079833Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.00740","last_updated":"2024-08-09T22:00:11Z","snapshot_observed_at":"2026-08-16T15:43:19.736565Z","submitted_at":"2023-04-03T06:24:10Z","title":"Inspecting and Editing Knowledge Representations in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.00740","snapshot_observed_at":"2026-08-15T20:57:38.083209Z","title":"Inspecting and editing knowledge represen- tations in language models.arXiv preprint arXiv:2304.00740, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.083209Z"},"links":{"cited_paper":"/paper/2304.00740","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:b9b3cc2f6196a2a58b42156465b563d5e1c7744c28402a220d35b0b36dbd94b9","observation_id":"5515d2ce-f803-4ad0-b457-2706ac816d63","resolution":{"observed_at":"2026-08-15T20:57:38.083209Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.613588Z","title":"Prospect theory: An analysis of decision under risk","venue":null,"work_id":"3c7e5958-2794-4917-86cf-01324803f048","year":1979},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.087197Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:98eab05a584560aae013d86f391274006636294d1c8495541788b23f75fec532","observation_id":"7e3f2d8a-9d18-4ab9-918c-4657f35ada45","resolution":{"observed_at":"2026-08-15T20:57:38.617228Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.603194Z","title":"World color survey","venue":null,"work_id":"5708395e-270c-4024-b55e-39f51264ac69","year":2023},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.091077Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:62e96c51edc14781e68b108052166ad1d7e202f421ab9db4c5bacf3572ff417c","observation_id":"43e183f1-e858-45c7-9423-b4309c5496f3","resolution":{"observed_at":"2026-08-15T20:57:38.607062Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.590971Z","title":"Uncovering category representations with linked mcmc with people","venue":null,"work_id":"a769262d-1d64-4991-a728-c2997b313de8","year":2020},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.094137Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:a744623a98b1cc082f50b72530eab148e80daf0c7495b17400ea72b8f63b18ab","observation_id":"f7541a21-716f-47b9-a744-bec146441142","resolution":{"observed_at":"2026-08-15T20:57:38.595487Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.098481Z","title":"Inference- time intervention: Eliciting truthful answers from a language model.Advances in Neural Information Processing Systems, 36:41451–41530, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.098481Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:03a8ef61bd806f22093644babff81ba2e2e4981eabb9544fdaed222080bead3b","observation_id":"f2d88693-a9b2-4ff5-9c35-473ace777cb5","resolution":{"observed_at":"2026-08-15T20:57:38.098481Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.17055","last_updated":"2025-03-10T17:42:37Z","snapshot_observed_at":"2026-08-19T03:38:15.565499Z","submitted_at":"2024-06-24T18:15:27Z","title":"Large Language Models Assume People are More Rational than We Really are","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.17055","snapshot_observed_at":"2026-08-15T20:57:38.101863Z","title":"Large language models assume people are more rational than we really are.arXiv preprint arXiv:2406.17055, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.101863Z"},"links":{"cited_paper":"/paper/2406.17055","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:759f9e4cf357d9cb09a720b18221affbcf876c151a6b0124cd25bfe301257043","observation_id":"e0138723-636a-467a-8fbd-03a409ea5537","resolution":{"observed_at":"2026-08-15T20:57:38.101863Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.105592Z","title":"Wiley New York, 1959","venue":null,"work_id":null,"year":1959},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.105592Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:406b38d99d001e471679b9fd391eac21ff587648b7d388aae2f2a78b4c87767e","observation_id":"4da69b50-0ac9-4237-833c-007cc27087da","resolution":{"observed_at":"2026-08-15T20:57:38.105592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.564239Z","title":"Expected utility","venue":null,"work_id":"c8c605d9-d1a3-43ba-9ae5-233b7a8536bf","year":1982},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.109305Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:aa1d24c32591e171b137079666af2dca0959b61d40614818c920c3f3a6fc708d","observation_id":"84af5b37-0bdf-4ff8-8973-1d5e35f0032c","resolution":{"observed_at":"2026-08-15T20:57:38.569098Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.552708Z","title":"Large language models predict human sensory judgments across six modalities.Scientific Reports, 14(1):21445, 2024","venue":null,"work_id":"6d1f99fb-0185-4268-86bc-29be168d79e9","year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.112066Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:d6183ad15a881b6bf83fc578e94689195dacf7f2cb320f8036644ef80db6cde2","observation_id":"4357e74d-122a-4df6-b377-42396c09ef44","resolution":{"observed_at":"2026-08-15T20:57:38.557290Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.541376Z","title":"Rational behavior, uncertain prospects, and measurable utility.Econometrica, 18(2):111–141, 1950","venue":null,"work_id":"a55e74ed-6e06-44a2-bee9-6877e58cefbf","year":1950},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.115726Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:103d8b88bdc14aa62c5cc50a018fab51a64ccf77afb830069c5b41b9776e9545","observation_id":"f7ab2a3e-7a5f-4b11-ab4d-cd67f32cd36a","resolution":{"observed_at":"2026-08-15T20:57:38.544815Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.08640","last_updated":"2025-02-19T06:48:30Z","snapshot_observed_at":"2026-08-16T12:58:16.280754Z","submitted_at":"2025-02-12T18:55:43Z","title":"Utility Engineering: Analyzing and Controlling Emergent Value Systems in AIs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.08640","snapshot_observed_at":"2026-08-15T20:57:38.118654Z","title":"Utility engineering: Analyzing and controlling emergent value systems in ais.arXiv preprint arXiv:2502.08640, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.118654Z"},"links":{"cited_paper":"/paper/2502.08640","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:3d4ecd6f7d1c45fef518d621163c024730e7b6528bbcf02fe97ef98bb2797d03","observation_id":"e1181c1e-6331-46d2-9a3b-8ba31c27d0c4","resolution":{"observed_at":"2026-08-15T20:57:38.118654Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.530511Z","title":"Gpt as a financial advisor.Available at SSRN 4384861, 2023","venue":null,"work_id":"434650db-467c-479a-82da-d8ed5fe4bc9b","year":2023},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.121881Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:fcd2d228724abd0d407b3369272fb7bc1e4c7a174a838ff8991aef6501c1fd6e","observation_id":"e98bc4ee-b57b-4b8f-95c5-093787e5cf02","resolution":{"observed_at":"2026-08-15T20:57:38.533807Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.518970Z","title":"Non-parametric estimation of the individual’s utility map","venue":null,"work_id":"126c70f7-aac8-482f-b6aa-ec0ff2019027","year":2013},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.125905Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:665630ce47f1977227e71b34ee2681eb6a5abfb81d21a02cce687046338ffdab","observation_id":"04ed913b-b03c-4bf5-b4b3-38238d8b3223","resolution":{"observed_at":"2026-08-15T20:57:38.523046Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.06681","last_updated":"2024-07-05T15:30:45Z","snapshot_observed_at":"2026-08-15T08:04:06.283165Z","submitted_at":"2023-12-09T04:40:46Z","title":"Steering Llama 2 via Contrastive Activation Addition","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.06681","snapshot_observed_at":"2026-08-15T20:57:38.128906Z","title":"Steering Llama 2 via contrastive activation addition.arXiv preprint arXiv:2312.06681, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.128906Z"},"links":{"cited_paper":"/paper/2312.06681","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:f1e1b6cdec0510e7452960e9df3de82b7d56a4c3a83d2148bcbaa5fc7fa18f19","observation_id":"435a058b-d534-436f-9316-e9769dc7ede7","resolution":{"observed_at":"2026-08-15T20:57:38.128906Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.03693","last_updated":"2023-10-05T17:12:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-05T17:12:17Z","title":"Fine-tuning Aligned Language Models Compromises Safety, Even When Users Do Not Intend To!","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.03693","snapshot_observed_at":"2026-08-15T20:57:38.132141Z","title":"Fine-tuning aligned language models compromises safety, even when users do not intend to! arXiv preprint arXiv:2310.03693, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.132141Z"},"links":{"cited_paper":"/paper/2310.03693","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:963be4e3ec06ef2a255d8d71ffb8eb6a8280de6733f8cc59eeacd770344d379f","observation_id":"a75824ed-ba71-43d9-96df-32aca8fa9bce","resolution":{"observed_at":"2026-08-15T20:57:38.132141Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.135468Z","title":"Direct preference optimization: Your language model is secretly a reward model","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.135468Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:5e4a946455ad5ac97756d16b8d34211436ef2dda6b8b8e89fb455f8819699745","observation_id":"52c8d9b4-edc8-4b89-a097-21daa6f8feb3","resolution":{"observed_at":"2026-08-15T20:57:38.135468Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.501532Z","title":"Viking, 2022","venue":null,"work_id":"ea3ce9b9-b72f-494c-be02-e15ccfea9e3b","year":2022},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.139073Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:a7ab6163d53f9f6283bf91d51eb50f7f64d13d3ad224ae3105135723ffe69de6","observation_id":"25c582c5-ad36-4528-90fe-e263ba9c3845","resolution":{"observed_at":"2026-08-15T20:57:38.504982Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.490619Z","title":"Markov chain Monte Carlo with people.Advances in Neural Information Processing Systems, 20, 2007","venue":null,"work_id":"1d9238cb-be84-42a7-bf4a-0fddf9e93110","year":2007},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.142968Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:54f6bae45132c499597e15b3e60905f7a8097ad4829fa11089e4a58d94dc878e","observation_id":"d6ef7362-a2a1-4a06-9aed-c207b34eeb73","resolution":{"observed_at":"2026-08-15T20:57:38.494618Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.480227Z","title":"Perception of risk.Science, 236(4799):280–285, 1987","venue":null,"work_id":"d128c058-8769-4abc-b4a2-de41f0901de8","year":1987},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.146811Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:4d3cd42a8f77514b82dbc315ad49840d2351f5e13dea4e421934f74ab8cdb212","observation_id":"72d29b75-4b77-4580-8533-03a80085e19f","resolution":{"observed_at":"2026-08-15T20:57:38.484075Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.13018","last_updated":"2024-11-26T10:31:41Z","snapshot_observed_at":"2026-08-16T14:50:58.892261Z","submitted_at":"2023-10-18T17:47:58Z","title":"Getting aligned on representational alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.13018","snapshot_observed_at":"2026-08-15T20:57:38.150112Z","title":"Getting aligned on representa- tional alignment.arXiv preprint arXiv:2310.13018, 2023","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.150112Z"},"links":{"cited_paper":"/paper/2310.13018","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:8621c1b8c8ff849872dfe324003a91570979f37baf08cc92b21368e5a4b6d77f","observation_id":"a79b56a1-4c54-4267-9223-c094bd670ff1","resolution":{"observed_at":"2026-08-15T20:57:38.150112Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00118","last_updated":"2024-10-02T15:22:49Z","snapshot_observed_at":"2026-08-21T05:15:28.840279Z","submitted_at":"2024-07-31T19:13:07Z","title":"Gemma 2: Improving Open Language Models at a Practical Size","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.00118","snapshot_observed_at":"2026-08-15T20:57:38.154334Z","title":"Gemma 2: Improving open language models at a practical size.arXiv preprint arXiv:2408.00118, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.154334Z"},"links":{"cited_paper":"/paper/2408.00118","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:969edab76fcadf2848b1fdc89212576dc3b1ef776646752660527061dd610d52","observation_id":"b32b2763-3609-4230-9030-784b61e93d04","resolution":{"observed_at":"2026-08-15T20:57:38.154334Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.10248","last_updated":"2024-10-10T13:20:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-08-20T12:21:05Z","title":"Steering Language Models With Activation Engineering","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.10248","snapshot_observed_at":"2026-08-15T20:57:38.158133Z","title":"Steering language models with activation engineering.arXiv preprint arXiv:2308.10248, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.158133Z"},"links":{"cited_paper":"/paper/2308.10248","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:d31017e05741ce2ba1ce6ffd3d8cf40d3ca893ae845612cbdf37ca1a6de8ea13","observation_id":"e33ccde9-2538-4967-ba28-375892034830","resolution":{"observed_at":"2026-08-15T20:57:38.158133Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.469925Z","title":"Advances in prospect theory: Cumulative representation of uncertainty.Journal of Risk and Uncertainty, 5:297–323, 1992","venue":null,"work_id":"2f77e34f-ff0c-436a-b308-e412ab4c7f41","year":1992},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.161695Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:ae4cae296e31e534a92572be5406d081affd61846facf61754b08a952307a74d","observation_id":"cb847dbe-b2fb-4a39-b5e5-c0067e34f3c7","resolution":{"observed_at":"2026-08-15T20:57:38.473540Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.458104Z","title":"The time course of visual processing: from early perception to decision-making.Journal of Cognitive Neuroscience, 13(4):454–461, 2001","venue":null,"work_id":"671334c5-11b9-48e8-a4e1-4ce12bc1fd1c","year":2001},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.165064Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:9edbde8a6937a38a353e8ea4bfc94dcdee6e81f889f1900e08a015a98726e79d","observation_id":"e2f8b3db-87c2-418a-816e-0380a13e67b2","resolution":{"observed_at":"2026-08-15T20:57:38.462495Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.446485Z","title":"Prince- ton University Press, 1947","venue":null,"work_id":"e2a7c3ab-a179-4e61-9be5-02a1454360f6","year":1947},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.168534Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:a1debb5b4e5f53fc5844de33d5cfc6c8825155ac82bd375d19fc7e78108d261f","observation_id":"9ad7c7e7-961e-4a70-8703-b25389b0604e","resolution":{"observed_at":"2026-08-15T20:57:38.450424Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.434543Z","title":"A domain-specific risk-attitude scale: Mea- suring risk perceptions and risk behaviors.Journal of Behavioral Decision Making, 15(4):263– 290, 2002","venue":null,"work_id":"88fdfeb8-204f-434a-b7b4-730beb4375c3","year":2002},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.172381Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:0555160f2e88ade49fe13cfa0c516ee32dd4bf1315647e7351c54e95b7a6a27a","observation_id":"eecdb662-3c26-464c-965a-adfcf7a37608","resolution":{"observed_at":"2026-08-15T20:57:38.439061Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.420058Z","title":"Common consequence conditions in decision making under risk.Journal of Risk and Uncertainty, 16:115–139, 1998","venue":null,"work_id":"62202133-9250-4cf1-ae34-ff5729c0922b","year":1998},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.175758Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:786c06803438b0c51490adbf623e5616d05441dffea2978eaafb55079dabb8aa","observation_id":"d4fc2a28-f052-4651-aa99-1c47be2cae07","resolution":{"observed_at":"2026-08-15T20:57:38.424881Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.406309Z","title":"Quickly recovering comprehensive individual mental representations of facial affect.OSF, 2024","venue":null,"work_id":"b4e0eefe-6238-42a7-b766-128ce6b0af9c","year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.179831Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:23c8f07594b199e9dcf0a7b6c9054f540742d810e46974161df62d63a4212b3d","observation_id":"1ea19ea2-96d9-4cc5-b246-25d5578a29e5","resolution":{"observed_at":"2026-08-15T20:57:38.410853Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.394882Z","title":"Tree of thoughts: Deliberate problem solving with large language models","venue":null,"work_id":"a4d5a063-581e-4904-bad7-aef9ecd9a8db","year":2023},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.184008Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:07cb015559c8297cbb6a7d69c3d1cdd5c8dd940984d2004fb40f2187db5c1145","observation_id":"8fdf5c73-441f-4d90-b413-4097c80db5b8","resolution":{"observed_at":"2026-08-15T20:57:38.398513Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.382421Z","title":"Steering large language models using ape","venue":null,"work_id":"c03540a0-1943-48c5-84d4-0eb897a8fba9","year":2022},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.187524Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:a45350f30fff8a3ee943739b927039223ae9ca83e335d164c1b20e10afef25a4","observation_id":"e797f235-811b-4e94-8d43-9ca317355cc2","resolution":{"observed_at":"2026-08-15T20:57:38.387232Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01860","last_updated":"2024-06-04T00:09:43Z","snapshot_observed_at":"2026-08-18T18:30:25.544405Z","submitted_at":"2024-06-04T00:09:43Z","title":"Eliciting the Priors of Large Language Models using Iterated In-Context Learning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01860","snapshot_observed_at":"2026-08-15T20:57:38.191765Z","title":"Eliciting the priors of large language models using iterated in-context learning.arXiv preprint arXiv:2406.01860, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.191765Z"},"links":{"cited_paper":"/paper/2406.01860","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:215d19a5374648831510abde79d9b7a444670c2561d1b44948a17d745725865e","observation_id":"3a402623-89e1-4557-8471-103c12706c1d","resolution":{"observed_at":"2026-08-15T20:57:38.191765Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.16646","last_updated":"2025-05-06T01:43:38Z","snapshot_observed_at":"2026-08-16T14:23:25.695905Z","submitted_at":"2024-01-30T00:40:49Z","title":"Incoherent Probability Judgments in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.16646","snapshot_observed_at":"2026-08-15T20:57:38.195091Z","title":"Incoherent probability judgments in large language models.arXiv preprint arXiv:2401.16646, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.195091Z"},"links":{"cited_paper":"/paper/2401.16646","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:bc3329b196c8cefb61fa57ba44b6f7e160e96b56eb9942f1f3cd19ae2b0ca478","observation_id":"f46f4d46-b1b5-4524-b971-d68d47f41c96","resolution":{"observed_at":"2026-08-15T20:57:38.195091Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.19313","last_updated":"2025-05-06T01:26:28Z","snapshot_observed_at":"2026-08-19T11:00:50.852330Z","submitted_at":"2024-05-29T17:37:14Z","title":"Language Models Trained to do Arithmetic Predict Human Risky and Intertemporal Choice","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.19313","snapshot_observed_at":"2026-08-15T20:57:38.198290Z","title":"Language models trained to do arithmetic predict human risky and intertemporal choice.arXiv preprint arXiv:2405.19313, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.198290Z"},"links":{"cited_paper":"/paper/2405.19313","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:49cd062ddbd3c0476268f10c434173e5f267cf9774b13f8106886fc27c257fdd","observation_id":"e965b5b9-f89c-4fc3-bd94-096c70433851","resolution":{"observed_at":"2026-08-15T20:57:38.198290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-15T20:57:38.369093Z","title":"Recovering mental representations from large language models with Markov Chain Monte Carlo","venue":null,"work_id":"d22467a9-f639-4235-b126-fd3f08b39bb2","year":2024},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.201777Z"},"links":{"citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:a0b0139cbce7adc708b9478e55c15c055e0d5d276b7b89bdce303c8c0659ec3f","observation_id":"2504a992-3743-4329-aebb-85137627ac10","resolution":{"observed_at":"2026-08-15T20:57:38.374592Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1909.08593","last_updated":"2020-01-08T23:02:36Z","snapshot_observed_at":"2026-08-16T00:15:57.597094Z","submitted_at":"2019-09-18T17:33:39Z","title":"Fine-Tuning Language Models from Human Preferences","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1909.08593","snapshot_observed_at":"2026-08-15T20:57:38.204999Z","title":"fanning out","venue":null,"work_id":null,"year":1909},"citing_paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-15T20:57:38.204999Z"},"links":{"cited_paper":"/paper/1909.08593","citing_paper":"/paper/2505.11615"},"observation_digest":"sha256:65bf903ae990627d64fc35bbe02b9ec67c6d704da963e6e773189a4aed863273","observation_id":"c440d9be-c103-4724-84d7-73f5569f3eb0","resolution":{"observed_at":"2026-08-15T20:57:38.204999Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.11615","last_updated":"2025-05-16T18:23:10Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-20T17:08:00.879884Z","submitted_at":"2025-05-16T18:23:10Z","title":"Steering Risk Preferences in Large Language Models by Aligning Behavioral and Neural Representations"},"reference_resolution":{"displayed":45,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":19,"verified_exact":0,"verified_fuzzy":26},"total_outbound_references":45},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 21 August 2026, this Paper Citation Record lists 45 of 45 outbound references and 4 inbound Pith citation observations for arXiv:2505.11615."}