{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2025:JYRBSRP2A354X3KJTPBJG57ZRL","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"a5bc879f4a2adb98010f830bf8138fbf3e7d0b0a7108c31593e9ee98ab25e18d","cross_cats_sorted":["cs.LG","cs.NA","stat.ML"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"math.NA","submitted_at":"2025-03-13T10:53:17Z","title_canon_sha256":"dc8e6e0c6e7b40d276f42363b543da657b2a2267411258c7d1823301a2ac664e"},"schema_version":"1.0","source":{"id":"2503.10251","kind":"arxiv","version":2}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2503.10251","created_at":"2026-06-23T03:13:44Z"},{"alias_kind":"arxiv_version","alias_value":"2503.10251v2","created_at":"2026-06-23T03:13:44Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2503.10251","created_at":"2026-06-23T03:13:44Z"},{"alias_kind":"pith_short_12","alias_value":"JYRBSRP2A354","created_at":"2026-06-23T03:13:44Z"},{"alias_kind":"pith_short_16","alias_value":"JYRBSRP2A354X3KJ","created_at":"2026-06-23T03:13:44Z"},{"alias_kind":"pith_short_8","alias_value":"JYRBSRP2","created_at":"2026-06-23T03:13:44Z"}],"graph_snapshots":[{"event_id":"sha256:ecb40462de7bcda1d5f2f69294776fb848e5e0eb4ced0182ab1bbd2a6d3b467b","target":"graph","created_at":"2026-06-23T03:13:44Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2503.10251/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Transformers are the state-of-the-art architecture for large language models, and a key to their scalability is the strategic usage of low-precision arithmetic. We develop a mixed-precision analysis of transformer inference, deriving bounds for the condition numbers and forward error of the architecture's constituent parts. Notably, we compare the numerical stability of LayerNorm and RMSNorm in the massive-outlier regime, tighten the error bound of softmax in the presence of attention sinks, and quantify the impact of its shifted evaluation on the sensitivity to perturbations. Furthermore, we ","authors_text":"Longbin Zeng, Philipp Petersen, Stanislav Budzinskiy, Wenyi Fang","cross_cats":["cs.LG","cs.NA","stat.ML"],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"math.NA","submitted_at":"2025-03-13T10:53:17Z","title":"Numerical stability analysis of large language models"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2503.10251","kind":"arxiv","version":2},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:a1fa776c42386046da2f08fae41550820a864807fd764915663f54b6832341fd","target":"record","created_at":"2026-06-23T03:13:44Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"a5bc879f4a2adb98010f830bf8138fbf3e7d0b0a7108c31593e9ee98ab25e18d","cross_cats_sorted":["cs.LG","cs.NA","stat.ML"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"math.NA","submitted_at":"2025-03-13T10:53:17Z","title_canon_sha256":"dc8e6e0c6e7b40d276f42363b543da657b2a2267411258c7d1823301a2ac664e"},"schema_version":"1.0","source":{"id":"2503.10251","kind":"arxiv","version":2}},"canonical_sha256":"4e221945fa06fbcbed499bc29377f98adaa430a320fdde3d978fda23c0ec3fdc","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"4e221945fa06fbcbed499bc29377f98adaa430a320fdde3d978fda23c0ec3fdc","first_computed_at":"2026-06-23T03:13:44.720592Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-06-23T03:13:44.720592Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"5tqxZgYpMAaJJ0GthcK6KYIBoqEty1OwXL6Zz/eDx+/gmAH/a40UK0c/7jMU8Rd69M34FgJPmjAhGzuCoDKxCg==","signature_status":"signed_v1","signed_at":"2026-06-23T03:13:44.721088Z","signed_message":"canonical_sha256_bytes"},"source_id":"2503.10251","source_kind":"arxiv","source_version":2}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:a1fa776c42386046da2f08fae41550820a864807fd764915663f54b6832341fd","sha256:ecb40462de7bcda1d5f2f69294776fb848e5e0eb4ced0182ab1bbd2a6d3b467b"],"state_sha256":"09ce41ba7ce65cd11e083d5dd05d22c5de8174d9794bd60ce887056080c01efe"}