<?xml version="1.0" encoding="utf-8" standalone="yes"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
  <channel>
    <title>Representations on Nalar</title>
    <link>https://nalar.dev/tags/representations/</link>
    <description>Recent content in Representations on Nalar</description>
    <generator>Hugo</generator>
    <language>en-us</language>
    <lastBuildDate>Tue, 15 Sep 2026 00:00:00 +0700</lastBuildDate>
    <atom:link href="https://nalar.dev/tags/representations/index.xml" rel="self" type="application/rss+xml" />
    <item>
      <title>Steer Transformer Activations with Residual Stream Vectors</title>
      <link>https://nalar.dev/steer-transformer-activations-with-residual-stream-vectors/</link>
      <pubDate>Tue, 15 Sep 2026 00:00:00 +0700</pubDate>
      <guid>https://nalar.dev/steer-transformer-activations-with-residual-stream-vectors/</guid>
      <description>&lt;p&gt;A transformer can produce different output behavior even when its weights and input tokens stay fixed. One way to cause that change is to alter an intermediate hidden state during the forward pass. Activation steering does this deliberately by adding a vector to a selected residual-stream position or set of positions.&lt;/p&gt;&#xA;&lt;p&gt;The mechanism is simple enough to express as an intervention, but its effect is not a global model setting. The chosen direction, coefficient, layer, token positions, and decoding setup all affect the result. Treating those choices as part of the inference configuration makes the behavior easier to reason about and test.&lt;/p&gt;</description>
    </item>
  </channel>
</rss>
