File size: 1,253 Bytes
6181f4e | 1 2 3 4 5 6 7 8 9 10 11 12 13 14 15 16 17 18 19 20 21 22 23 24 25 26 27 28 29 30 31 32 33 34 35 36 37 38 39 40 41 42 43 44 45 46 47 48 49 50 51 52 53 54 55 56 57 58 59 60 61 62 63 | {
"model": "Qwen/Qwen2.5-1.5B-Instruct",
"d_model": 1536,
"num_layers": 28,
"num_heads": 12,
"inject_layer": 9,
"best_lambda": 1.0,
"n_incontext": 4,
"n_target_prompts": 15,
"top_heads": [
{
"layer": 24,
"head": 3,
"aie": 0.004586321767419577
},
{
"layer": 20,
"head": 9,
"aie": 0.0027352971956133842
},
{
"layer": 0,
"head": 6,
"aie": 0.0017912477487698197
},
{
"layer": 17,
"head": 7,
"aie": 0.0017713563283905387
},
{
"layer": 20,
"head": 8,
"aie": 0.0016093882732093334
},
{
"layer": 21,
"head": 3,
"aie": 0.0014112978242337704
},
{
"layer": 18,
"head": 10,
"aie": 0.001295118359848857
},
{
"layer": 17,
"head": 3,
"aie": 0.0012491021770983934
},
{
"layer": 15,
"head": 1,
"aie": 0.00092243158724159
},
{
"layer": 19,
"head": 11,
"aie": 0.0007095862529240549
}
],
"sampling_strategy": "Indices 0-29 used as validation (lambda tuning). Indices 30-44 as target queries. Indices 45+ as in-context pool. N=4 in-context examples per prompt, sampled without replacement."
} |