Download code/kernels/sp_desc/dh_reader.cpp from changh95/superpoint-p150: direct link, hf CLI and curl.
- Browser
- Download file 1.28 kB
-
https://huggingface.co/changh95/superpoint-p150/resolve/main/code/kernels/sp_desc/dh_reader.cpp
- Command line
-
hf download hf://changh95/superpoint-p150/code/kernels/sp_desc/dh_reader.cpp
-
curl -L -o dh_reader.cpp https://huggingface.co/changh95/superpoint-p150/resolve/main/code/kernels/sp_desc/dh_reader.cpp
1.28 kB
| // SPDX-FileCopyrightText: © 2026 Tenstorrent USA, Inc. | |
| // SPDX-License-Identifier: Apache-2.0 | |
| // | |
| // Descriptor head (1x1 conv + L2 norm + untilize, models/tt/desc_head.py), PROC 0: publish the local | |
| // input shard (CB_IN) and the core's L1-resident copy of the 1x1 weights + bias (CB_W, both bound to | |
| // sharded tensors), and build the reduce scaler tile (1.0 in row 0 of each face). | |
| void kernel_main() { | |
| constexpr uint32_t cb_in = get_compile_time_arg_val(0); | |
| constexpr uint32_t cb_one = get_compile_time_arg_val(1); | |
| constexpr uint32_t NT = get_compile_time_arg_val(2); | |
| constexpr uint32_t cb_w = get_compile_time_arg_val(3); | |
| constexpr uint32_t NW = get_compile_time_arg_val(4); | |
| cb_reserve_back(cb_in, NT); | |
| cb_push_back(cb_in, NT); | |
| cb_reserve_back(cb_w, NW); | |
| cb_push_back(cb_w, NW); | |
| cb_reserve_back(cb_one, 1); | |
| volatile tt_l1_ptr uint32_t* p = reinterpret_cast<volatile tt_l1_ptr uint32_t*>(get_write_ptr(cb_one)); | |
| for (uint32_t i = 0; i < 512; ++i) { | |
| p[i] = 0; | |
| } | |
| for (uint32_t f = 0; f < 4; ++f) { | |
| for (uint32_t j = 0; j < 8; ++j) { | |
| p[f * 128 + j] = 0x3F803F80u; | |
| } | |
| } | |
| (void)p[511]; | |
| cb_push_back(cb_one, 1); | |
| } | |