162306a36Sopenharmony_ci/* SPDX-License-Identifier: GPL-2.0-or-later */ 262306a36Sopenharmony_ci/* 362306a36Sopenharmony_ci * Shared glue code for 128bit block ciphers, AVX2 assembler macros 462306a36Sopenharmony_ci * 562306a36Sopenharmony_ci * Copyright © 2012-2013 Jussi Kivilinna <jussi.kivilinna@mbnet.fi> 662306a36Sopenharmony_ci */ 762306a36Sopenharmony_ci 862306a36Sopenharmony_ci#define load_16way(src, x0, x1, x2, x3, x4, x5, x6, x7) \ 962306a36Sopenharmony_ci vmovdqu (0*32)(src), x0; \ 1062306a36Sopenharmony_ci vmovdqu (1*32)(src), x1; \ 1162306a36Sopenharmony_ci vmovdqu (2*32)(src), x2; \ 1262306a36Sopenharmony_ci vmovdqu (3*32)(src), x3; \ 1362306a36Sopenharmony_ci vmovdqu (4*32)(src), x4; \ 1462306a36Sopenharmony_ci vmovdqu (5*32)(src), x5; \ 1562306a36Sopenharmony_ci vmovdqu (6*32)(src), x6; \ 1662306a36Sopenharmony_ci vmovdqu (7*32)(src), x7; 1762306a36Sopenharmony_ci 1862306a36Sopenharmony_ci#define store_16way(dst, x0, x1, x2, x3, x4, x5, x6, x7) \ 1962306a36Sopenharmony_ci vmovdqu x0, (0*32)(dst); \ 2062306a36Sopenharmony_ci vmovdqu x1, (1*32)(dst); \ 2162306a36Sopenharmony_ci vmovdqu x2, (2*32)(dst); \ 2262306a36Sopenharmony_ci vmovdqu x3, (3*32)(dst); \ 2362306a36Sopenharmony_ci vmovdqu x4, (4*32)(dst); \ 2462306a36Sopenharmony_ci vmovdqu x5, (5*32)(dst); \ 2562306a36Sopenharmony_ci vmovdqu x6, (6*32)(dst); \ 2662306a36Sopenharmony_ci vmovdqu x7, (7*32)(dst); 2762306a36Sopenharmony_ci 2862306a36Sopenharmony_ci#define store_cbc_16way(src, dst, x0, x1, x2, x3, x4, x5, x6, x7, t0) \ 2962306a36Sopenharmony_ci vpxor t0, t0, t0; \ 3062306a36Sopenharmony_ci vinserti128 $1, (src), t0, t0; \ 3162306a36Sopenharmony_ci vpxor t0, x0, x0; \ 3262306a36Sopenharmony_ci vpxor (0*32+16)(src), x1, x1; \ 3362306a36Sopenharmony_ci vpxor (1*32+16)(src), x2, x2; \ 3462306a36Sopenharmony_ci vpxor (2*32+16)(src), x3, x3; \ 3562306a36Sopenharmony_ci vpxor (3*32+16)(src), x4, x4; \ 3662306a36Sopenharmony_ci vpxor (4*32+16)(src), x5, x5; \ 3762306a36Sopenharmony_ci vpxor (5*32+16)(src), x6, x6; \ 3862306a36Sopenharmony_ci vpxor (6*32+16)(src), x7, x7; \ 3962306a36Sopenharmony_ci store_16way(dst, x0, x1, x2, x3, x4, x5, x6, x7); 40