162306a36Sopenharmony_ci/* SPDX-License-Identifier: GPL-2.0-or-later */
262306a36Sopenharmony_ci/*
362306a36Sopenharmony_ci * Shared glue code for 128bit block ciphers, AVX2 assembler macros
462306a36Sopenharmony_ci *
562306a36Sopenharmony_ci * Copyright © 2012-2013 Jussi Kivilinna <jussi.kivilinna@mbnet.fi>
662306a36Sopenharmony_ci */
762306a36Sopenharmony_ci
862306a36Sopenharmony_ci#define load_16way(src, x0, x1, x2, x3, x4, x5, x6, x7) \
962306a36Sopenharmony_ci	vmovdqu (0*32)(src), x0; \
1062306a36Sopenharmony_ci	vmovdqu (1*32)(src), x1; \
1162306a36Sopenharmony_ci	vmovdqu (2*32)(src), x2; \
1262306a36Sopenharmony_ci	vmovdqu (3*32)(src), x3; \
1362306a36Sopenharmony_ci	vmovdqu (4*32)(src), x4; \
1462306a36Sopenharmony_ci	vmovdqu (5*32)(src), x5; \
1562306a36Sopenharmony_ci	vmovdqu (6*32)(src), x6; \
1662306a36Sopenharmony_ci	vmovdqu (7*32)(src), x7;
1762306a36Sopenharmony_ci
1862306a36Sopenharmony_ci#define store_16way(dst, x0, x1, x2, x3, x4, x5, x6, x7) \
1962306a36Sopenharmony_ci	vmovdqu x0, (0*32)(dst); \
2062306a36Sopenharmony_ci	vmovdqu x1, (1*32)(dst); \
2162306a36Sopenharmony_ci	vmovdqu x2, (2*32)(dst); \
2262306a36Sopenharmony_ci	vmovdqu x3, (3*32)(dst); \
2362306a36Sopenharmony_ci	vmovdqu x4, (4*32)(dst); \
2462306a36Sopenharmony_ci	vmovdqu x5, (5*32)(dst); \
2562306a36Sopenharmony_ci	vmovdqu x6, (6*32)(dst); \
2662306a36Sopenharmony_ci	vmovdqu x7, (7*32)(dst);
2762306a36Sopenharmony_ci
2862306a36Sopenharmony_ci#define store_cbc_16way(src, dst, x0, x1, x2, x3, x4, x5, x6, x7, t0) \
2962306a36Sopenharmony_ci	vpxor t0, t0, t0; \
3062306a36Sopenharmony_ci	vinserti128 $1, (src), t0, t0; \
3162306a36Sopenharmony_ci	vpxor t0, x0, x0; \
3262306a36Sopenharmony_ci	vpxor (0*32+16)(src), x1, x1; \
3362306a36Sopenharmony_ci	vpxor (1*32+16)(src), x2, x2; \
3462306a36Sopenharmony_ci	vpxor (2*32+16)(src), x3, x3; \
3562306a36Sopenharmony_ci	vpxor (3*32+16)(src), x4, x4; \
3662306a36Sopenharmony_ci	vpxor (4*32+16)(src), x5, x5; \
3762306a36Sopenharmony_ci	vpxor (5*32+16)(src), x6, x6; \
3862306a36Sopenharmony_ci	vpxor (6*32+16)(src), x7, x7; \
3962306a36Sopenharmony_ci	store_16way(dst, x0, x1, x2, x3, x4, x5, x6, x7);
40