glue_helper-asm-avx2.S 1.2 KB

123456789101112131415161718192021222324252627282930313233343536373839
  1. /* SPDX-License-Identifier: GPL-2.0-or-later */
  2. /*
  3. * Shared glue code for 128bit block ciphers, AVX2 assembler macros
  4. *
  5. * Copyright © 2012-2013 Jussi Kivilinna <[email protected]>
  6. */
  7. #define load_16way(src, x0, x1, x2, x3, x4, x5, x6, x7) \
  8. vmovdqu (0*32)(src), x0; \
  9. vmovdqu (1*32)(src), x1; \
  10. vmovdqu (2*32)(src), x2; \
  11. vmovdqu (3*32)(src), x3; \
  12. vmovdqu (4*32)(src), x4; \
  13. vmovdqu (5*32)(src), x5; \
  14. vmovdqu (6*32)(src), x6; \
  15. vmovdqu (7*32)(src), x7;
  16. #define store_16way(dst, x0, x1, x2, x3, x4, x5, x6, x7) \
  17. vmovdqu x0, (0*32)(dst); \
  18. vmovdqu x1, (1*32)(dst); \
  19. vmovdqu x2, (2*32)(dst); \
  20. vmovdqu x3, (3*32)(dst); \
  21. vmovdqu x4, (4*32)(dst); \
  22. vmovdqu x5, (5*32)(dst); \
  23. vmovdqu x6, (6*32)(dst); \
  24. vmovdqu x7, (7*32)(dst);
  25. #define store_cbc_16way(src, dst, x0, x1, x2, x3, x4, x5, x6, x7, t0) \
  26. vpxor t0, t0, t0; \
  27. vinserti128 $1, (src), t0, t0; \
  28. vpxor t0, x0, x0; \
  29. vpxor (0*32+16)(src), x1, x1; \
  30. vpxor (1*32+16)(src), x2, x2; \
  31. vpxor (2*32+16)(src), x3, x3; \
  32. vpxor (3*32+16)(src), x4, x4; \
  33. vpxor (4*32+16)(src), x5, x5; \
  34. vpxor (5*32+16)(src), x6, x6; \
  35. vpxor (6*32+16)(src), x7, x7; \
  36. store_16way(dst, x0, x1, x2, x3, x4, x5, x6, x7);