Marat Dukhan | f1a6ed3 | 2021-09-26 13:40:19 -0700 | [diff] [blame] | 1 | #!/bin/sh |
| 2 | # Copyright 2021 Google LLC |
| 3 | # |
| 4 | # This source code is licensed under the BSD-style license found in the |
| 5 | # LICENSE file in the root directory of this source tree. |
| 6 | |
Marat Dukhan | 8ff372c | 2021-09-28 14:43:17 -0700 | [diff] [blame] | 7 | ################################## ARM NEON ################################### |
Marat Dukhan | 322ed6f | 2021-10-16 17:44:16 -0700 | [diff] [blame] | 8 | tools/xngen src/f16-f32-vcvt/neon-int16.c.in -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-neon-int16-x8.c & |
| 9 | tools/xngen src/f16-f32-vcvt/neon-int16.c.in -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-neon-int16-x16.c & |
| 10 | tools/xngen src/f16-f32-vcvt/neon-int16.c.in -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-neon-int16-x24.c & |
| 11 | tools/xngen src/f16-f32-vcvt/neon-int16.c.in -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-neon-int16-x32.c & |
| 12 | |
| 13 | tools/xngen src/f16-f32-vcvt/neon-int32.c.in -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-neon-int32-x8.c & |
| 14 | tools/xngen src/f16-f32-vcvt/neon-int32.c.in -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-neon-int32-x16.c & |
| 15 | tools/xngen src/f16-f32-vcvt/neon-int32.c.in -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-neon-int32-x24.c & |
| 16 | tools/xngen src/f16-f32-vcvt/neon-int32.c.in -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-neon-int32-x32.c & |
| 17 | |
Marat Dukhan | 8ff372c | 2021-09-28 14:43:17 -0700 | [diff] [blame] | 18 | tools/xngen src/f16-f32-vcvt/neonfp16.c.in -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-neonfp16-x8.c & |
| 19 | tools/xngen src/f16-f32-vcvt/neonfp16.c.in -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-neonfp16-x16.c & |
| 20 | |
Marat Dukhan | 1227adb | 2021-10-16 17:33:51 -0700 | [diff] [blame] | 21 | ################################# x86 128-bit ################################# |
| 22 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=2 -D AVX=0 -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-sse2-int16-x8.c & |
| 23 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=2 -D AVX=0 -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-sse2-int16-x16.c & |
| 24 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=2 -D AVX=0 -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-sse2-int16-x24.c & |
| 25 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=2 -D AVX=0 -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-sse2-int16-x32.c & |
| 26 | |
| 27 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=2 -D AVX=0 -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-sse2-int32-x8.c & |
| 28 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=2 -D AVX=0 -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-sse2-int32-x16.c & |
| 29 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=2 -D AVX=0 -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-sse2-int32-x24.c & |
| 30 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=2 -D AVX=0 -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-sse2-int32-x32.c & |
| 31 | |
| 32 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=4 -D AVX=0 -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-sse41-int16-x8.c & |
| 33 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=4 -D AVX=0 -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-sse41-int16-x16.c & |
| 34 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=4 -D AVX=0 -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-sse41-int16-x24.c & |
| 35 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=4 -D AVX=0 -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-sse41-int16-x32.c & |
| 36 | |
| 37 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=4 -D AVX=0 -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-sse41-int32-x8.c & |
| 38 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=4 -D AVX=0 -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-sse41-int32-x16.c & |
| 39 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=4 -D AVX=0 -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-sse41-int32-x24.c & |
| 40 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=4 -D AVX=0 -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-sse41-int32-x32.c & |
| 41 | |
| 42 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=4 -D AVX=1 -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-avx-int16-x8.c & |
| 43 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=4 -D AVX=1 -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-avx-int16-x16.c & |
| 44 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=4 -D AVX=1 -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-avx-int16-x24.c & |
| 45 | tools/xngen src/f16-f32-vcvt/sse-int16.c.in -D SSE=4 -D AVX=1 -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-avx-int16-x32.c & |
| 46 | |
| 47 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=4 -D AVX=1 -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-avx-int32-x8.c & |
| 48 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=4 -D AVX=1 -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-avx-int32-x16.c & |
| 49 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=4 -D AVX=1 -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-avx-int32-x24.c & |
| 50 | tools/xngen src/f16-f32-vcvt/sse-int32.c.in -D SSE=4 -D AVX=1 -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-avx-int32-x32.c & |
| 51 | |
Marat Dukhan | f1a6ed3 | 2021-09-26 13:40:19 -0700 | [diff] [blame] | 52 | ################################# x86 256-bit ################################# |
| 53 | tools/xngen src/f16-f32-vcvt/f16c.c.in -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-f16c-x8.c & |
| 54 | tools/xngen src/f16-f32-vcvt/f16c.c.in -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-f16c-x16.c & |
| 55 | |
Marat Dukhan | 79c76ab | 2021-09-26 20:26:39 -0700 | [diff] [blame] | 56 | ################################# x86 512-bit ################################# |
| 57 | tools/xngen src/f16-f32-vcvt/avx512skx.c.in -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-avx512skx-x16.c & |
| 58 | tools/xngen src/f16-f32-vcvt/avx512skx.c.in -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-avx512skx-x32.c & |
| 59 | |
Marat Dukhan | f6507f8 | 2021-10-16 18:13:04 -0700 | [diff] [blame] | 60 | ################################## WAsm SIMD ################################## |
| 61 | tools/xngen src/f16-f32-vcvt/wasmsimd-int16.c.in -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-wasmsimd-int16-x8.c & |
| 62 | tools/xngen src/f16-f32-vcvt/wasmsimd-int16.c.in -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-wasmsimd-int16-x16.c & |
| 63 | tools/xngen src/f16-f32-vcvt/wasmsimd-int16.c.in -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-wasmsimd-int16-x24.c & |
| 64 | tools/xngen src/f16-f32-vcvt/wasmsimd-int16.c.in -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-wasmsimd-int16-x32.c & |
| 65 | |
| 66 | tools/xngen src/f16-f32-vcvt/wasmsimd-int32.c.in -D BATCH_TILE=8 -o src/f16-f32-vcvt/gen/vcvt-wasmsimd-int32-x8.c & |
| 67 | tools/xngen src/f16-f32-vcvt/wasmsimd-int32.c.in -D BATCH_TILE=16 -o src/f16-f32-vcvt/gen/vcvt-wasmsimd-int32-x16.c & |
| 68 | tools/xngen src/f16-f32-vcvt/wasmsimd-int32.c.in -D BATCH_TILE=24 -o src/f16-f32-vcvt/gen/vcvt-wasmsimd-int32-x24.c & |
| 69 | tools/xngen src/f16-f32-vcvt/wasmsimd-int32.c.in -D BATCH_TILE=32 -o src/f16-f32-vcvt/gen/vcvt-wasmsimd-int32-x32.c & |
| 70 | |
Marat Dukhan | e2c0001 | 2021-10-17 22:02:35 -0700 | [diff] [blame] | 71 | #################################### Scalar ################################### |
Marat Dukhan | 134f984 | 2021-12-29 19:57:31 -0800 | [diff] [blame] | 72 | tools/xngen src/f16-f32-vcvt/scalar.c.in -D BATCH_TILE=1 -o src/f16-f32-vcvt/gen/vcvt-scalar-x1.c & |
| 73 | tools/xngen src/f16-f32-vcvt/scalar.c.in -D BATCH_TILE=2 -o src/f16-f32-vcvt/gen/vcvt-scalar-x2.c & |
| 74 | tools/xngen src/f16-f32-vcvt/scalar.c.in -D BATCH_TILE=3 -o src/f16-f32-vcvt/gen/vcvt-scalar-x3.c & |
| 75 | tools/xngen src/f16-f32-vcvt/scalar.c.in -D BATCH_TILE=4 -o src/f16-f32-vcvt/gen/vcvt-scalar-x4.c & |
Marat Dukhan | e2c0001 | 2021-10-17 22:02:35 -0700 | [diff] [blame] | 76 | |
Marat Dukhan | f1a6ed3 | 2021-09-26 13:40:19 -0700 | [diff] [blame] | 77 | ################################## Unit tests ################################# |
| 78 | tools/generate-vcvt-test.py --spec test/f16-f32-vcvt.yaml --output test/f16-f32-vcvt.cc & |
| 79 | |
| 80 | wait |