xref: /aosp_15_r20/external/XNNPACK/scripts/generate-f16-vsigmoid.sh (revision 4bdc94577ba0e567308109d787f7fec7b531ce36)
1#!/bin/sh
2# Copyright 2022 Google LLC
3#
4# This source code is licensed under the BSD-style license found in the
5# LICENSE file in the root directory of this source tree.
6
7################################### ARM NEON ##################################
8tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=8  -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x8.c &
9tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=16 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x16.c &
10tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=24 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x24.c &
11tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=32 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x32.c &
12tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=40 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x40.c &
13tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=48 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x48.c &
14tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=56 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x56.c &
15tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=64 -D DIV_ALGO=DIV -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-div-x64.c &
16
17tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=8  -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x8.c &
18tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=16 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x16.c &
19tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=24 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x24.c &
20tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=32 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x32.c &
21tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=40 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x40.c &
22tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=48 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x48.c &
23tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=56 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x56.c &
24tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=64 -D DIV_ALGO=NR1FMA -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1fma-x64.c &
25
26tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=8  -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x8.c &
27tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=16 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x16.c &
28tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=24 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x24.c &
29tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=32 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x32.c &
30tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=40 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x40.c &
31tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=48 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x48.c &
32tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=56 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x56.c &
33tools/xngen src/f16-vsigmoid/neonfp16arith.c.in -D BATCH_TILE=64 -D DIV_ALGO=NR1RECPS -o src/f16-vsigmoid/gen/vsigmoid-neonfp16arith-rr2-p2-nr1recps-x64.c &
34
35################################### x86 AVX2 ##################################
36tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=8  -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x8.c &
37tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=16 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x16.c &
38tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=24 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x24.c &
39tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=32 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x32.c &
40tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=40 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x40.c &
41tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=48 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x48.c &
42tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=56 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x56.c &
43tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=64 -D DIV_ALGO=div -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-div-x64.c &
44
45tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=8  -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x8.c &
46tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=16 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x16.c &
47tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=24 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x24.c &
48tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=32 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x32.c &
49tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=40 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x40.c &
50tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=48 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x48.c &
51tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=56 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x56.c &
52tools/xngen src/f16-vsigmoid/avx2.c.in -D BATCH_TILE=64 -D DIV_ALGO=rcp -o src/f16-vsigmoid/gen/vsigmoid-avx2-rr1-p2-rcp-x64.c &
53
54################################## Unit tests #################################
55tools/generate-vunary-test.py --spec test/f16-vsigmoid.yaml --output test/f16-vsigmoid.cc &
56
57wait
58