xref: /libCEED/rust/libceed-sys/c-src/backends/opt/ceed-opt-operator.c (revision e15f9bd09af0280c89b79924fa9af7dd2e3e30be)
189c6efa4Sjeremylt // Copyright (c) 2017-2018, Lawrence Livermore National Security, LLC.
289c6efa4Sjeremylt // Produced at the Lawrence Livermore National Laboratory. LLNL-CODE-734707.
389c6efa4Sjeremylt // All Rights reserved. See files LICENSE and NOTICE for details.
489c6efa4Sjeremylt //
589c6efa4Sjeremylt // This file is part of CEED, a collection of benchmarks, miniapps, software
689c6efa4Sjeremylt // libraries and APIs for efficient high-order finite element and spectral
789c6efa4Sjeremylt // element discretizations for exascale applications. For more information and
889c6efa4Sjeremylt // source code availability see http://github.com/ceed.
989c6efa4Sjeremylt //
1089c6efa4Sjeremylt // The CEED research is supported by the Exascale Computing Project 17-SC-20-SC,
1189c6efa4Sjeremylt // a collaborative effort of two U.S. Department of Energy organizations (Office
1289c6efa4Sjeremylt // of Science and the National Nuclear Security Administration) responsible for
1389c6efa4Sjeremylt // the planning and preparation of a capable exascale ecosystem, including
1489c6efa4Sjeremylt // software, applications, hardware, advanced system engineering and early
1589c6efa4Sjeremylt // testbed platforms, in support of the nation's exascale computing imperative.
1689c6efa4Sjeremylt 
173d576824SJeremy L Thompson #include <ceed.h>
183d576824SJeremy L Thompson #include <ceed-backend.h>
193d576824SJeremy L Thompson #include <stdbool.h>
203d576824SJeremy L Thompson #include <stdint.h>
2189c6efa4Sjeremylt #include <string.h>
2289c6efa4Sjeremylt #include "ceed-opt.h"
2389c6efa4Sjeremylt 
24f10650afSjeremylt //------------------------------------------------------------------------------
25f10650afSjeremylt // Setup Input/Output Fields
26f10650afSjeremylt //------------------------------------------------------------------------------
2789c6efa4Sjeremylt static int CeedOperatorSetupFields_Opt(CeedQFunction qf, CeedOperator op,
2889c6efa4Sjeremylt                                        bool inOrOut, const CeedInt blksize,
2989c6efa4Sjeremylt                                        CeedElemRestriction *blkrestr,
3089c6efa4Sjeremylt                                        CeedVector *fullevecs, CeedVector *evecs,
3189c6efa4Sjeremylt                                        CeedVector *qvecs, CeedInt starte,
3289c6efa4Sjeremylt                                        CeedInt numfields, CeedInt Q) {
334d537eeaSYohann   CeedInt dim, ierr, ncomp, size, P;
3489c6efa4Sjeremylt   Ceed ceed;
35*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetCeed(op, &ceed); CeedChkBackend(ierr);
3689c6efa4Sjeremylt   CeedBasis basis;
3789c6efa4Sjeremylt   CeedElemRestriction r;
3889c6efa4Sjeremylt   CeedOperatorField *opfields;
3989c6efa4Sjeremylt   CeedQFunctionField *qffields;
4089c6efa4Sjeremylt   if (inOrOut) {
4189c6efa4Sjeremylt     ierr = CeedOperatorGetFields(op, NULL, &opfields);
42*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
4389c6efa4Sjeremylt     ierr = CeedQFunctionGetFields(qf, NULL, &qffields);
44*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
4589c6efa4Sjeremylt   } else {
4689c6efa4Sjeremylt     ierr = CeedOperatorGetFields(op, &opfields, NULL);
47*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
4889c6efa4Sjeremylt     ierr = CeedQFunctionGetFields(qf, &qffields, NULL);
49*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
5089c6efa4Sjeremylt   }
5189c6efa4Sjeremylt 
5289c6efa4Sjeremylt   // Loop over fields
5389c6efa4Sjeremylt   for (CeedInt i=0; i<numfields; i++) {
5489c6efa4Sjeremylt     CeedEvalMode emode;
55*e15f9bd0SJeremy L Thompson     ierr = CeedQFunctionFieldGetEvalMode(qffields[i], &emode); CeedChkBackend(ierr);
5689c6efa4Sjeremylt 
5789c6efa4Sjeremylt     if (emode != CEED_EVAL_WEIGHT) {
5889c6efa4Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opfields[i], &r);
59*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
6089c6efa4Sjeremylt       Ceed ceed;
61*e15f9bd0SJeremy L Thompson       ierr = CeedElemRestrictionGetCeed(r, &ceed); CeedChkBackend(ierr);
62d979a051Sjeremylt       CeedInt nelem, elemsize, lsize, compstride;
63*e15f9bd0SJeremy L Thompson       ierr = CeedElemRestrictionGetNumElements(r, &nelem); CeedChkBackend(ierr);
64*e15f9bd0SJeremy L Thompson       ierr = CeedElemRestrictionGetElementSize(r, &elemsize); CeedChkBackend(ierr);
65*e15f9bd0SJeremy L Thompson       ierr = CeedElemRestrictionGetLVectorSize(r, &lsize); CeedChkBackend(ierr);
66*e15f9bd0SJeremy L Thompson       ierr = CeedElemRestrictionGetNumComponents(r, &ncomp); CeedChkBackend(ierr);
67bd33150aSjeremylt 
683ac43b2cSJeremy L Thompson       bool strided;
69*e15f9bd0SJeremy L Thompson       ierr = CeedElemRestrictionIsStrided(r, &strided); CeedChkBackend(ierr);
703ac43b2cSJeremy L Thompson       if (strided) {
713ac43b2cSJeremy L Thompson         CeedInt strides[3];
72*e15f9bd0SJeremy L Thompson         ierr = CeedElemRestrictionGetStrides(r, &strides); CeedChkBackend(ierr);
733ac43b2cSJeremy L Thompson         ierr = CeedElemRestrictionCreateBlockedStrided(ceed, nelem, elemsize,
743ac43b2cSJeremy L Thompson                blksize, ncomp, lsize, strides, &blkrestr[i+starte]);
75*e15f9bd0SJeremy L Thompson         CeedChkBackend(ierr);
763ac43b2cSJeremy L Thompson       } else {
77bd33150aSjeremylt         const CeedInt *offsets = NULL;
78bd33150aSjeremylt         ierr = CeedElemRestrictionGetOffsets(r, CEED_MEM_HOST, &offsets);
79*e15f9bd0SJeremy L Thompson         CeedChkBackend(ierr);
80*e15f9bd0SJeremy L Thompson         ierr = CeedElemRestrictionGetCompStride(r, &compstride); CeedChkBackend(ierr);
81d979a051Sjeremylt         ierr = CeedElemRestrictionCreateBlocked(ceed, nelem, elemsize,
82d979a051Sjeremylt                                                 blksize, ncomp, compstride,
83d979a051Sjeremylt                                                 lsize, CEED_MEM_HOST,
84bd33150aSjeremylt                                                 CEED_COPY_VALUES, offsets,
857509a596Sjeremylt                                                 &blkrestr[i+starte]);
86*e15f9bd0SJeremy L Thompson         CeedChkBackend(ierr);
87*e15f9bd0SJeremy L Thompson         ierr = CeedElemRestrictionRestoreOffsets(r, &offsets); CeedChkBackend(ierr);
883ac43b2cSJeremy L Thompson       }
8989c6efa4Sjeremylt       ierr = CeedElemRestrictionCreateVector(blkrestr[i+starte], NULL,
9089c6efa4Sjeremylt                                              &fullevecs[i+starte]);
91*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
9289c6efa4Sjeremylt     }
9389c6efa4Sjeremylt 
9489c6efa4Sjeremylt     switch(emode) {
9589c6efa4Sjeremylt     case CEED_EVAL_NONE:
96*e15f9bd0SJeremy L Thompson       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChkBackend(ierr);
97*e15f9bd0SJeremy L Thompson       ierr = CeedVectorCreate(ceed, Q*size*blksize, &evecs[i]); CeedChkBackend(ierr);
98*e15f9bd0SJeremy L Thompson       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChkBackend(ierr);
9989c6efa4Sjeremylt       break;
10089c6efa4Sjeremylt     case CEED_EVAL_INTERP:
101*e15f9bd0SJeremy L Thompson       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChkBackend(ierr);
10289c6efa4Sjeremylt       ierr = CeedElemRestrictionGetElementSize(r, &P);
103*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
104*e15f9bd0SJeremy L Thompson       ierr = CeedVectorCreate(ceed, P*size*blksize, &evecs[i]); CeedChkBackend(ierr);
105*e15f9bd0SJeremy L Thompson       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChkBackend(ierr);
10689c6efa4Sjeremylt       break;
10789c6efa4Sjeremylt     case CEED_EVAL_GRAD:
108*e15f9bd0SJeremy L Thompson       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChkBackend(ierr);
109*e15f9bd0SJeremy L Thompson       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChkBackend(ierr);
110*e15f9bd0SJeremy L Thompson       ierr = CeedBasisGetDimension(basis, &dim); CeedChkBackend(ierr);
11189c6efa4Sjeremylt       ierr = CeedElemRestrictionGetElementSize(r, &P);
112*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
113*e15f9bd0SJeremy L Thompson       ierr = CeedVectorCreate(ceed, P*size/dim*blksize, &evecs[i]);
114*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
115*e15f9bd0SJeremy L Thompson       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChkBackend(ierr);
11689c6efa4Sjeremylt       break;
11789c6efa4Sjeremylt     case CEED_EVAL_WEIGHT: // Only on input fields
118*e15f9bd0SJeremy L Thompson       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChkBackend(ierr);
119*e15f9bd0SJeremy L Thompson       ierr = CeedVectorCreate(ceed, Q*blksize, &qvecs[i]); CeedChkBackend(ierr);
12089c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
121a7b7f929Sjeremylt                             CEED_EVAL_WEIGHT, CEED_VECTOR_NONE, qvecs[i]);
122*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
12389c6efa4Sjeremylt 
12489c6efa4Sjeremylt       break;
12589c6efa4Sjeremylt     case CEED_EVAL_DIV:
1265f67fadeSJeremy L Thompson       break; // Not implemented
12789c6efa4Sjeremylt     case CEED_EVAL_CURL:
1285f67fadeSJeremy L Thompson       break; // Not implemented
12989c6efa4Sjeremylt     }
13089c6efa4Sjeremylt   }
131*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
13289c6efa4Sjeremylt }
13389c6efa4Sjeremylt 
134f10650afSjeremylt //------------------------------------------------------------------------------
135f10650afSjeremylt // Setup Operator
136f10650afSjeremylt //------------------------------------------------------------------------------
13789c6efa4Sjeremylt static int CeedOperatorSetup_Opt(CeedOperator op) {
13889c6efa4Sjeremylt   int ierr;
13989c6efa4Sjeremylt   bool setupdone;
140*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorIsSetupDone(op, &setupdone); CeedChkBackend(ierr);
141*e15f9bd0SJeremy L Thompson   if (setupdone) return CEED_ERROR_SUCCESS;
14289c6efa4Sjeremylt   Ceed ceed;
143*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetCeed(op, &ceed); CeedChkBackend(ierr);
14489c6efa4Sjeremylt   Ceed_Opt *ceedimpl;
145*e15f9bd0SJeremy L Thompson   ierr = CeedGetData(ceed, &ceedimpl); CeedChkBackend(ierr);
14689c6efa4Sjeremylt   const CeedInt blksize = ceedimpl->blksize;
14789c6efa4Sjeremylt   CeedOperator_Opt *impl;
148*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetData(op, &impl); CeedChkBackend(ierr);
14989c6efa4Sjeremylt   CeedQFunction qf;
150*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetQFunction(op, &qf); CeedChkBackend(ierr);
15189c6efa4Sjeremylt   CeedInt Q, numinputfields, numoutputfields;
152*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChkBackend(ierr);
153*e15f9bd0SJeremy L Thompson   ierr = CeedQFunctionIsIdentity(qf, &impl->identityqf); CeedChkBackend(ierr);
15489c6efa4Sjeremylt   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
155*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
15689c6efa4Sjeremylt   CeedOperatorField *opinputfields, *opoutputfields;
15789c6efa4Sjeremylt   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
158*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
15989c6efa4Sjeremylt   CeedQFunctionField *qfinputfields, *qfoutputfields;
16089c6efa4Sjeremylt   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
161*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
16289c6efa4Sjeremylt 
16389c6efa4Sjeremylt   // Allocate
16489c6efa4Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->blkrestr);
165*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
16689c6efa4Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->evecs);
167*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
16889c6efa4Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->edata);
169*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
17089c6efa4Sjeremylt 
171*e15f9bd0SJeremy L Thompson   ierr = CeedCalloc(16, &impl->inputstate); CeedChkBackend(ierr);
172*e15f9bd0SJeremy L Thompson   ierr = CeedCalloc(16, &impl->evecsin); CeedChkBackend(ierr);
173*e15f9bd0SJeremy L Thompson   ierr = CeedCalloc(16, &impl->evecsout); CeedChkBackend(ierr);
174*e15f9bd0SJeremy L Thompson   ierr = CeedCalloc(16, &impl->qvecsin); CeedChkBackend(ierr);
175*e15f9bd0SJeremy L Thompson   ierr = CeedCalloc(16, &impl->qvecsout); CeedChkBackend(ierr);
17689c6efa4Sjeremylt 
17789c6efa4Sjeremylt   impl->numein = numinputfields; impl->numeout = numoutputfields;
17889c6efa4Sjeremylt 
17989c6efa4Sjeremylt   // Set up infield and outfield pointer arrays
18089c6efa4Sjeremylt   // Infields
18189c6efa4Sjeremylt   ierr = CeedOperatorSetupFields_Opt(qf, op, 0, blksize, impl->blkrestr,
18289c6efa4Sjeremylt                                      impl->evecs, impl->evecsin,
18389c6efa4Sjeremylt                                      impl->qvecsin, 0,
18489c6efa4Sjeremylt                                      numinputfields, Q);
185*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
18689c6efa4Sjeremylt   // Outfields
18789c6efa4Sjeremylt   ierr = CeedOperatorSetupFields_Opt(qf, op, 1, blksize, impl->blkrestr,
18889c6efa4Sjeremylt                                      impl->evecs, impl->evecsout,
18989c6efa4Sjeremylt                                      impl->qvecsout, numinputfields,
19089c6efa4Sjeremylt                                      numoutputfields, Q);
191*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
19289c6efa4Sjeremylt 
19316911fdaSjeremylt   // Identity QFunctions
19416911fdaSjeremylt   if (impl->identityqf) {
19516911fdaSjeremylt     CeedEvalMode inmode, outmode;
19616911fdaSjeremylt     CeedQFunctionField *infields, *outfields;
197*e15f9bd0SJeremy L Thompson     ierr = CeedQFunctionGetFields(qf, &infields, &outfields); CeedChkBackend(ierr);
19816911fdaSjeremylt 
19916911fdaSjeremylt     for (CeedInt i=0; i<numinputfields; i++) {
20016911fdaSjeremylt       ierr = CeedQFunctionFieldGetEvalMode(infields[i], &inmode);
201*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
20216911fdaSjeremylt       ierr = CeedQFunctionFieldGetEvalMode(outfields[i], &outmode);
203*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
20416911fdaSjeremylt 
205*e15f9bd0SJeremy L Thompson       ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChkBackend(ierr);
20616911fdaSjeremylt       impl->qvecsout[i] = impl->qvecsin[i];
207*e15f9bd0SJeremy L Thompson       ierr = CeedVectorAddReference(impl->qvecsin[i]); CeedChkBackend(ierr);
20816911fdaSjeremylt     }
20916911fdaSjeremylt   }
21016911fdaSjeremylt 
211*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorSetSetupDone(op); CeedChkBackend(ierr);
21289c6efa4Sjeremylt 
213*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
21489c6efa4Sjeremylt }
21589c6efa4Sjeremylt 
216f10650afSjeremylt //------------------------------------------------------------------------------
217f10650afSjeremylt // Setup Input Fields
218f10650afSjeremylt //------------------------------------------------------------------------------
2191d102b48SJeremy L Thompson static inline int CeedOperatorSetupInputs_Opt(CeedInt numinputfields,
2201d102b48SJeremy L Thompson     CeedQFunctionField *qfinputfields, CeedOperatorField *opinputfields,
2211d102b48SJeremy L Thompson     CeedVector invec, CeedOperator_Opt *impl, CeedRequest *request) {
2221d102b48SJeremy L Thompson   CeedInt ierr;
22389c6efa4Sjeremylt   CeedEvalMode emode;
22489c6efa4Sjeremylt   CeedVector vec;
22589c6efa4Sjeremylt   uint64_t state;
22689c6efa4Sjeremylt 
22789c6efa4Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
22889c6efa4Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
229*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
23089c6efa4Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
23189c6efa4Sjeremylt     } else {
23289c6efa4Sjeremylt       // Get input vector
233*e15f9bd0SJeremy L Thompson       ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChkBackend(ierr);
23489c6efa4Sjeremylt       if (vec != CEED_VECTOR_ACTIVE) {
23589c6efa4Sjeremylt         // Restrict
236*e15f9bd0SJeremy L Thompson         ierr = CeedVectorGetState(vec, &state); CeedChkBackend(ierr);
23789c6efa4Sjeremylt         if (state != impl->inputstate[i]) {
23889c6efa4Sjeremylt           ierr = CeedElemRestrictionApply(impl->blkrestr[i], CEED_NOTRANSPOSE,
239a8d32208Sjeremylt                                           vec, impl->evecs[i], request);
240*e15f9bd0SJeremy L Thompson           CeedChkBackend(ierr);
24189c6efa4Sjeremylt           impl->inputstate[i] = state;
24289c6efa4Sjeremylt         }
24389c6efa4Sjeremylt       } else {
24489c6efa4Sjeremylt         // Set Qvec for CEED_EVAL_NONE
24589c6efa4Sjeremylt         if (emode == CEED_EVAL_NONE) {
24689c6efa4Sjeremylt           ierr = CeedVectorGetArray(impl->evecsin[i], CEED_MEM_HOST,
247*e15f9bd0SJeremy L Thompson                                     &impl->edata[i]); CeedChkBackend(ierr);
24889c6efa4Sjeremylt           ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST,
24989c6efa4Sjeremylt                                     CEED_USE_POINTER,
250*e15f9bd0SJeremy L Thompson                                     impl->edata[i]); CeedChkBackend(ierr);
25189c6efa4Sjeremylt           ierr = CeedVectorRestoreArray(impl->evecsin[i],
252*e15f9bd0SJeremy L Thompson                                         &impl->edata[i]); CeedChkBackend(ierr);
25389c6efa4Sjeremylt         }
25489c6efa4Sjeremylt       }
25589c6efa4Sjeremylt       // Get evec
25689c6efa4Sjeremylt       ierr = CeedVectorGetArrayRead(impl->evecs[i], CEED_MEM_HOST,
25789c6efa4Sjeremylt                                     (const CeedScalar **) &impl->edata[i]);
258*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
25989c6efa4Sjeremylt     }
26089c6efa4Sjeremylt   }
261*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
2621d102b48SJeremy L Thompson }
26389c6efa4Sjeremylt 
264f10650afSjeremylt //------------------------------------------------------------------------------
265f10650afSjeremylt // Input Basis Action
266f10650afSjeremylt //------------------------------------------------------------------------------
2671d102b48SJeremy L Thompson static inline int CeedOperatorInputBasis_Opt(CeedInt e, CeedInt Q,
2681d102b48SJeremy L Thompson     CeedQFunctionField *qfinputfields, CeedOperatorField *opinputfields,
2691d102b48SJeremy L Thompson     CeedInt numinputfields, CeedInt blksize, CeedVector invec, bool skipactive,
2701d102b48SJeremy L Thompson     CeedOperator_Opt *impl, CeedRequest *request) {
2711d102b48SJeremy L Thompson   CeedInt ierr;
2721d102b48SJeremy L Thompson   CeedInt dim, elemsize, size;
2731d102b48SJeremy L Thompson   CeedElemRestriction Erestrict;
2741d102b48SJeremy L Thompson   CeedEvalMode emode;
2751d102b48SJeremy L Thompson   CeedBasis basis;
2761d102b48SJeremy L Thompson   CeedVector vec;
27789c6efa4Sjeremylt 
27889c6efa4Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
279*e15f9bd0SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChkBackend(ierr);
2801d102b48SJeremy L Thompson     // Skip active input
2811d102b48SJeremy L Thompson     if (skipactive) {
2821d102b48SJeremy L Thompson       if (vec == CEED_VECTOR_ACTIVE)
2831d102b48SJeremy L Thompson         continue;
2841d102b48SJeremy L Thompson     }
2851d102b48SJeremy L Thompson 
28689c6efa4Sjeremylt     CeedInt activein = 0;
2874d537eeaSYohann     // Get elemsize, emode, size
28889c6efa4Sjeremylt     ierr = CeedOperatorFieldGetElemRestriction(opinputfields[i], &Erestrict);
289*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
29089c6efa4Sjeremylt     ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
291*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
29289c6efa4Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
293*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
294*e15f9bd0SJeremy L Thompson     ierr = CeedQFunctionFieldGetSize(qfinputfields[i], &size); CeedChkBackend(ierr);
29589c6efa4Sjeremylt     // Restrict block active input
29689c6efa4Sjeremylt     if (vec == CEED_VECTOR_ACTIVE) {
29789c6efa4Sjeremylt       ierr = CeedElemRestrictionApplyBlock(impl->blkrestr[i], e/blksize,
298a8d32208Sjeremylt                                            CEED_NOTRANSPOSE, invec,
29989c6efa4Sjeremylt                                            impl->evecsin[i], request);
300*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
30189c6efa4Sjeremylt       activein = 1;
30289c6efa4Sjeremylt     }
30389c6efa4Sjeremylt     // Basis action
30489c6efa4Sjeremylt     switch(emode) {
30589c6efa4Sjeremylt     case CEED_EVAL_NONE:
30689c6efa4Sjeremylt       if (!activein) {
30789c6efa4Sjeremylt         ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST,
30889c6efa4Sjeremylt                                   CEED_USE_POINTER,
309*e15f9bd0SJeremy L Thompson                                   &impl->edata[i][e*Q*size]); CeedChkBackend(ierr);
31089c6efa4Sjeremylt       }
31189c6efa4Sjeremylt       break;
31289c6efa4Sjeremylt     case CEED_EVAL_INTERP:
31389c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis);
314*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
31589c6efa4Sjeremylt       if (!activein) {
31689c6efa4Sjeremylt         ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
31789c6efa4Sjeremylt                                   CEED_USE_POINTER,
3184d537eeaSYohann                                   &impl->edata[i][e*elemsize*size]);
319*e15f9bd0SJeremy L Thompson         CeedChkBackend(ierr);
32089c6efa4Sjeremylt       }
32189c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
32289c6efa4Sjeremylt                             CEED_EVAL_INTERP, impl->evecsin[i],
323*e15f9bd0SJeremy L Thompson                             impl->qvecsin[i]); CeedChkBackend(ierr);
32489c6efa4Sjeremylt       break;
32589c6efa4Sjeremylt     case CEED_EVAL_GRAD:
32689c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis);
327*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
32889c6efa4Sjeremylt       if (!activein) {
329*e15f9bd0SJeremy L Thompson         ierr = CeedBasisGetDimension(basis, &dim); CeedChkBackend(ierr);
33089c6efa4Sjeremylt         ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
33189c6efa4Sjeremylt                                   CEED_USE_POINTER,
3324d537eeaSYohann                                   &impl->edata[i][e*elemsize*size/dim]);
333*e15f9bd0SJeremy L Thompson         CeedChkBackend(ierr);
33489c6efa4Sjeremylt       }
33589c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
33689c6efa4Sjeremylt                             CEED_EVAL_GRAD, impl->evecsin[i],
337*e15f9bd0SJeremy L Thompson                             impl->qvecsin[i]); CeedChkBackend(ierr);
33889c6efa4Sjeremylt       break;
33989c6efa4Sjeremylt     case CEED_EVAL_WEIGHT:
34089c6efa4Sjeremylt       break;  // No action
341bbfacfcdSjeremylt     // LCOV_EXCL_START
34289c6efa4Sjeremylt     case CEED_EVAL_DIV:
3431d102b48SJeremy L Thompson     case CEED_EVAL_CURL: {
3441d102b48SJeremy L Thompson       ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis);
345*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
3461d102b48SJeremy L Thompson       Ceed ceed;
347*e15f9bd0SJeremy L Thompson       ierr = CeedBasisGetCeed(basis, &ceed); CeedChkBackend(ierr);
348*e15f9bd0SJeremy L Thompson       return CeedError(ceed, CEED_ERROR_BACKEND,
349*e15f9bd0SJeremy L Thompson                        "Ceed evaluation mode not implemented");
350bbfacfcdSjeremylt       // LCOV_EXCL_STOP
3511d102b48SJeremy L Thompson     }
3521d102b48SJeremy L Thompson     }
3531d102b48SJeremy L Thompson   }
354*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
3551d102b48SJeremy L Thompson }
35689c6efa4Sjeremylt 
357f10650afSjeremylt //------------------------------------------------------------------------------
358f10650afSjeremylt // Output Basis Action
359f10650afSjeremylt //------------------------------------------------------------------------------
3601d102b48SJeremy L Thompson static inline int CeedOperatorOutputBasis_Opt(CeedInt e, CeedInt Q,
3611d102b48SJeremy L Thompson     CeedQFunctionField *qfoutputfields, CeedOperatorField *opoutputfields,
3621d102b48SJeremy L Thompson     CeedInt blksize, CeedInt numinputfields, CeedInt numoutputfields,
3631d102b48SJeremy L Thompson     CeedOperator op, CeedVector outvec, CeedOperator_Opt *impl,
3641d102b48SJeremy L Thompson     CeedRequest *request) {
3651d102b48SJeremy L Thompson   CeedInt ierr;
3661d102b48SJeremy L Thompson   CeedElemRestriction Erestrict;
3671d102b48SJeremy L Thompson   CeedEvalMode emode;
3681d102b48SJeremy L Thompson   CeedBasis basis;
3691d102b48SJeremy L Thompson   CeedVector vec;
3701d102b48SJeremy L Thompson 
37189c6efa4Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
3724d537eeaSYohann     // Get elemsize, emode, size
37389c6efa4Sjeremylt     ierr = CeedOperatorFieldGetElemRestriction(opoutputfields[i], &Erestrict);
374*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
37589c6efa4Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
376*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
37789c6efa4Sjeremylt     // Basis action
37889c6efa4Sjeremylt     switch(emode) {
37989c6efa4Sjeremylt     case CEED_EVAL_NONE:
38089c6efa4Sjeremylt       break; // No action
38189c6efa4Sjeremylt     case CEED_EVAL_INTERP:
38289c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
383*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
38489c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
38589c6efa4Sjeremylt                             CEED_EVAL_INTERP, impl->qvecsout[i],
386*e15f9bd0SJeremy L Thompson                             impl->evecsout[i]); CeedChkBackend(ierr);
38789c6efa4Sjeremylt       break;
38889c6efa4Sjeremylt     case CEED_EVAL_GRAD:
38989c6efa4Sjeremylt       ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
390*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
39189c6efa4Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
39289c6efa4Sjeremylt                             CEED_EVAL_GRAD, impl->qvecsout[i],
393*e15f9bd0SJeremy L Thompson                             impl->evecsout[i]); CeedChkBackend(ierr);
39489c6efa4Sjeremylt       break;
395c042f62fSJeremy L Thompson     // LCOV_EXCL_START
396bbfacfcdSjeremylt     case CEED_EVAL_WEIGHT: {
39789c6efa4Sjeremylt       Ceed ceed;
398*e15f9bd0SJeremy L Thompson       ierr = CeedOperatorGetCeed(op, &ceed); CeedChkBackend(ierr);
399*e15f9bd0SJeremy L Thompson       return CeedError(ceed, CEED_ERROR_BACKEND,
400*e15f9bd0SJeremy L Thompson                        "CEED_EVAL_WEIGHT cannot be an output "
4011d102b48SJeremy L Thompson                        "evaluation mode");
40289c6efa4Sjeremylt     }
40389c6efa4Sjeremylt     case CEED_EVAL_DIV:
4041d102b48SJeremy L Thompson     case CEED_EVAL_CURL: {
4051d102b48SJeremy L Thompson       Ceed ceed;
406*e15f9bd0SJeremy L Thompson       ierr = CeedOperatorGetCeed(op, &ceed); CeedChkBackend(ierr);
407*e15f9bd0SJeremy L Thompson       return CeedError(ceed, CEED_ERROR_BACKEND,
408*e15f9bd0SJeremy L Thompson                        "Ceed evaluation mode not implemented");
409bbfacfcdSjeremylt       // LCOV_EXCL_STOP
4101d102b48SJeremy L Thompson     }
41189c6efa4Sjeremylt     }
41289c6efa4Sjeremylt     // Restrict output block
41389c6efa4Sjeremylt     // Get output vector
414*e15f9bd0SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec);
415*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
41689c6efa4Sjeremylt     if (vec == CEED_VECTOR_ACTIVE)
41789c6efa4Sjeremylt       vec = outvec;
41889c6efa4Sjeremylt     // Restrict
41989c6efa4Sjeremylt     ierr = CeedElemRestrictionApplyBlock(impl->blkrestr[i+impl->numein],
42089c6efa4Sjeremylt                                          e/blksize, CEED_TRANSPOSE,
421a8d32208Sjeremylt                                          impl->evecsout[i], vec, request);
422*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
42389c6efa4Sjeremylt   }
424*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
42589c6efa4Sjeremylt }
42689c6efa4Sjeremylt 
427f10650afSjeremylt //------------------------------------------------------------------------------
428f10650afSjeremylt // Restore Input Vectors
429f10650afSjeremylt //------------------------------------------------------------------------------
4301d102b48SJeremy L Thompson static inline int CeedOperatorRestoreInputs_Opt(CeedInt numinputfields,
4311d102b48SJeremy L Thompson     CeedQFunctionField *qfinputfields, CeedOperatorField *opinputfields,
432187168c7SJeremy L Thompson     CeedOperator_Opt *impl) {
4331d102b48SJeremy L Thompson   CeedInt ierr;
4341d102b48SJeremy L Thompson   CeedEvalMode emode;
4351d102b48SJeremy L Thompson 
43689c6efa4Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
43789c6efa4Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
438*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
43989c6efa4Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
44089c6efa4Sjeremylt     } else {
44189c6efa4Sjeremylt       ierr = CeedVectorRestoreArrayRead(impl->evecs[i],
44289c6efa4Sjeremylt                                         (const CeedScalar **) &impl->edata[i]);
443*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
44489c6efa4Sjeremylt     }
44589c6efa4Sjeremylt   }
446*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
4471d102b48SJeremy L Thompson }
4481d102b48SJeremy L Thompson 
449f10650afSjeremylt //------------------------------------------------------------------------------
450f10650afSjeremylt // Operator Apply
451f10650afSjeremylt //------------------------------------------------------------------------------
45269af5e5fSJeremy L Thompson static int CeedOperatorApplyAdd_Opt(CeedOperator op, CeedVector invec,
4531d102b48SJeremy L Thompson                                     CeedVector outvec, CeedRequest *request) {
4541d102b48SJeremy L Thompson   int ierr;
4551d102b48SJeremy L Thompson   Ceed ceed;
456*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetCeed(op, &ceed); CeedChkBackend(ierr);
4571d102b48SJeremy L Thompson   Ceed_Opt *ceedimpl;
458*e15f9bd0SJeremy L Thompson   ierr = CeedGetData(ceed, &ceedimpl); CeedChkBackend(ierr);
4591d102b48SJeremy L Thompson   CeedInt blksize = ceedimpl->blksize;
4601d102b48SJeremy L Thompson   CeedOperator_Opt *impl;
461*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetData(op, &impl); CeedChkBackend(ierr);
4621d102b48SJeremy L Thompson   CeedInt Q, numinputfields, numoutputfields, numelements;
463*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetNumElements(op, &numelements); CeedChkBackend(ierr);
464*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChkBackend(ierr);
4651d102b48SJeremy L Thompson   CeedInt nblks = (numelements/blksize) + !!(numelements%blksize);
4661d102b48SJeremy L Thompson   CeedQFunction qf;
467*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetQFunction(op, &qf); CeedChkBackend(ierr);
4681d102b48SJeremy L Thompson   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
469*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
4701d102b48SJeremy L Thompson   CeedOperatorField *opinputfields, *opoutputfields;
4711d102b48SJeremy L Thompson   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
472*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
4731d102b48SJeremy L Thompson   CeedQFunctionField *qfinputfields, *qfoutputfields;
4741d102b48SJeremy L Thompson   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
475*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
4761d102b48SJeremy L Thompson   CeedEvalMode emode;
4771d102b48SJeremy L Thompson 
4781d102b48SJeremy L Thompson   // Setup
479*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorSetup_Opt(op); CeedChkBackend(ierr);
4801d102b48SJeremy L Thompson 
4811d102b48SJeremy L Thompson   // Input Evecs and Restriction
4821d102b48SJeremy L Thompson   ierr = CeedOperatorSetupInputs_Opt(numinputfields, qfinputfields,
48316911fdaSjeremylt                                      opinputfields, invec, impl, request);
484*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
4851d102b48SJeremy L Thompson 
4861d102b48SJeremy L Thompson   // Output Lvecs, Evecs, and Qvecs
4871d102b48SJeremy L Thompson   for (CeedInt i=0; i<numoutputfields; i++) {
4881d102b48SJeremy L Thompson     // Set Qvec if needed
4891d102b48SJeremy L Thompson     ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
490*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
4911d102b48SJeremy L Thompson     if (emode == CEED_EVAL_NONE) {
4921d102b48SJeremy L Thompson       // Set qvec to single block evec
4931d102b48SJeremy L Thompson       ierr = CeedVectorGetArray(impl->evecsout[i], CEED_MEM_HOST,
4941d102b48SJeremy L Thompson                                 &impl->edata[i + numinputfields]);
495*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
4961d102b48SJeremy L Thompson       ierr = CeedVectorSetArray(impl->qvecsout[i], CEED_MEM_HOST,
4971d102b48SJeremy L Thompson                                 CEED_USE_POINTER,
498*e15f9bd0SJeremy L Thompson                                 impl->edata[i + numinputfields]); CeedChkBackend(ierr);
4991d102b48SJeremy L Thompson       ierr = CeedVectorRestoreArray(impl->evecsout[i],
5001d102b48SJeremy L Thompson                                     &impl->edata[i + numinputfields]);
501*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
5021d102b48SJeremy L Thompson     }
5031d102b48SJeremy L Thompson   }
5041d102b48SJeremy L Thompson 
5051d102b48SJeremy L Thompson   // Loop through elements
5061d102b48SJeremy L Thompson   for (CeedInt e=0; e<nblks*blksize; e+=blksize) {
5071d102b48SJeremy L Thompson     // Input basis apply
5081d102b48SJeremy L Thompson     ierr = CeedOperatorInputBasis_Opt(e, Q, qfinputfields, opinputfields,
5091d102b48SJeremy L Thompson                                       numinputfields, blksize, invec, false,
510*e15f9bd0SJeremy L Thompson                                       impl, request); CeedChkBackend(ierr);
5111d102b48SJeremy L Thompson 
5121d102b48SJeremy L Thompson     // Q function
51316911fdaSjeremylt     if (!impl->identityqf) {
5141d102b48SJeremy L Thompson       ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout);
515*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
51616911fdaSjeremylt     }
5171d102b48SJeremy L Thompson 
5181d102b48SJeremy L Thompson     // Output basis apply and restrict
5191d102b48SJeremy L Thompson     ierr = CeedOperatorOutputBasis_Opt(e, Q, qfoutputfields, opoutputfields,
5207f823360Sjeremylt                                        blksize, numinputfields, numoutputfields,
5217f823360Sjeremylt                                        op, outvec, impl, request);
522*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
5231d102b48SJeremy L Thompson   }
5241d102b48SJeremy L Thompson 
5251d102b48SJeremy L Thompson   // Restore input arrays
5261d102b48SJeremy L Thompson   ierr = CeedOperatorRestoreInputs_Opt(numinputfields, qfinputfields,
527187168c7SJeremy L Thompson                                        opinputfields, impl);
528*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
52989c6efa4Sjeremylt 
530*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
53189c6efa4Sjeremylt }
53289c6efa4Sjeremylt 
533f10650afSjeremylt //------------------------------------------------------------------------------
5341d102b48SJeremy L Thompson // Assemble Linear QFunction
535f10650afSjeremylt //------------------------------------------------------------------------------
53680ac2e43SJeremy L Thompson static int CeedOperatorLinearAssembleQFunction_Opt(CeedOperator op,
5371d102b48SJeremy L Thompson     CeedVector *assembled, CeedElemRestriction *rstr, CeedRequest *request) {
5381d102b48SJeremy L Thompson   int ierr;
5391d102b48SJeremy L Thompson   Ceed ceed;
540*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetCeed(op, &ceed); CeedChkBackend(ierr);
5411d102b48SJeremy L Thompson   Ceed_Opt *ceedimpl;
542*e15f9bd0SJeremy L Thompson   ierr = CeedGetData(ceed, &ceedimpl); CeedChkBackend(ierr);
5431d102b48SJeremy L Thompson   const CeedInt blksize = ceedimpl->blksize;
5441d102b48SJeremy L Thompson   CeedOperator_Opt *impl;
545*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetData(op, &impl); CeedChkBackend(ierr);
5461d102b48SJeremy L Thompson   CeedInt Q, numinputfields, numoutputfields, numelements, size;
547*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetNumElements(op, &numelements); CeedChkBackend(ierr);
548*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChkBackend(ierr);
5491d102b48SJeremy L Thompson   CeedInt nblks = (numelements/blksize) + !!(numelements%blksize);
5501d102b48SJeremy L Thompson   CeedQFunction qf;
551*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetQFunction(op, &qf); CeedChkBackend(ierr);
5521d102b48SJeremy L Thompson   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
553*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
5541d102b48SJeremy L Thompson   CeedOperatorField *opinputfields, *opoutputfields;
5551d102b48SJeremy L Thompson   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
556*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
5571d102b48SJeremy L Thompson   CeedQFunctionField *qfinputfields, *qfoutputfields;
5581d102b48SJeremy L Thompson   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
559*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
5601d102b48SJeremy L Thompson   CeedVector vec, lvec;
5611d102b48SJeremy L Thompson   CeedInt numactivein = 0, numactiveout = 0;
56242ea3801Sjeremylt   CeedVector *activein = NULL;
5631d102b48SJeremy L Thompson   CeedScalar *a, *tmp;
5641d102b48SJeremy L Thompson 
5651d102b48SJeremy L Thompson   // Setup
566*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorSetup_Opt(op); CeedChkBackend(ierr);
5671d102b48SJeremy L Thompson 
56816911fdaSjeremylt   // Check for identity
56916911fdaSjeremylt   if (impl->identityqf)
57016911fdaSjeremylt     // LCOV_EXCL_START
571*e15f9bd0SJeremy L Thompson     return CeedError(ceed, CEED_ERROR_BACKEND,
572*e15f9bd0SJeremy L Thompson                      "Assembling identity qfunctions not supported");
57316911fdaSjeremylt   // LCOV_EXCL_STOP
57416911fdaSjeremylt 
5751d102b48SJeremy L Thompson   // Input Evecs and Restriction
5761d102b48SJeremy L Thompson   ierr = CeedOperatorSetupInputs_Opt(numinputfields, qfinputfields,
5771d102b48SJeremy L Thompson                                      opinputfields, NULL, impl, request);
578*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
5791d102b48SJeremy L Thompson 
5801d102b48SJeremy L Thompson   // Count number of active input fields
5811d102b48SJeremy L Thompson   for (CeedInt i=0; i<numinputfields; i++) {
5821d102b48SJeremy L Thompson     // Get input vector
583*e15f9bd0SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChkBackend(ierr);
5841d102b48SJeremy L Thompson     // Check if active input
5851d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
586*e15f9bd0SJeremy L Thompson       ierr = CeedQFunctionFieldGetSize(qfinputfields[i], &size); CeedChkBackend(ierr);
587*e15f9bd0SJeremy L Thompson       ierr = CeedVectorSetValue(impl->qvecsin[i], 0.0); CeedChkBackend(ierr);
5881d102b48SJeremy L Thompson       ierr = CeedVectorGetArray(impl->qvecsin[i], CEED_MEM_HOST, &tmp);
589*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
590*e15f9bd0SJeremy L Thompson       ierr = CeedRealloc(numactivein + size, &activein); CeedChkBackend(ierr);
5911d102b48SJeremy L Thompson       for (CeedInt field=0; field<size; field++) {
59242ea3801Sjeremylt         ierr = CeedVectorCreate(ceed, Q*blksize, &activein[numactivein+field]);
593*e15f9bd0SJeremy L Thompson         CeedChkBackend(ierr);
59442ea3801Sjeremylt         ierr = CeedVectorSetArray(activein[numactivein+field], CEED_MEM_HOST,
59542ea3801Sjeremylt                                   CEED_USE_POINTER, &tmp[field*Q*blksize]);
596*e15f9bd0SJeremy L Thompson         CeedChkBackend(ierr);
5971d102b48SJeremy L Thompson       }
5981d102b48SJeremy L Thompson       numactivein += size;
599*e15f9bd0SJeremy L Thompson       ierr = CeedVectorRestoreArray(impl->qvecsin[i], &tmp); CeedChkBackend(ierr);
6001d102b48SJeremy L Thompson     }
60189c6efa4Sjeremylt   }
60289c6efa4Sjeremylt 
6031d102b48SJeremy L Thompson   // Count number of active output fields
6041d102b48SJeremy L Thompson   for (CeedInt i=0; i<numoutputfields; i++) {
6051d102b48SJeremy L Thompson     // Get output vector
606*e15f9bd0SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec);
607*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
6081d102b48SJeremy L Thompson     // Check if active output
6091d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
610*e15f9bd0SJeremy L Thompson       ierr = CeedQFunctionFieldGetSize(qfoutputfields[i], &size);
611*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
6121d102b48SJeremy L Thompson       numactiveout += size;
6131d102b48SJeremy L Thompson     }
6141d102b48SJeremy L Thompson   }
6151d102b48SJeremy L Thompson 
6161d102b48SJeremy L Thompson   // Check sizes
6171d102b48SJeremy L Thompson   if (!numactivein || !numactiveout)
6181d102b48SJeremy L Thompson     // LCOV_EXCL_START
619*e15f9bd0SJeremy L Thompson     return CeedError(ceed, CEED_ERROR_BACKEND,
620*e15f9bd0SJeremy L Thompson                      "Cannot assemble QFunction without active inputs "
6211d102b48SJeremy L Thompson                      "and outputs");
6221d102b48SJeremy L Thompson   // LCOV_EXCL_STOP
6231d102b48SJeremy L Thompson 
6241d102b48SJeremy L Thompson   // Setup lvec
6251d102b48SJeremy L Thompson   ierr = CeedVectorCreate(ceed, nblks*blksize*Q*numactivein*numactiveout,
626*e15f9bd0SJeremy L Thompson                           &lvec); CeedChkBackend(ierr);
627*e15f9bd0SJeremy L Thompson   ierr = CeedVectorGetArray(lvec, CEED_MEM_HOST, &a); CeedChkBackend(ierr);
6281d102b48SJeremy L Thompson 
6291d102b48SJeremy L Thompson   // Create output restriction
6307509a596Sjeremylt   CeedInt strides[3] = {1, Q, numactivein *numactiveout*Q};
6317509a596Sjeremylt   ierr = CeedElemRestrictionCreateStrided(ceed, numelements, Q,
632d979a051Sjeremylt                                           numactivein*numactiveout,
633d979a051Sjeremylt                                           numactivein*numactiveout*numelements*Q,
634*e15f9bd0SJeremy L Thompson                                           strides, rstr); CeedChkBackend(ierr);
6351d102b48SJeremy L Thompson   // Create assembled vector
6361d102b48SJeremy L Thompson   ierr = CeedVectorCreate(ceed, numelements*Q*numactivein*numactiveout,
637*e15f9bd0SJeremy L Thompson                           assembled); CeedChkBackend(ierr);
6381d102b48SJeremy L Thompson 
6391d102b48SJeremy L Thompson   // Loop through elements
6401d102b48SJeremy L Thompson   for (CeedInt e=0; e<nblks*blksize; e+=blksize) {
6411d102b48SJeremy L Thompson     // Input basis apply
6421d102b48SJeremy L Thompson     ierr = CeedOperatorInputBasis_Opt(e, Q, qfinputfields, opinputfields,
6431d102b48SJeremy L Thompson                                       numinputfields, blksize, NULL, true,
644*e15f9bd0SJeremy L Thompson                                       impl, request); CeedChkBackend(ierr);
6451d102b48SJeremy L Thompson 
6461d102b48SJeremy L Thompson     // Assemble QFunction
6471d102b48SJeremy L Thompson     for (CeedInt in=0; in<numactivein; in++) {
6481d102b48SJeremy L Thompson       // Set Inputs
649*e15f9bd0SJeremy L Thompson       ierr = CeedVectorSetValue(activein[in], 1.0); CeedChkBackend(ierr);
65042ea3801Sjeremylt       if (numactivein > 1) {
65142ea3801Sjeremylt         ierr = CeedVectorSetValue(activein[(in+numactivein-1)%numactivein],
652*e15f9bd0SJeremy L Thompson                                   0.0); CeedChkBackend(ierr);
65342ea3801Sjeremylt       }
6541d102b48SJeremy L Thompson       // Set Outputs
6551d102b48SJeremy L Thompson       for (CeedInt out=0; out<numoutputfields; out++) {
6561d102b48SJeremy L Thompson         // Get output vector
6571d102b48SJeremy L Thompson         ierr = CeedOperatorFieldGetVector(opoutputfields[out], &vec);
658*e15f9bd0SJeremy L Thompson         CeedChkBackend(ierr);
6591d102b48SJeremy L Thompson         // Check if active output
6601d102b48SJeremy L Thompson         if (vec == CEED_VECTOR_ACTIVE) {
6611d102b48SJeremy L Thompson           CeedVectorSetArray(impl->qvecsout[out], CEED_MEM_HOST,
662*e15f9bd0SJeremy L Thompson                              CEED_USE_POINTER, a); CeedChkBackend(ierr);
6631d102b48SJeremy L Thompson           ierr = CeedQFunctionFieldGetSize(qfoutputfields[out], &size);
664*e15f9bd0SJeremy L Thompson           CeedChkBackend(ierr);
6651d102b48SJeremy L Thompson           a += size*Q*blksize; // Advance the pointer by the size of the output
6661d102b48SJeremy L Thompson         }
6671d102b48SJeremy L Thompson       }
6681d102b48SJeremy L Thompson       // Apply QFunction
6691d102b48SJeremy L Thompson       ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout);
670*e15f9bd0SJeremy L Thompson       CeedChkBackend(ierr);
6711d102b48SJeremy L Thompson     }
6721d102b48SJeremy L Thompson   }
6731d102b48SJeremy L Thompson 
6741d102b48SJeremy L Thompson   // Un-set output Qvecs to prevent accidental overwrite of Assembled
6751d102b48SJeremy L Thompson   for (CeedInt out=0; out<numoutputfields; out++) {
6761d102b48SJeremy L Thompson     // Get output vector
6771d102b48SJeremy L Thompson     ierr = CeedOperatorFieldGetVector(opoutputfields[out], &vec);
678*e15f9bd0SJeremy L Thompson     CeedChkBackend(ierr);
6791d102b48SJeremy L Thompson     // Check if active output
6801d102b48SJeremy L Thompson     if (vec == CEED_VECTOR_ACTIVE) {
6811d102b48SJeremy L Thompson       CeedVectorSetArray(impl->qvecsout[out], CEED_MEM_HOST, CEED_COPY_VALUES,
682*e15f9bd0SJeremy L Thompson                          NULL); CeedChkBackend(ierr);
6831d102b48SJeremy L Thompson     }
6841d102b48SJeremy L Thompson   }
6851d102b48SJeremy L Thompson 
6861d102b48SJeremy L Thompson   // Restore input arrays
6871d102b48SJeremy L Thompson   ierr = CeedOperatorRestoreInputs_Opt(numinputfields, qfinputfields,
688187168c7SJeremy L Thompson                                        opinputfields, impl);
689*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
6901d102b48SJeremy L Thompson 
6911d102b48SJeremy L Thompson   // Output blocked restriction
692*e15f9bd0SJeremy L Thompson   ierr = CeedVectorRestoreArray(lvec, &a); CeedChkBackend(ierr);
693*e15f9bd0SJeremy L Thompson   ierr = CeedVectorSetValue(*assembled, 0.0); CeedChkBackend(ierr);
6941d102b48SJeremy L Thompson   CeedElemRestriction blkrstr;
6957509a596Sjeremylt   ierr = CeedElemRestrictionCreateBlockedStrided(ceed, numelements, Q, blksize,
696d979a051Sjeremylt          numactivein*numactiveout, numactivein*numactiveout*numelements*Q,
697*e15f9bd0SJeremy L Thompson          strides, &blkrstr); CeedChkBackend(ierr);
698a8d32208Sjeremylt   ierr = CeedElemRestrictionApply(blkrstr, CEED_TRANSPOSE, lvec, *assembled,
699*e15f9bd0SJeremy L Thompson                                   request); CeedChkBackend(ierr);
7001d102b48SJeremy L Thompson 
7011d102b48SJeremy L Thompson   // Cleanup
70242ea3801Sjeremylt   for (CeedInt i=0; i<numactivein; i++) {
703*e15f9bd0SJeremy L Thompson     ierr = CeedVectorDestroy(&activein[i]); CeedChkBackend(ierr);
70442ea3801Sjeremylt   }
705*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&activein); CeedChkBackend(ierr);
706*e15f9bd0SJeremy L Thompson   ierr = CeedVectorDestroy(&lvec); CeedChkBackend(ierr);
707*e15f9bd0SJeremy L Thompson   ierr = CeedElemRestrictionDestroy(&blkrstr); CeedChkBackend(ierr);
7081d102b48SJeremy L Thompson 
709*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
71089c6efa4Sjeremylt }
71189c6efa4Sjeremylt 
712f10650afSjeremylt //------------------------------------------------------------------------------
713f10650afSjeremylt // Operator Destroy
714f10650afSjeremylt //------------------------------------------------------------------------------
715f10650afSjeremylt static int CeedOperatorDestroy_Opt(CeedOperator op) {
716f10650afSjeremylt   int ierr;
717f10650afSjeremylt   CeedOperator_Opt *impl;
718*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetData(op, &impl); CeedChkBackend(ierr);
719f10650afSjeremylt 
720f10650afSjeremylt   for (CeedInt i=0; i<impl->numein+impl->numeout; i++) {
721*e15f9bd0SJeremy L Thompson     ierr = CeedElemRestrictionDestroy(&impl->blkrestr[i]); CeedChkBackend(ierr);
722*e15f9bd0SJeremy L Thompson     ierr = CeedVectorDestroy(&impl->evecs[i]); CeedChkBackend(ierr);
723f10650afSjeremylt   }
724*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&impl->blkrestr); CeedChkBackend(ierr);
725*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&impl->evecs); CeedChkBackend(ierr);
726*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&impl->edata); CeedChkBackend(ierr);
727*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&impl->inputstate); CeedChkBackend(ierr);
728f10650afSjeremylt 
729f10650afSjeremylt   for (CeedInt i=0; i<impl->numein; i++) {
730*e15f9bd0SJeremy L Thompson     ierr = CeedVectorDestroy(&impl->evecsin[i]); CeedChkBackend(ierr);
731*e15f9bd0SJeremy L Thompson     ierr = CeedVectorDestroy(&impl->qvecsin[i]); CeedChkBackend(ierr);
732f10650afSjeremylt   }
733*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&impl->evecsin); CeedChkBackend(ierr);
734*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&impl->qvecsin); CeedChkBackend(ierr);
735f10650afSjeremylt 
736f10650afSjeremylt   for (CeedInt i=0; i<impl->numeout; i++) {
737*e15f9bd0SJeremy L Thompson     ierr = CeedVectorDestroy(&impl->evecsout[i]); CeedChkBackend(ierr);
738*e15f9bd0SJeremy L Thompson     ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChkBackend(ierr);
739f10650afSjeremylt   }
740*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&impl->evecsout); CeedChkBackend(ierr);
741*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&impl->qvecsout); CeedChkBackend(ierr);
742f10650afSjeremylt 
743*e15f9bd0SJeremy L Thompson   ierr = CeedFree(&impl); CeedChkBackend(ierr);
744*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
745f10650afSjeremylt }
746f10650afSjeremylt 
747f10650afSjeremylt //------------------------------------------------------------------------------
748f10650afSjeremylt // Operator Create
749f10650afSjeremylt //------------------------------------------------------------------------------
75089c6efa4Sjeremylt int CeedOperatorCreate_Opt(CeedOperator op) {
75189c6efa4Sjeremylt   int ierr;
75289c6efa4Sjeremylt   Ceed ceed;
753*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorGetCeed(op, &ceed); CeedChkBackend(ierr);
75489c6efa4Sjeremylt   Ceed_Opt *ceedimpl;
755*e15f9bd0SJeremy L Thompson   ierr = CeedGetData(ceed, &ceedimpl); CeedChkBackend(ierr);
75689c6efa4Sjeremylt   CeedInt blksize = ceedimpl->blksize;
75789c6efa4Sjeremylt   CeedOperator_Opt *impl;
75889c6efa4Sjeremylt 
759*e15f9bd0SJeremy L Thompson   ierr = CeedCalloc(1, &impl); CeedChkBackend(ierr);
760*e15f9bd0SJeremy L Thompson   ierr = CeedOperatorSetData(op, impl); CeedChkBackend(ierr);
76189c6efa4Sjeremylt 
76282946b17Sjeremylt   if (blksize != 1 && blksize != 8)
76382946b17Sjeremylt     // LCOV_EXCL_START
764*e15f9bd0SJeremy L Thompson     return CeedError(ceed, CEED_ERROR_BACKEND,
765*e15f9bd0SJeremy L Thompson                      "Opt backend cannot use blocksize: %d", blksize);
76682946b17Sjeremylt   // LCOV_EXCL_STOP
76782946b17Sjeremylt 
76880ac2e43SJeremy L Thompson   ierr = CeedSetBackendFunction(ceed, "Operator", op, "LinearAssembleQFunction",
76980ac2e43SJeremy L Thompson                                 CeedOperatorLinearAssembleQFunction_Opt);
770*e15f9bd0SJeremy L Thompson   CeedChkBackend(ierr);
771cae8b89aSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "ApplyAdd",
772*e15f9bd0SJeremy L Thompson                                 CeedOperatorApplyAdd_Opt); CeedChkBackend(ierr);
77389c6efa4Sjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Destroy",
774*e15f9bd0SJeremy L Thompson                                 CeedOperatorDestroy_Opt); CeedChkBackend(ierr);
775*e15f9bd0SJeremy L Thompson   return CEED_ERROR_SUCCESS;
77689c6efa4Sjeremylt }
777f10650afSjeremylt //------------------------------------------------------------------------------
778