xref: /libCEED/rust/libceed-sys/c-src/backends/blocked/ceed-blocked-operator.c (revision c042f62f62e10e1321eb699b116e67a6568d5716) !
14a2e7687Sjeremylt // Copyright (c) 2017-2018, Lawrence Livermore National Security, LLC.
24a2e7687Sjeremylt // Produced at the Lawrence Livermore National Laboratory. LLNL-CODE-734707.
34a2e7687Sjeremylt // All Rights reserved. See files LICENSE and NOTICE for details.
44a2e7687Sjeremylt //
54a2e7687Sjeremylt // This file is part of CEED, a collection of benchmarks, miniapps, software
64a2e7687Sjeremylt // libraries and APIs for efficient high-order finite element and spectral
74a2e7687Sjeremylt // element discretizations for exascale applications. For more information and
84a2e7687Sjeremylt // source code availability see http://github.com/ceed.
94a2e7687Sjeremylt //
104a2e7687Sjeremylt // The CEED research is supported by the Exascale Computing Project 17-SC-20-SC,
114a2e7687Sjeremylt // a collaborative effort of two U.S. Department of Energy organizations (Office
124a2e7687Sjeremylt // of Science and the National Nuclear Security Administration) responsible for
134a2e7687Sjeremylt // the planning and preparation of a capable exascale ecosystem, including
144a2e7687Sjeremylt // software, applications, hardware, advanced system engineering and early
154a2e7687Sjeremylt // testbed platforms, in support of the nation's exascale computing imperative.
164a2e7687Sjeremylt 
174a2e7687Sjeremylt #include "ceed-blocked.h"
184a2e7687Sjeremylt #include "../ref/ceed-ref.h"
194a2e7687Sjeremylt 
204a2e7687Sjeremylt static int CeedOperatorDestroy_Blocked(CeedOperator op) {
214a2e7687Sjeremylt   int ierr;
224ce2993fSjeremylt   CeedOperator_Blocked *impl;
234ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
244a2e7687Sjeremylt 
254a2e7687Sjeremylt   for (CeedInt i=0; i<impl->numein+impl->numeout; i++) {
264a2e7687Sjeremylt     ierr = CeedElemRestrictionDestroy(&impl->blkrestr[i]); CeedChk(ierr);
274a2e7687Sjeremylt     ierr = CeedVectorDestroy(&impl->evecs[i]); CeedChk(ierr);
284a2e7687Sjeremylt   }
294a2e7687Sjeremylt   ierr = CeedFree(&impl->blkrestr); CeedChk(ierr);
304a2e7687Sjeremylt   ierr = CeedFree(&impl->evecs); CeedChk(ierr);
314a2e7687Sjeremylt   ierr = CeedFree(&impl->edata); CeedChk(ierr);
3216c359e6Sjeremylt   ierr = CeedFree(&impl->inputstate); CeedChk(ierr);
334a2e7687Sjeremylt 
34aedaa0e5Sjeremylt   for (CeedInt i=0; i<impl->numein; i++) {
3591703d3fSjeremylt     ierr = CeedVectorDestroy(&impl->evecsin[i]); CeedChk(ierr);
36aedaa0e5Sjeremylt     ierr = CeedVectorDestroy(&impl->qvecsin[i]); CeedChk(ierr);
374a2e7687Sjeremylt   }
3891703d3fSjeremylt   ierr = CeedFree(&impl->evecsin); CeedChk(ierr);
39aedaa0e5Sjeremylt   ierr = CeedFree(&impl->qvecsin); CeedChk(ierr);
404a2e7687Sjeremylt 
41aedaa0e5Sjeremylt   for (CeedInt i=0; i<impl->numeout; i++) {
4291703d3fSjeremylt     ierr = CeedVectorDestroy(&impl->evecsout[i]); CeedChk(ierr);
43aedaa0e5Sjeremylt     ierr = CeedVectorDestroy(&impl->qvecsout[i]); CeedChk(ierr);
44aedaa0e5Sjeremylt   }
4591703d3fSjeremylt   ierr = CeedFree(&impl->evecsout); CeedChk(ierr);
46aedaa0e5Sjeremylt   ierr = CeedFree(&impl->qvecsout); CeedChk(ierr);
474a2e7687Sjeremylt 
48fe2413ffSjeremylt   ierr = CeedFree(&impl); CeedChk(ierr);
494a2e7687Sjeremylt   return 0;
504a2e7687Sjeremylt }
514a2e7687Sjeremylt 
524a2e7687Sjeremylt /*
534a2e7687Sjeremylt   Setup infields or outfields
544a2e7687Sjeremylt  */
5589c6efa4Sjeremylt static int CeedOperatorSetupFields_Blocked(CeedQFunction qf,
5689c6efa4Sjeremylt     CeedOperator op, bool inOrOut,
574a2e7687Sjeremylt     CeedElemRestriction *blkrestr,
5891703d3fSjeremylt     CeedVector *fullevecs, CeedVector *evecs,
59aedaa0e5Sjeremylt     CeedVector *qvecs, CeedInt starte,
604a2e7687Sjeremylt     CeedInt numfields, CeedInt Q) {
614d537eeaSYohann   CeedInt dim, ierr, ncomp, size, P;
62aedaa0e5Sjeremylt   Ceed ceed;
63aedaa0e5Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
64d1bcdac9Sjeremylt   CeedBasis basis;
65d1bcdac9Sjeremylt   CeedElemRestriction r;
66aedaa0e5Sjeremylt   CeedOperatorField *opfields;
67aedaa0e5Sjeremylt   CeedQFunctionField *qffields;
68fe2413ffSjeremylt   if (inOrOut) {
69aedaa0e5Sjeremylt     ierr = CeedOperatorGetFields(op, NULL, &opfields);
70fe2413ffSjeremylt     CeedChk(ierr);
71aedaa0e5Sjeremylt     ierr = CeedQFunctionGetFields(qf, NULL, &qffields);
72fe2413ffSjeremylt     CeedChk(ierr);
73fe2413ffSjeremylt   } else {
74aedaa0e5Sjeremylt     ierr = CeedOperatorGetFields(op, &opfields, NULL);
75fe2413ffSjeremylt     CeedChk(ierr);
76aedaa0e5Sjeremylt     ierr = CeedQFunctionGetFields(qf, &qffields, NULL);
77fe2413ffSjeremylt     CeedChk(ierr);
78fe2413ffSjeremylt   }
794a2e7687Sjeremylt   const CeedInt blksize = 8;
804a2e7687Sjeremylt 
814a2e7687Sjeremylt   // Loop over fields
824a2e7687Sjeremylt   for (CeedInt i=0; i<numfields; i++) {
83d1bcdac9Sjeremylt     CeedEvalMode emode;
84aedaa0e5Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qffields[i], &emode); CeedChk(ierr);
854a2e7687Sjeremylt 
864a2e7687Sjeremylt     if (emode != CEED_EVAL_WEIGHT) {
87aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opfields[i], &r);
88d1bcdac9Sjeremylt       CeedChk(ierr);
89fe2413ffSjeremylt       CeedElemRestriction_Ref *data;
90de686571SJeremy L Thompson       ierr = CeedElemRestrictionGetData(r, (void *)&data); CeedChk(ierr);
914ce2993fSjeremylt       Ceed ceed;
924ce2993fSjeremylt       ierr = CeedElemRestrictionGetCeed(r, &ceed); CeedChk(ierr);
938795c945Sjeremylt       CeedInt nelem, elemsize, nnodes;
944ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumElements(r, &nelem); CeedChk(ierr);
954ce2993fSjeremylt       ierr = CeedElemRestrictionGetElementSize(r, &elemsize); CeedChk(ierr);
968795c945Sjeremylt       ierr = CeedElemRestrictionGetNumNodes(r, &nnodes); CeedChk(ierr);
974ce2993fSjeremylt       ierr = CeedElemRestrictionGetNumComponents(r, &ncomp); CeedChk(ierr);
984ce2993fSjeremylt       ierr = CeedElemRestrictionCreateBlocked(ceed, nelem, elemsize,
998795c945Sjeremylt                                               blksize, nnodes, ncomp,
1004a2e7687Sjeremylt                                               CEED_MEM_HOST, CEED_COPY_VALUES,
101aedaa0e5Sjeremylt                                               data->indices, &blkrestr[i+starte]);
1024a2e7687Sjeremylt       CeedChk(ierr);
103aedaa0e5Sjeremylt       ierr = CeedElemRestrictionCreateVector(blkrestr[i+starte], NULL,
10491703d3fSjeremylt                                              &fullevecs[i+starte]);
1054a2e7687Sjeremylt       CeedChk(ierr);
1064a2e7687Sjeremylt     }
1074a2e7687Sjeremylt 
1084a2e7687Sjeremylt     switch(emode) {
1094a2e7687Sjeremylt     case CEED_EVAL_NONE:
1104d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChk(ierr);
1114d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChk(ierr);
112aedaa0e5Sjeremylt       break;
113aedaa0e5Sjeremylt     case CEED_EVAL_INTERP:
1144d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChk(ierr);
1154d1cd9fcSJeremy L Thompson       ierr = CeedElemRestrictionGetElementSize(r, &P);
1164d1cd9fcSJeremy L Thompson       CeedChk(ierr);
1174d537eeaSYohann       ierr = CeedVectorCreate(ceed, P*size*blksize, &evecs[i]); CeedChk(ierr);
1184d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChk(ierr);
1194a2e7687Sjeremylt       break;
1204a2e7687Sjeremylt     case CEED_EVAL_GRAD:
121aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
1224d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qffields[i], &size); CeedChk(ierr);
123d1bcdac9Sjeremylt       ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
1244d1cd9fcSJeremy L Thompson       ierr = CeedElemRestrictionGetElementSize(r, &P);
1254d1cd9fcSJeremy L Thompson       CeedChk(ierr);
1264d537eeaSYohann       ierr = CeedVectorCreate(ceed, P*size/dim*blksize, &evecs[i]); CeedChk(ierr);
1274d537eeaSYohann       ierr = CeedVectorCreate(ceed, Q*size*blksize, &qvecs[i]); CeedChk(ierr);
1284a2e7687Sjeremylt       break;
1294a2e7687Sjeremylt     case CEED_EVAL_WEIGHT: // Only on input fields
130aedaa0e5Sjeremylt       ierr = CeedOperatorFieldGetBasis(opfields[i], &basis); CeedChk(ierr);
131aedaa0e5Sjeremylt       ierr = CeedVectorCreate(ceed, Q*blksize, &qvecs[i]); CeedChk(ierr);
132d1bcdac9Sjeremylt       ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
133aedaa0e5Sjeremylt                             CEED_EVAL_WEIGHT, NULL, qvecs[i]); CeedChk(ierr);
134aedaa0e5Sjeremylt 
1354a2e7687Sjeremylt       break;
1364a2e7687Sjeremylt     case CEED_EVAL_DIV:
1374d537eeaSYohann       break; // Not implemented
1384a2e7687Sjeremylt     case CEED_EVAL_CURL:
1394d537eeaSYohann       break; // Not implemented
1404a2e7687Sjeremylt     }
1414a2e7687Sjeremylt   }
1424a2e7687Sjeremylt   return 0;
1434a2e7687Sjeremylt }
1444a2e7687Sjeremylt 
1454a2e7687Sjeremylt /*
1464a2e7687Sjeremylt   CeedOperator needs to connect all the named fields (be they active or passive)
1474a2e7687Sjeremylt   to the named inputs and outputs of its CeedQFunction.
1484a2e7687Sjeremylt  */
1494a2e7687Sjeremylt static int CeedOperatorSetup_Blocked(CeedOperator op) {
1504a2e7687Sjeremylt   int ierr;
1514ce2993fSjeremylt   bool setupdone;
1524ce2993fSjeremylt   ierr = CeedOperatorGetSetupStatus(op, &setupdone); CeedChk(ierr);
1534ce2993fSjeremylt   if (setupdone) return 0;
154aedaa0e5Sjeremylt   Ceed ceed;
155aedaa0e5Sjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
1564ce2993fSjeremylt   CeedOperator_Blocked *impl;
1574ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
1584ce2993fSjeremylt   CeedQFunction qf;
1594ce2993fSjeremylt   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
1604ce2993fSjeremylt   CeedInt Q, numinputfields, numoutputfields;
1614ce2993fSjeremylt   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
1624a2e7687Sjeremylt   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
1634a2e7687Sjeremylt   CeedChk(ierr);
164d1bcdac9Sjeremylt   CeedOperatorField *opinputfields, *opoutputfields;
165d1bcdac9Sjeremylt   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
166d1bcdac9Sjeremylt   CeedChk(ierr);
167d1bcdac9Sjeremylt   CeedQFunctionField *qfinputfields, *qfoutputfields;
168d1bcdac9Sjeremylt   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
169d1bcdac9Sjeremylt   CeedChk(ierr);
1704a2e7687Sjeremylt 
1714a2e7687Sjeremylt   // Allocate
172aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->blkrestr);
1734a2e7687Sjeremylt   CeedChk(ierr);
174aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->evecs);
1754a2e7687Sjeremylt   CeedChk(ierr);
176aedaa0e5Sjeremylt   ierr = CeedCalloc(numinputfields + numoutputfields, &impl->edata);
1774a2e7687Sjeremylt   CeedChk(ierr);
1784a2e7687Sjeremylt 
17916c359e6Sjeremylt   ierr = CeedCalloc(16, &impl->inputstate); CeedChk(ierr);
18091703d3fSjeremylt   ierr = CeedCalloc(16, &impl->evecsin); CeedChk(ierr);
18191703d3fSjeremylt   ierr = CeedCalloc(16, &impl->evecsout); CeedChk(ierr);
182aedaa0e5Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsin); CeedChk(ierr);
183aedaa0e5Sjeremylt   ierr = CeedCalloc(16, &impl->qvecsout); CeedChk(ierr);
1844a2e7687Sjeremylt 
185aedaa0e5Sjeremylt   impl->numein = numinputfields; impl->numeout = numoutputfields;
186aedaa0e5Sjeremylt 
1874a2e7687Sjeremylt   // Set up infield and outfield pointer arrays
1884a2e7687Sjeremylt   // Infields
189aedaa0e5Sjeremylt   ierr = CeedOperatorSetupFields_Blocked(qf, op, 0, impl->blkrestr,
19091703d3fSjeremylt                                          impl->evecs, impl->evecsin,
19191703d3fSjeremylt                                          impl->qvecsin, 0,
192aedaa0e5Sjeremylt                                          numinputfields, Q);
1934a2e7687Sjeremylt   CeedChk(ierr);
1944a2e7687Sjeremylt   // Outfields
195aedaa0e5Sjeremylt   ierr = CeedOperatorSetupFields_Blocked(qf, op, 1, impl->blkrestr,
19691703d3fSjeremylt                                          impl->evecs, impl->evecsout,
19791703d3fSjeremylt                                          impl->qvecsout, numinputfields,
19891703d3fSjeremylt                                          numoutputfields, Q);
1994a2e7687Sjeremylt   CeedChk(ierr);
200aedaa0e5Sjeremylt 
2014ce2993fSjeremylt   ierr = CeedOperatorSetSetupDone(op); CeedChk(ierr);
2024a2e7687Sjeremylt 
2034a2e7687Sjeremylt   return 0;
2044a2e7687Sjeremylt }
2054a2e7687Sjeremylt 
2064a2e7687Sjeremylt static int CeedOperatorApply_Blocked(CeedOperator op, CeedVector invec,
20789c6efa4Sjeremylt                                      CeedVector outvec,
20889c6efa4Sjeremylt                                      CeedRequest *request) {
2094a2e7687Sjeremylt   int ierr;
2104ce2993fSjeremylt   CeedOperator_Blocked *impl;
2114ce2993fSjeremylt   ierr = CeedOperatorGetData(op, (void *)&impl); CeedChk(ierr);
2124ce2993fSjeremylt   const CeedInt blksize = 8;
2134d537eeaSYohann   CeedInt Q, elemsize, numinputfields, numoutputfields, numelements, size, dim;
2144ce2993fSjeremylt   ierr = CeedOperatorGetNumElements(op, &numelements); CeedChk(ierr);
2154ce2993fSjeremylt   ierr = CeedOperatorGetNumQuadraturePoints(op, &Q); CeedChk(ierr);
2164ce2993fSjeremylt   CeedInt nblks = (numelements/blksize) + !!(numelements%blksize);
2174ce2993fSjeremylt   CeedQFunction qf;
2184ce2993fSjeremylt   ierr = CeedOperatorGetQFunction(op, &qf); CeedChk(ierr);
2194ce2993fSjeremylt   ierr= CeedQFunctionGetNumArgs(qf, &numinputfields, &numoutputfields);
2204ce2993fSjeremylt   CeedChk(ierr);
2214dccadb6Sjeremylt   CeedTransposeMode lmode;
222d1bcdac9Sjeremylt   CeedOperatorField *opinputfields, *opoutputfields;
223d1bcdac9Sjeremylt   ierr = CeedOperatorGetFields(op, &opinputfields, &opoutputfields);
224d1bcdac9Sjeremylt   CeedChk(ierr);
225d1bcdac9Sjeremylt   CeedQFunctionField *qfinputfields, *qfoutputfields;
226d1bcdac9Sjeremylt   ierr = CeedQFunctionGetFields(qf, &qfinputfields, &qfoutputfields);
227d1bcdac9Sjeremylt   CeedChk(ierr);
228d1bcdac9Sjeremylt   CeedEvalMode emode;
229d1bcdac9Sjeremylt   CeedVector vec;
230d1bcdac9Sjeremylt   CeedBasis basis;
231d1bcdac9Sjeremylt   CeedElemRestriction Erestrict;
23216c359e6Sjeremylt   uint64_t state;
2334a2e7687Sjeremylt 
2344a2e7687Sjeremylt   // Setup
2354a2e7687Sjeremylt   ierr = CeedOperatorSetup_Blocked(op); CeedChk(ierr);
2364a2e7687Sjeremylt 
2374a2e7687Sjeremylt   // Input Evecs and Restriction
2384a2e7687Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
239d1bcdac9Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
240d1bcdac9Sjeremylt     CeedChk(ierr);
2414a2e7687Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
2424a2e7687Sjeremylt     } else {
243d1bcdac9Sjeremylt       // Get input vector
244d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetVector(opinputfields[i], &vec); CeedChk(ierr);
24589c6efa4Sjeremylt       if (vec == CEED_VECTOR_ACTIVE)
24689c6efa4Sjeremylt         vec = invec;
2474a2e7687Sjeremylt       // Restrict
24816c359e6Sjeremylt       ierr = CeedVectorGetState(vec, &state); CeedChk(ierr);
24989c6efa4Sjeremylt       if (state != impl->inputstate[i] || vec == invec) {
25089c6efa4Sjeremylt         ierr = CeedOperatorFieldGetLMode(opinputfields[i], &lmode); CeedChk(ierr);
2514a2e7687Sjeremylt         ierr = CeedElemRestrictionApply(impl->blkrestr[i], CEED_NOTRANSPOSE,
25289c6efa4Sjeremylt                                         lmode, vec, impl->evecs[i],
25389c6efa4Sjeremylt                                         request); CeedChk(ierr); CeedChk(ierr);
25416c359e6Sjeremylt         impl->inputstate[i] = state;
25516c359e6Sjeremylt       }
2564a2e7687Sjeremylt       // Get evec
2574a2e7687Sjeremylt       ierr = CeedVectorGetArrayRead(impl->evecs[i], CEED_MEM_HOST,
2584a2e7687Sjeremylt                                     (const CeedScalar **) &impl->edata[i]);
2594a2e7687Sjeremylt       CeedChk(ierr);
2604a2e7687Sjeremylt     }
2614a2e7687Sjeremylt   }
2624a2e7687Sjeremylt 
26389c6efa4Sjeremylt   // Output Evecs
2644a2e7687Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
26589c6efa4Sjeremylt     ierr = CeedVectorGetArray(impl->evecs[i+impl->numein], CEED_MEM_HOST,
26689c6efa4Sjeremylt                               &impl->edata[i + numinputfields]); CeedChk(ierr);
2674a2e7687Sjeremylt   }
2684a2e7687Sjeremylt 
2694a2e7687Sjeremylt   // Loop through elements
2704a2e7687Sjeremylt   for (CeedInt e=0; e<nblks*blksize; e+=blksize) {
2714a2e7687Sjeremylt     // Input basis apply if needed
2724a2e7687Sjeremylt     for (CeedInt i=0; i<numinputfields; i++) {
2734d537eeaSYohann       // Get elemsize, emode, size
274d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opinputfields[i], &Erestrict);
275d1bcdac9Sjeremylt       CeedChk(ierr);
276d1bcdac9Sjeremylt       ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
277d1bcdac9Sjeremylt       CeedChk(ierr);
278d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
279d1bcdac9Sjeremylt       CeedChk(ierr);
2804d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qfinputfields[i], &size); CeedChk(ierr);
2814a2e7687Sjeremylt       // Basis action
2824a2e7687Sjeremylt       switch(emode) {
2834a2e7687Sjeremylt       case CEED_EVAL_NONE:
284aedaa0e5Sjeremylt         ierr = CeedVectorSetArray(impl->qvecsin[i], CEED_MEM_HOST,
285aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
2864d537eeaSYohann                                   &impl->edata[i][e*Q*size]); CeedChk(ierr);
2874a2e7687Sjeremylt         break;
2884a2e7687Sjeremylt       case CEED_EVAL_INTERP:
28989c6efa4Sjeremylt         ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr);
29091703d3fSjeremylt         ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
291aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
2924d537eeaSYohann                                   &impl->edata[i][e*elemsize*size]);
2934a2e7687Sjeremylt         CeedChk(ierr);
294d1bcdac9Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
29591703d3fSjeremylt                               CEED_EVAL_INTERP, impl->evecsin[i],
296aedaa0e5Sjeremylt                               impl->qvecsin[i]); CeedChk(ierr);
2974a2e7687Sjeremylt         break;
2984a2e7687Sjeremylt       case CEED_EVAL_GRAD:
29989c6efa4Sjeremylt         ierr = CeedOperatorFieldGetBasis(opinputfields[i], &basis); CeedChk(ierr);
3004d537eeaSYohann         ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
30191703d3fSjeremylt         ierr = CeedVectorSetArray(impl->evecsin[i], CEED_MEM_HOST,
302aedaa0e5Sjeremylt                                   CEED_USE_POINTER,
3034d537eeaSYohann                                   &impl->edata[i][e*elemsize*size/dim]);
3044a2e7687Sjeremylt         CeedChk(ierr);
305d1bcdac9Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_NOTRANSPOSE,
30691703d3fSjeremylt                               CEED_EVAL_GRAD, impl->evecsin[i],
307aedaa0e5Sjeremylt                               impl->qvecsin[i]); CeedChk(ierr);
3084a2e7687Sjeremylt         break;
3094a2e7687Sjeremylt       case CEED_EVAL_WEIGHT:
3104a2e7687Sjeremylt         break;  // No action
3114a2e7687Sjeremylt       case CEED_EVAL_DIV:
3128c91a0c9SJeremy L Thompson         break; // Not implemented
3134a2e7687Sjeremylt       case CEED_EVAL_CURL:
3148c91a0c9SJeremy L Thompson         break; // Not implemented
3154a2e7687Sjeremylt       }
3164a2e7687Sjeremylt     }
3174a2e7687Sjeremylt 
31889c6efa4Sjeremylt     // Output pointers
31989c6efa4Sjeremylt     for (CeedInt i=0; i<numoutputfields; i++) {
32089c6efa4Sjeremylt       ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
32189c6efa4Sjeremylt       CeedChk(ierr);
32289c6efa4Sjeremylt       if (emode == CEED_EVAL_NONE) {
3234d537eeaSYohann         ierr = CeedQFunctionFieldGetSize(qfoutputfields[i], &size);
32489c6efa4Sjeremylt         CeedChk(ierr);
32589c6efa4Sjeremylt         ierr = CeedVectorSetArray(impl->qvecsout[i], CEED_MEM_HOST,
32689c6efa4Sjeremylt                                   CEED_USE_POINTER,
3274d537eeaSYohann                                   &impl->edata[i + numinputfields][e*Q*size]);
32889c6efa4Sjeremylt         CeedChk(ierr);
32989c6efa4Sjeremylt       }
33089c6efa4Sjeremylt     }
3314a2e7687Sjeremylt     // Q function
332aedaa0e5Sjeremylt     ierr = CeedQFunctionApply(qf, Q*blksize, impl->qvecsin, impl->qvecsout);
333aedaa0e5Sjeremylt     CeedChk(ierr);
3344a2e7687Sjeremylt 
33589c6efa4Sjeremylt     // Output basis apply if needed
3364a2e7687Sjeremylt     for (CeedInt i=0; i<numoutputfields; i++) {
3374d537eeaSYohann       // Get elemsize, emode, size
338d1bcdac9Sjeremylt       ierr = CeedOperatorFieldGetElemRestriction(opoutputfields[i], &Erestrict);
339d1bcdac9Sjeremylt       CeedChk(ierr);
34089c6efa4Sjeremylt       ierr = CeedElemRestrictionGetElementSize(Erestrict, &elemsize);
34189c6efa4Sjeremylt       CeedChk(ierr);
342d1bcdac9Sjeremylt       ierr = CeedQFunctionFieldGetEvalMode(qfoutputfields[i], &emode);
343d1bcdac9Sjeremylt       CeedChk(ierr);
3444d537eeaSYohann       ierr = CeedQFunctionFieldGetSize(qfoutputfields[i], &size); CeedChk(ierr);
3454a2e7687Sjeremylt       // Basis action
3464a2e7687Sjeremylt       switch(emode) {
3474a2e7687Sjeremylt       case CEED_EVAL_NONE:
3484a2e7687Sjeremylt         break; // No action
3494a2e7687Sjeremylt       case CEED_EVAL_INTERP:
350d1bcdac9Sjeremylt         ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
351d1bcdac9Sjeremylt         CeedChk(ierr);
35289c6efa4Sjeremylt         ierr = CeedVectorSetArray(impl->evecsout[i], CEED_MEM_HOST,
35389c6efa4Sjeremylt                                   CEED_USE_POINTER,
3544d537eeaSYohann                                   &impl->edata[i + numinputfields][e*elemsize*size]);
35589c6efa4Sjeremylt         CeedChk(ierr);
356aedaa0e5Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
357aedaa0e5Sjeremylt                               CEED_EVAL_INTERP, impl->qvecsout[i],
35891703d3fSjeremylt                               impl->evecsout[i]); CeedChk(ierr);
3594a2e7687Sjeremylt         break;
3604a2e7687Sjeremylt       case CEED_EVAL_GRAD:
361d1bcdac9Sjeremylt         ierr = CeedOperatorFieldGetBasis(opoutputfields[i], &basis);
362d1bcdac9Sjeremylt         CeedChk(ierr);
3634d537eeaSYohann         ierr = CeedBasisGetDimension(basis, &dim); CeedChk(ierr);
36489c6efa4Sjeremylt         ierr = CeedVectorSetArray(impl->evecsout[i], CEED_MEM_HOST,
36589c6efa4Sjeremylt                                   CEED_USE_POINTER,
3664d537eeaSYohann                                   &impl->edata[i + numinputfields][e*elemsize*size/dim]);
36789c6efa4Sjeremylt         CeedChk(ierr);
368d1bcdac9Sjeremylt         ierr = CeedBasisApply(basis, blksize, CEED_TRANSPOSE,
369aedaa0e5Sjeremylt                               CEED_EVAL_GRAD, impl->qvecsout[i],
37091703d3fSjeremylt                               impl->evecsout[i]); CeedChk(ierr);
3714a2e7687Sjeremylt         break;
3724ce2993fSjeremylt       case CEED_EVAL_WEIGHT: {
373*c042f62fSJeremy L Thompson         // LCOV_EXCL_START
3744ce2993fSjeremylt         Ceed ceed;
3754ce2993fSjeremylt         ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
3764ce2993fSjeremylt         return CeedError(ceed, 1,
3774a2e7687Sjeremylt                          "CEED_EVAL_WEIGHT cannot be an output evaluation mode");
378*c042f62fSJeremy L Thompson         // LCOV_EXCL_STOP
3794a2e7687Sjeremylt         break; // Should not occur
3804ce2993fSjeremylt       }
3814a2e7687Sjeremylt       case CEED_EVAL_DIV:
3828c91a0c9SJeremy L Thompson         break; // Not implemented
3834a2e7687Sjeremylt       case CEED_EVAL_CURL:
3848c91a0c9SJeremy L Thompson         break; // Not implemented
3854a2e7687Sjeremylt       }
38689c6efa4Sjeremylt     }
38789c6efa4Sjeremylt   }
38889c6efa4Sjeremylt 
38989c6efa4Sjeremylt   // Zero lvecs
39089c6efa4Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
39189c6efa4Sjeremylt     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
39289c6efa4Sjeremylt     if (vec == CEED_VECTOR_ACTIVE) {
39389c6efa4Sjeremylt       if (!impl->add) {
39489c6efa4Sjeremylt         vec = outvec;
39589c6efa4Sjeremylt         ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr);
39689c6efa4Sjeremylt       }
39789c6efa4Sjeremylt     } else {
39889c6efa4Sjeremylt       ierr = CeedVectorSetValue(vec, 0.0); CeedChk(ierr);
39989c6efa4Sjeremylt     }
40089c6efa4Sjeremylt   }
40189c6efa4Sjeremylt   impl->add = false;
40289c6efa4Sjeremylt 
40389c6efa4Sjeremylt   // Output restriction
40489c6efa4Sjeremylt   for (CeedInt i=0; i<numoutputfields; i++) {
40589c6efa4Sjeremylt     // Restore evec
40689c6efa4Sjeremylt     ierr = CeedVectorRestoreArray(impl->evecs[i+impl->numein],
40789c6efa4Sjeremylt                                   &impl->edata[i + numinputfields]); CeedChk(ierr);
408d1bcdac9Sjeremylt     // Get output vector
409d1bcdac9Sjeremylt     ierr = CeedOperatorFieldGetVector(opoutputfields[i], &vec); CeedChk(ierr);
41089c6efa4Sjeremylt     // Active
411d1bcdac9Sjeremylt     if (vec == CEED_VECTOR_ACTIVE)
412d1bcdac9Sjeremylt       vec = outvec;
4134a2e7687Sjeremylt     // Restrict
41489c6efa4Sjeremylt     ierr = CeedOperatorFieldGetLMode(opoutputfields[i], &lmode); CeedChk(ierr);
41589c6efa4Sjeremylt     ierr = CeedElemRestrictionApply(impl->blkrestr[i+impl->numein], CEED_TRANSPOSE,
41689c6efa4Sjeremylt                                     lmode, impl->evecs[i+impl->numein], vec,
41789c6efa4Sjeremylt                                     request); CeedChk(ierr);
41889c6efa4Sjeremylt 
4194a2e7687Sjeremylt   }
4204a2e7687Sjeremylt 
4214a2e7687Sjeremylt   // Restore input arrays
4224a2e7687Sjeremylt   for (CeedInt i=0; i<numinputfields; i++) {
423d1bcdac9Sjeremylt     ierr = CeedQFunctionFieldGetEvalMode(qfinputfields[i], &emode);
424d1bcdac9Sjeremylt     CeedChk(ierr);
4254a2e7687Sjeremylt     if (emode == CEED_EVAL_WEIGHT) { // Skip
4264a2e7687Sjeremylt     } else {
4274a2e7687Sjeremylt       ierr = CeedVectorRestoreArrayRead(impl->evecs[i],
428d1bcdac9Sjeremylt                                         (const CeedScalar **) &impl->edata[i]);
429d1bcdac9Sjeremylt       CeedChk(ierr);
4304a2e7687Sjeremylt     }
4314a2e7687Sjeremylt   }
4324a2e7687Sjeremylt 
4334a2e7687Sjeremylt   return 0;
4344a2e7687Sjeremylt }
4354a2e7687Sjeremylt 
4364a2e7687Sjeremylt int CeedOperatorCreate_Blocked(CeedOperator op) {
4374a2e7687Sjeremylt   int ierr;
438fe2413ffSjeremylt   Ceed ceed;
439fe2413ffSjeremylt   ierr = CeedOperatorGetCeed(op, &ceed); CeedChk(ierr);
4404ce2993fSjeremylt   CeedOperator_Blocked *impl;
4414a2e7687Sjeremylt 
4424a2e7687Sjeremylt   ierr = CeedCalloc(1, &impl); CeedChk(ierr);
443de686571SJeremy L Thompson   ierr = CeedOperatorSetData(op, (void *)&impl); CeedChk(ierr);
444fe2413ffSjeremylt 
445fe2413ffSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Apply",
446fe2413ffSjeremylt                                 CeedOperatorApply_Blocked); CeedChk(ierr);
447fe2413ffSjeremylt   ierr = CeedSetBackendFunction(ceed, "Operator", op, "Destroy",
448fe2413ffSjeremylt                                 CeedOperatorDestroy_Blocked); CeedChk(ierr);
4494a2e7687Sjeremylt   return 0;
4504a2e7687Sjeremylt }
451